{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,4,5]],"date-time":"2025-04-05T04:19:02Z","timestamp":1743826742924,"version":"3.40.3"},"publisher-location":"Cham","reference-count":36,"publisher":"Springer Nature Switzerland","isbn-type":[{"value":"9783031887109","type":"print"},{"value":"9783031887116","type":"electronic"}],"license":[{"start":{"date-parts":[[2025,1,1]],"date-time":"2025-01-01T00:00:00Z","timestamp":1735689600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,1,1]],"date-time":"2025-01-01T00:00:00Z","timestamp":1735689600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2025]]},"DOI":"10.1007\/978-3-031-88711-6_10","type":"book-chapter","created":{"date-parts":[[2025,4,4]],"date-time":"2025-04-04T17:15:06Z","timestamp":1743786906000},"page":"153-168","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Evaluating LLM Abilities to\u00a0Understand Tabular Electronic Health Records: A Comprehensive Study of\u00a0Patient Data Extraction and\u00a0Retrieval"],"prefix":"10.1007","author":[{"ORCID":"https:\/\/orcid.org\/0000-0001-6243-0864","authenticated-orcid":false,"given":"Jes\u00fas","family":"Lov\u00f3n-Melgarejo","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Martin","family":"Mouysset","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Jo","family":"Oleiwan","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-8852-5797","authenticated-orcid":false,"given":"Jose G.","family":"Moreno","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-5018-0108","authenticated-orcid":false,"given":"Christine","family":"Damase-Michel","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-3615-8032","authenticated-orcid":false,"given":"Lynda","family":"Tamine","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,4,4]]},"reference":[{"key":"10_CR1","unstructured":"Chen, N., et al.: Bridge the gap between language models and tabular understanding. arXiv preprint arXiv:2302.09302 (2023)"},{"key":"10_CR2","doi-asserted-by":"crossref","unstructured":"Chen, W.: Large language models are few(1)-shot table reasoners. In: Findings of EACL (2023)","DOI":"10.18653\/v1\/2023.findings-eacl.83"},{"key":"10_CR3","doi-asserted-by":"publisher","unstructured":"Deng, X., Bashlovkina, V., Han, F., Baumgartner, S., Bendersky, M.: What do LLMs know about financial markets? A case study on reddit market sentiment analysis. In: Companion Proceedings of the ACM Web Conference 2023, pp. 107\u2013110. WWW \u201923 Companion, Association for Computing Machinery, New York, NY, USA (2023). https:\/\/doi.org\/10.1145\/3543873.3587324","DOI":"10.1145\/3543873.3587324"},{"issue":"1","key":"10_CR4","doi-asserted-by":"publisher","first-page":"33","DOI":"10.1145\/3542700.3542709","volume":"51","author":"X Deng","year":"2022","unstructured":"Deng, X., Sun, H., Lees, A., Wu, Y., Yu, C.: TURL: table understanding through representation learning. ACM SIGMOD Rec. 51(1), 33\u201340 (2022)","journal-title":"ACM SIGMOD Rec."},{"key":"10_CR5","unstructured":"Fang, X., et al.: Large language models on tabular data\u2013a survey. arXiv preprint arXiv:2402.17944 (2024)"},{"key":"10_CR6","doi-asserted-by":"publisher","unstructured":"Gong, H., et al.: TableGPT: Few-shot table-to-text generation with table structure reconstruction and content matching. In: Scott, D., Bel, N., Zong, C. (eds.) Proceedings of the 28th International Conference on Computational Linguistics, pp. 1978\u20131988. International Committee on Computational Linguistics, Barcelona, Spain (Online) (2020). https:\/\/doi.org\/10.18653\/v1\/2020.coling-main.179","DOI":"10.18653\/v1\/2020.coling-main.179"},{"key":"10_CR7","doi-asserted-by":"publisher","first-page":"611","DOI":"10.1007\/978-3-319-76941-7_52","volume-title":"Advances in Information Retrieval","author":"T Haug","year":"2018","unstructured":"Haug, T., Ganea, O.E., Grnarova, P.: Neural multi-step reasoning for question answering on semi-structured tables. In: Pasi, G., Piwowarski, B., Azzopardi, L., Hanbury, A. (eds.) Advances in Information Retrieval, pp. 611\u2013617. Springer International Publishing, Cham (2018)"},{"key":"10_CR8","unstructured":"Hegselmann, S., Buendia, A., Lang, H., Agrawal, M., Jiang, X., Sontag, D.A.: TabLLM: Few-shot classification of tabular data with large language models. In: AISTATG. vol. abs\/2210.10723 (2022)"},{"key":"10_CR9","doi-asserted-by":"crossref","unstructured":"Hou, Y., et al.: Large language models are zero-shot rankers for recommender systems. In: European Conference on Information Retrieval, pp. 364\u2013381 (2024)","DOI":"10.1007\/978-3-031-56060-6_24"},{"key":"10_CR10","doi-asserted-by":"crossref","unstructured":"Johnson, A., et al.: MIMIC-III, a freely accessible critical care database. Sci Data 3 (2016)","DOI":"10.1038\/sdata.2016.35"},{"key":"10_CR11","doi-asserted-by":"crossref","unstructured":"Lin, J., Ma, X., Lin, S.C., Yang, J.H., Pradeep, R., Nogueira, R.: Pyserini: a Python toolkit for reproducible information retrieval research with sparse and dense representations. In: Proceedings of the 44th Annual International ACM SIGIR Conference on Research and Development in Information Retrieval (SIGIR 2021), pp. 2356\u20132362 (2021)","DOI":"10.1145\/3404835.3463238"},{"key":"10_CR12","doi-asserted-by":"crossref","unstructured":"Lin, S.C., et al.: How to train your dragon: diverse augmentation towards generalizable dense retrieval. In: Findings of the Association for Computational Linguistics: EMNLP 2023, pp. 6385\u20136400 (2023)","DOI":"10.18653\/v1\/2023.findings-emnlp.423"},{"key":"10_CR13","unstructured":"Lin, X.V., et\u00a0al.: RA-DIT: retrieval-augmented dual instruction tuning. In: The Twelfth International Conference on Learning Representations (2023)"},{"key":"10_CR14","doi-asserted-by":"publisher","unstructured":"Liu, S.C., et al.: JarviX: a LLM no code platform for tabular data analysis and optimization. In: Wang, M., Zitouni, I. (eds.) Proceedings of the 2023 Conference on Empirical Methods in Natural Language Processing: Industry Track, pp. 622\u2013630. Association for Computational Linguistics, Singapore (2023). https:\/\/doi.org\/10.18653\/v1\/2023.emnlp-industry.59","DOI":"10.18653\/v1\/2023.emnlp-industry.59"},{"key":"10_CR15","doi-asserted-by":"crossref","unstructured":"Liu, T., Wang, K., Sha, L., Chang, B., Sui, Z.: Table-to-text generation by structure-aware seq2seq learning. In: AAAI Conference on Artificial Intelligence (2017). https:\/\/api.semanticscholar.org\/CorpusID:7672408","DOI":"10.1609\/aaai.v32i1.11925"},{"key":"10_CR16","unstructured":"Lovon-Melgarejo, J., Ben-Haddi, T., Di\u00a0Scala, J., Moreno, J.G., Tamine, L.: Revisiting the MIMIC-IV benchmark: Experiments using language models for electronic health records. In: Demner-Fushman, D., Ananiadou, S., Thompson, P., Ondov, B. (eds.) Proceedings of the First Workshop on Patient-Oriented Language Processing (CL4Health) @ LREC-COLING 2024, pp. 189\u2013196. ELRA and ICCL, Torino, Italia (2024). https:\/\/aclanthology.org\/2024.cl4health-1.23"},{"key":"10_CR17","doi-asserted-by":"publisher","unstructured":"Puduppully, R., Dong, L., Lapata, M.: Data-to-text generation with entity modeling. In: Korhonen, A., Traum, D., M\u00e0rquez, L. (eds.) Proceedings of the 57th Annual Meeting of the Association for Computational Linguistics, pp. 2023\u20132035. Association for Computational Linguistics, Florence, Italy (2019). https:\/\/doi.org\/10.18653\/v1\/P19-1195","DOI":"10.18653\/v1\/P19-1195"},{"key":"10_CR18","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"65","DOI":"10.1007\/978-3-030-45439-5_5","volume-title":"Advances in Information Retrieval","author":"C Rebuffel","year":"2020","unstructured":"Rebuffel, C., Soulier, L., Scoutheeten, G., Gallinari, P.: A hierarchical model for data-to-text generation. In: Jose, J.M., et al. (eds.) ECIR 2020. LNCS, vol. 12035, pp. 65\u201380. Springer, Cham (2020). https:\/\/doi.org\/10.1007\/978-3-030-45439-5_5"},{"key":"10_CR19","doi-asserted-by":"publisher","unstructured":"Rubin, O., Herzig, J., Berant, J.: Learning to retrieve prompts for in-context learning. In: Carpuat, M., de\u00a0Marneffe, M.C., Meza\u00a0Ruiz, I.V. (eds.) Proceedings of the 2022 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies, pp. 2655\u20132671. Association for Computational Linguistics, Seattle, United States (2022). https:\/\/doi.org\/10.18653\/v1\/2022.naacl-main.191","DOI":"10.18653\/v1\/2022.naacl-main.191"},{"key":"10_CR20","unstructured":"Sarkar, S., Lausen, L.: Testing the limits of unified sequence to sequence LLM pretraining on diverse table data tasks. In: NeurIPS 2023 Second Table Representation Learning Workshop (2023)"},{"key":"10_CR21","unstructured":"Singha, A., Cambronero, J., Gulwani, S., Le, V., Parnin, C.: Tabular representation, noisy operators, and impacts on table structure understanding tasks in LLMs. In: Table Representation Learning Workshop at NeurIPS 2023 (2023)"},{"key":"10_CR22","unstructured":"Slack, D., Singh, S.: Tablet: Learning from instructions for tabular data. arXiv preprint arXiv:2304.13188 (2023)"},{"key":"10_CR23","doi-asserted-by":"publisher","DOI":"10.1016\/j.jbi.2020.103637","volume":"113","author":"E Steinberg","year":"2021","unstructured":"Steinberg, E., Jung, K., Fries, J.A., Corbin, C.K., Pfohl, S.R., Shah, N.H.: Language models are an effective representation learning technique for electronic health record data. J. Biomed. Inform. 113, 103637 (2021). https:\/\/doi.org\/10.1016\/j.jbi.2020.103637","journal-title":"J. Biomed. Inform."},{"key":"10_CR24","doi-asserted-by":"publisher","unstructured":"Sui, Y., Zhou, M., Zhou, M., Han, S., Zhang, D.: Table meets LLM: can large language models understand structured table data? A benchmark and empirical study. In: Proceedings of the 17th ACM International Conference on Web Search and Data Mining, pp. 645\u2013654. WSDM \u201924, Association for Computing Machinery, New York, NY, USA (2024). https:\/\/doi.org\/10.1145\/3616855.3635752","DOI":"10.1145\/3616855.3635752"},{"key":"10_CR25","doi-asserted-by":"crossref","unstructured":"Sui, Y., et al.: TAP4LLM: table provider on sampling, augmenting, and packing semi-structured data for large language model reasoning. In: Findings of the Association for Computational Linguistics: EMNLP 2024 (2024)","DOI":"10.18653\/v1\/2024.findings-emnlp.603"},{"key":"10_CR26","doi-asserted-by":"publisher","unstructured":"Trabelsi, M., Chen, Z., Zhang, S., Davison, B.D., Heflin, J.: StruBERT: structure-aware BERT for table search and matching. In: Proceedings of the ACM Web Conference 2022, pp. 442\u2013451. WWW \u201922, Association for Computing Machinery, New York, NY, USA (2022). https:\/\/doi.org\/10.1145\/3485447.3511972","DOI":"10.1145\/3485447.3511972"},{"key":"10_CR27","doi-asserted-by":"publisher","unstructured":"Wang, P., Shi, T., Reddy, C.K.: Text-to-SQL generation for question answering on electronic medical records. In: Proceedings of The Web Conference 2020, pp. 350\u2013361. WWW \u201920, Association for Computing Machinery, New York, NY, USA (2020). https:\/\/doi.org\/10.1145\/3366423.3380120","DOI":"10.1145\/3366423.3380120"},{"key":"10_CR28","doi-asserted-by":"publisher","unstructured":"Wang, Y., et al.: Self-instruct: aligning language models with self-generated instructions. In: Rogers, A., Boyd-Graber, J., Okazaki, N. (eds.) Proceedings of the 61st Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers), pp. 13484\u201313508. Association for Computational Linguistics, Toronto, Canada (2023). https:\/\/doi.org\/10.18653\/v1\/2023.acl-long.754","DOI":"10.18653\/v1\/2023.acl-long.754"},{"key":"10_CR29","doi-asserted-by":"publisher","unstructured":"Wang, Z., et al.: TUTA: tree-based transformers for generally structured table pre-training. In: Proceedings of the 27th ACM SIGKDD Conference on Knowledge Discovery & Data Mining, pp. 1780\u20131790. KDD \u201921, Association for Computing Machinery, New York, NY, USA (2021). https:\/\/doi.org\/10.1145\/3447548.3467434","DOI":"10.1145\/3447548.3467434"},{"key":"10_CR30","unstructured":"Wei, J., et al.: Emergent abilities of large language models. arXiv preprint arXiv:2206.07682 (2022)"},{"key":"10_CR31","doi-asserted-by":"crossref","unstructured":"Yang, H., Zhang, Y., Xu, J., Lu, H., Heng, P.A., Lam, W.: Unveiling the generalization power of fine-tuned large language models. In: Proceedings of the 2024 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies (Volume 1: Long Papers), pp. 884\u2013899 (2024)","DOI":"10.18653\/v1\/2024.naacl-long.51"},{"key":"10_CR32","doi-asserted-by":"publisher","unstructured":"Ye, Y., Hui, B., Yang, M., Li, B., Huang, F., Li, Y.: Large language models are versatile decomposers: Decomposing evidence and questions for table-based reasoning. In: Proceedings of the 46th International ACM SIGIR Conference on Research and Development in Information Retrieval, pp. 174\u2013184. SIGIR \u201923, Association for Computing Machinery, New York, NY, USA (2023). https:\/\/doi.org\/10.1145\/3539618.3591708","DOI":"10.1145\/3539618.3591708"},{"key":"10_CR33","unstructured":"Yu, W., et al.: Generate rather than retrieve: Large language models are strong context generators. In: The Eleventh International Conference on Learning Representations (2023)"},{"key":"10_CR34","doi-asserted-by":"publisher","unstructured":"Zhang, Y., Feng, S., Tan, C.: Active example selection for in-context learning. In: Goldberg, Y., Kozareva, Z., Zhang, Y. (eds.) Proceedings of the 2022 Conference on Empirical Methods in Natural Language Processing, pp. 9134\u20139148. Association for Computational Linguistics, Abu Dhabi, United Arab Emirates (2022). https:\/\/doi.org\/10.18653\/v1\/2022.emnlp-main.622","DOI":"10.18653\/v1\/2022.emnlp-main.622"},{"key":"10_CR35","doi-asserted-by":"crossref","unstructured":"Zhao, B., et al.: Large language models are complex table parsers. In: Conference on Empirical Methods in Natural Language Processing (2023)","DOI":"10.18653\/v1\/2023.emnlp-main.914"},{"key":"10_CR36","doi-asserted-by":"crossref","unstructured":"Zhuang, S., Zhuang, H., Koopman, B., Zuccon, G.: A setwise approach for effective and highly efficient zero-shot ranking with large language models. In: Proceedings of the 47th International ACM SIGIR Conference on Research and Development in Information Retrieval, pp. 38\u201347 (2024)","DOI":"10.1145\/3626772.3657813"}],"container-title":["Lecture Notes in Computer Science","Advances in Information Retrieval"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/978-3-031-88711-6_10","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,4,4]],"date-time":"2025-04-04T17:15:21Z","timestamp":1743786921000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/978-3-031-88711-6_10"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025]]},"ISBN":["9783031887109","9783031887116"],"references-count":36,"URL":"https:\/\/doi.org\/10.1007\/978-3-031-88711-6_10","relation":{},"ISSN":["0302-9743","1611-3349"],"issn-type":[{"value":"0302-9743","type":"print"},{"value":"1611-3349","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025]]},"assertion":[{"value":"4 April 2025","order":1,"name":"first_online","label":"First Online","group":{"name":"ChapterHistory","label":"Chapter History"}},{"value":"The authors have no competing interests to declare that are relevant to the content of this article.","order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Disclosure of Interests"}},{"value":"ECIR","order":1,"name":"conference_acronym","label":"Conference Acronym","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"European Conference on Information Retrieval","order":2,"name":"conference_name","label":"Conference Name","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Lucca","order":3,"name":"conference_city","label":"Conference City","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Italy","order":4,"name":"conference_country","label":"Conference Country","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"2025","order":5,"name":"conference_year","label":"Conference Year","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"7 April 2025","order":7,"name":"conference_start_date","label":"Conference Start Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"11 April 2025","order":8,"name":"conference_end_date","label":"Conference End Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"47","order":9,"name":"conference_number","label":"Conference Number","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"ecir2025","order":10,"name":"conference_id","label":"Conference ID","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"https:\/\/ecir2025.eu\/","order":11,"name":"conference_url","label":"Conference URL","group":{"name":"ConferenceInfo","label":"Conference Information"}}]}}