{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,13]],"date-time":"2026-06-13T00:55:58Z","timestamp":1781312158638,"version":"3.54.1"},"reference-count":39,"publisher":"Elsevier BV","license":[{"start":{"date-parts":[[2026,8,1]],"date-time":"2026-08-01T00:00:00Z","timestamp":1785542400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/tdm\/userlicense\/1.0\/"},{"start":{"date-parts":[[2026,8,1]],"date-time":"2026-08-01T00:00:00Z","timestamp":1785542400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/legal\/tdmrep-license"},{"start":{"date-parts":[[2026,8,1]],"date-time":"2026-08-01T00:00:00Z","timestamp":1785542400000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-017"},{"start":{"date-parts":[[2026,8,1]],"date-time":"2026-08-01T00:00:00Z","timestamp":1785542400000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-037"},{"start":{"date-parts":[[2026,8,1]],"date-time":"2026-08-01T00:00:00Z","timestamp":1785542400000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-012"},{"start":{"date-parts":[[2026,8,1]],"date-time":"2026-08-01T00:00:00Z","timestamp":1785542400000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-029"},{"start":{"date-parts":[[2026,8,1]],"date-time":"2026-08-01T00:00:00Z","timestamp":1785542400000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-004"}],"content-domain":{"domain":["elsevier.com","sciencedirect.com"],"crossmark-restriction":true},"short-container-title":["Expert Systems with Applications"],"published-print":{"date-parts":[[2026,8]]},"DOI":"10.1016\/j.eswa.2026.132359","type":"journal-article","created":{"date-parts":[[2026,4,7]],"date-time":"2026-04-07T15:24:17Z","timestamp":1775575457000},"page":"132359","update-policy":"https:\/\/doi.org\/10.1016\/elsevier_cm_policy","source":"Crossref","is-referenced-by-count":0,"special_numbering":"C","title":["MARM: Medical adaptive reasoning model"],"prefix":"10.1016","volume":"322","author":[{"ORCID":"https:\/\/orcid.org\/0009-0009-7435-9108","authenticated-orcid":false,"given":"Ruihui","family":"Hou","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0005-5807-9131","authenticated-orcid":false,"given":"Ziyue","family":"Huai","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0009-5785-2750","authenticated-orcid":false,"given":"Yinan","family":"Wu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-3546-8338","authenticated-orcid":false,"given":"Tong","family":"Ruan","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"78","reference":[{"key":"10.1016\/j.eswa.2026.132359_bib0001","unstructured":"Alomrani, M. A., Zhang, Y., Li, D., Sun, Q., Pal, S., Zhang, Z., Hu, Y., Ajwani, R. D., Valkanas, A., Karimi, R. et al. (2025). Reasoning on a budget: A survey of adaptive and controllable test-time compute in LLMs. arXiv preprint arXiv: 2507.02076."},{"key":"10.1016\/j.eswa.2026.132359_bib0002","unstructured":"Chen, J., Cai, Z., Ji, K., Wang, X., Liu, W., Wang, R., Hou, J., & Wang, B. (2024a). Huatuogpt-o1, towards medical complex reasoning with LLMs. arXiv preprint arXiv: 2412.18925."},{"key":"10.1016\/j.eswa.2026.132359_sbref0004","series-title":"Proceedings of the 31st international conference on computational linguistics","first-page":"10183","article-title":"AI hospital: Benchmarking large language models in a multi-agent medical interaction simulator","author":"Fan","year":"2025"},{"key":"10.1016\/j.eswa.2026.132359_bib0004","unstructured":"Guo, D., Yang, D., Zhang, H., Song, J., Zhang, R., Xu, R., Zhu, Q., Ma, S., Wang, P., Bi, X. et al. (2025). Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning. arXiv preprint arXiv: 2501.12948."},{"key":"10.1016\/j.eswa.2026.132359_sbref0006","doi-asserted-by":"crossref","DOI":"10.1016\/j.eswa.2025.129024","article-title":"Act-LLM: A whole-process chain for character-centric role-playing with large language models","volume":"296","author":"Han","year":"2026","journal-title":"Expert Systems with Applications"},{"key":"10.1016\/j.eswa.2026.132359_sbref0007","doi-asserted-by":"crossref","DOI":"10.1016\/j.knosys.2025.114524","article-title":"Msdiagnosis: A benchmark and framework for evaluating large language models in multi-step clinical diagnosis","volume":"330","author":"Hou","year":"2025","journal-title":"Knowledge-Based Systems"},{"key":"10.1016\/j.eswa.2026.132359_sbref0008","doi-asserted-by":"crossref","DOI":"10.1016\/j.eswa.2026.131806","article-title":"Cdaflow: Enhancing llm clinical decision-making through agentic workflow","volume":"316","author":"Hou","year":"2026","journal-title":"Expert Systems with Applications"},{"key":"10.1016\/j.eswa.2026.132359_bib0008","unstructured":"Huang, K., Altosaar, J., & Ranganath, R. (2019). Clinicalbert: Modeling clinical notes and predicting hospital readmission. arXiv preprint arXiv: 1904.05342."},{"issue":"14","key":"10.1016\/j.eswa.2026.132359_bib0009","doi-asserted-by":"crossref","first-page":"6421","DOI":"10.3390\/app11146421","article-title":"What disease does this patient have? A large-scale open domain question answering dataset from medical exams","volume":"11","author":"Jin","year":"2021","journal-title":"Applied Sciences"},{"key":"10.1016\/j.eswa.2026.132359_bib0010","unstructured":"Jin, Q., Dhingra, B., Liu, Z., Cohen, W. W., & Lu, X. (2019). Pubmedqa: A dataset for biomedical research question answering. arXiv preprint arXiv: 1909.06146."},{"key":"10.1016\/j.eswa.2026.132359_sbref0012","doi-asserted-by":"crossref","DOI":"10.1016\/j.eswa.2025.128736","article-title":"Llm-augmented hierarchical reinforcement learning for human-like decision-making of autonomous driving","volume":"294","author":"Li","year":"2025","journal-title":"Expert Systems with Applications"},{"key":"10.1016\/j.eswa.2026.132359_sbref0013","series-title":"The twelfth international conference on learning representations","article-title":"Let\u2019s verify step by step","author":"Lightman","year":"2024"},{"key":"10.1016\/j.eswa.2026.132359_bib0013","series-title":"Conference on health, inference, and learning","first-page":"248","article-title":"Medmcqa: A large-scale multi-subject multi-choice dataset for medical domain question answering","author":"Pal","year":"2022"},{"key":"10.1016\/j.eswa.2026.132359_bib0014","doi-asserted-by":"crossref","unstructured":"Qiu, P., Wu, C., Liu, S., Zhao, W., Chen, Z., Gu, H., Peng, C., Zhang, Y., Wang, Y., & Xie, W. (2025). Quantifying the reasoning abilities of llms on real-world clinical cases. arXiv preprint arXiv: 2503.04691.","DOI":"10.1038\/s41467-025-64769-1"},{"key":"10.1016\/j.eswa.2026.132359_bib0015","unstructured":"Rui, S., Chen, K., Ma, W., & Wang, X. (2025). Adathink-med: Medical adaptive thinking with uncertainty-guided length calibration. arXiv preprint arXiv: 2509.24560."},{"key":"10.1016\/j.eswa.2026.132359_bib0016","unstructured":"Saab, K., Tu, T., Weng, W.-H., Tanno, R., Stutz, D., Wulczyn, E., Zhang, F., Strother, T., Park, C., Vedadi, E. et al. (2024). Capabilities of gemini models in medicine. arXiv preprint arXiv: 2404.18416."},{"key":"10.1016\/j.eswa.2026.132359_bib0017","unstructured":"Schmidgall, S., Ziaei, R., Harris, C., Reis, E., Jopling, J., & Moor, M. (2024). Agentclinic: a multimodal agent benchmark to evaluate AI in simulated clinical environments. arXiv preprint arXiv: 2405.07960."},{"key":"10.1016\/j.eswa.2026.132359_bib0018","unstructured":"Schulman, J., Wolski, F., Dhariwal, P., Radford, A., & Klimov, O. (2017). Proximal policy optimization algorithms. arXiv preprint arXiv: 1707.06347."},{"key":"10.1016\/j.eswa.2026.132359_bib0019","unstructured":"Sellergren, A., Kazemzadeh, S., Jaroensri, T., Kiraly, A., Traverse, M., Kohlberger, T., Xu, S., Jamil, F., Hughes, C., Lau, C. et al. (2025). Medgemma technical report. arXiv preprint arXiv: 2507.05201."},{"key":"10.1016\/j.eswa.2026.132359_sbref0021","series-title":"The thirteenth international conference on learning representations","article-title":"Rewarding progress: Scaling automated process verifiers for LLM reasoning","author":"Setlur","year":"2025"},{"key":"10.1016\/j.eswa.2026.132359_bib0021","unstructured":"Shao, Z., Wang, P., Zhu, Q., Xu, R., Song, J., Bi, X., Zhang, H., Zhang, M., Li, Y. K., Wu, Y. et al. (2024). Deepseekmath: Pushing the limits of mathematical reasoning in open language models. arXiv preprint arXiv: 2402.03300."},{"key":"10.1016\/j.eswa.2026.132359_bib0022","unstructured":"Tu, S., Lin, J., Zhang, Q., Tian, X., Li, L., Lan, X., & Zhao, D. (2025). Learning when to think: Shaping adaptive reasoning in r1-style models via multi-stage RL. arXiv preprint arXiv: 2505.10832."},{"key":"10.1016\/j.eswa.2026.132359_bib0023","first-page":"95266","article-title":"Mmlu-pro: A more robust and challenging multi-task language understanding benchmark","volume":"37","author":"Wang","year":"2024","journal-title":"Advances in Neural Information Processing Systems"},{"key":"10.1016\/j.eswa.2026.132359_bib0024","first-page":"24824","article-title":"Chain-of-thought prompting elicits reasoning in large language models","volume":"35","author":"Wei","year":"2022","journal-title":"Advances in Neural Information Processing Systems"},{"key":"10.1016\/j.eswa.2026.132359_bib0025","unstructured":"Wu, K., Wu, E., Thapa, R., Wei, K., Zhang, A., Suresh, A., Tao, J. J., Sun, M. W., Lozano, A., & Zou, J. (2025a). Medcasereasoning: Evaluating and learning diagnostic reasoning from clinical case reports. arXiv preprint arXiv: 2505.11733."},{"key":"10.1016\/j.eswa.2026.132359_bib0026","unstructured":"Wu, S., Xie, J., Zhang, Y., Chen, A., Zhang, K., Su, Y., & Xiao, Y. (2025b). Arm: Adaptive reasoning model. arXiv preprint arXiv: 2505.20258."},{"key":"10.1016\/j.eswa.2026.132359_bib0027","first-page":"87621","article-title":"Medjourney: Benchmark and evaluation of large language models over patient clinical journey","volume":"37","author":"Wu","year":"2024","journal-title":"Advances in Neural Information Processing Systems"},{"key":"10.1016\/j.eswa.2026.132359_bib0028","unstructured":"Xu, S., Zhou, Y., Liu, Z., Wu, Z., Zhong, T., Zhao, H., Li, Y., Jiang, H., Pan, Y., Chen, J. et al. (2024). Towards next-generation medical agent: How o1 is reshaping decision-making in medical scenarios. arXiv preprint arXiv: 2411.14461."},{"key":"10.1016\/j.eswa.2026.132359_sbref0030","doi-asserted-by":"crossref","DOI":"10.1016\/j.eswa.2025.127582","article-title":"Staf-llm: A scalable and task-adaptive fine-tuning framework for large language models in medical domain","volume":"281","author":"Xu","year":"2025","journal-title":"Expert Systems with Applications"},{"key":"10.1016\/j.eswa.2026.132359_bib0030","unstructured":"Yan, W., Liu, H., Wu, T., Chen, Q., Wang, W., Chai, H., Wang, J., Zhao, W., Zhang, Y., Zhang, R. et al. (2024). Clinicallab: Aligning agents for multi-departmental clinical diagnostics in the real world. arXiv preprint arXiv: 2406.13890."},{"key":"10.1016\/j.eswa.2026.132359_bib0031","unstructured":"Yang, A., Li, A., Yang, B., Zhang, B., Hui, B., Zheng, B., Yu, B., Gao, C., Huang, C., Lv, C. et al. (2025a). Qwen3 technical report. arXiv preprint arXiv: 2505.09388."},{"key":"10.1016\/j.eswa.2026.132359_bib0032","unstructured":"Yang, A., Yang, B., Zhang, B., Hui, B., Zheng, B., Yu, B., Li, C., Liu, D., Huang, F., Wei, H., Lin, H., Yang, J., Tu, J., Zhang, J., Yang, J., Yang, J., Zhou, J., Lin, J., Dang, K., Lu, K., Bao, K., Yang, K., Yu, L., Li, M., Xue, M., Zhang, P., Zhu, Q., Men, R., Lin, R., Li, T., Xia, T., Ren, X., Ren, X., Fan, Y., Su, Y., Zhang, Y., Wan, Y., Liu, Y., Cui, Z., Zhang, Z., & Qiu, Z. (2024). Qwen2.5 technical report. arXiv preprint arXiv: 2412.15115."},{"key":"10.1016\/j.eswa.2026.132359_bib0033","unstructured":"Yang, Z., Qian, J., Peng, Z., Zhang, H., & Huang, Z.-A. (2025b). Med-REFL: Medical reasoning enhancement via self-corrected fine-grained reflection. arXiv preprint arXiv: 2506.13793."},{"key":"10.1016\/j.eswa.2026.132359_bib0034","doi-asserted-by":"crossref","unstructured":"Yun, J., Sohn, J., Park, J., Kim, H., Tang, X., Shao, Y., Koo, Y., Ko, M., Chen, Q., Gerstein, M. et al. (2025a). Med-PRM: Medical reasoning models with stepwise, guideline-verified process rewards. arXiv preprint arXiv: 2506.11474.","DOI":"10.18653\/v1\/2025.emnlp-main.837"},{"key":"10.1016\/j.eswa.2026.132359_bib0035","doi-asserted-by":"crossref","unstructured":"Yun, J., Sohn, J., Park, J., Kim, H., Tang, X., Shao, Y., Koo, Y., Ko, M., Chen, Q., Gerstein, M. et al. (2025b). Med-PRM: Medical reasoning models with stepwise, guideline-verified process rewards. arXiv preprint arXiv: 2506.11474.","DOI":"10.18653\/v1\/2025.emnlp-main.837"},{"key":"10.1016\/j.eswa.2026.132359_bib0036","unstructured":"Zhang, R., Xiao, C., & Cao, Y. (2025a). Long or short cot? Investigating instance-level switch of large reasoning models. arXiv preprint arXiv: 2506.04182."},{"key":"10.1016\/j.eswa.2026.132359_bib0037","unstructured":"Zhang, X., Wang, Y., Feng, Z., Chen, R., Zhou, Z., Zhang, Y., Xu, H., Wu, J., & Liu, Z. (2025b). Med-u1: Incentivizing unified medical reasoning in LLMs via large-scale reinforcement learning. arXiv preprint arXiv: 2506.12307."},{"key":"10.1016\/j.eswa.2026.132359_sbref0039","doi-asserted-by":"crossref","DOI":"10.1016\/j.eswa.2025.129423","article-title":"Tape: A multi-agent framework for task-adaptive planning and execution in resource-constrained environments","volume":"297","author":"Zhang","year":"2026","journal-title":"Expert Systems with Applications"},{"key":"10.1016\/j.eswa.2026.132359_bib0039","unstructured":"Zhao, Z., Jin, Q., Chen, F., Peng, T., & Yu, S. (2022). PMC-patients: A large-scale dataset of patient summaries and relations for benchmarking retrieval-based clinical decision support systems. arXiv preprint arXiv: 2202.13876."}],"container-title":["Expert Systems with Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0957417426012728?httpAccept=text\/xml","content-type":"text\/xml","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0957417426012728?httpAccept=text\/plain","content-type":"text\/plain","content-version":"vor","intended-application":"text-mining"}],"deposited":{"date-parts":[[2026,6,13]],"date-time":"2026-06-13T00:01:54Z","timestamp":1781308914000},"score":1,"resource":{"primary":{"URL":"https:\/\/linkinghub.elsevier.com\/retrieve\/pii\/S0957417426012728"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,8]]},"references-count":39,"alternative-id":["S0957417426012728"],"URL":"https:\/\/doi.org\/10.1016\/j.eswa.2026.132359","relation":{},"ISSN":["0957-4174"],"issn-type":[{"value":"0957-4174","type":"print"}],"subject":[],"published":{"date-parts":[[2026,8]]},"assertion":[{"value":"Elsevier","name":"publisher","label":"This article is maintained by"},{"value":"MARM: Medical adaptive reasoning model","name":"articletitle","label":"Article Title"},{"value":"Expert Systems with Applications","name":"journaltitle","label":"Journal Title"},{"value":"https:\/\/doi.org\/10.1016\/j.eswa.2026.132359","name":"articlelink","label":"CrossRef DOI link to publisher maintained version"},{"value":"article","name":"content_type","label":"Content Type"},{"value":"\u00a9 2026 Elsevier Ltd. All rights are reserved, including those for text and data mining, AI training, and similar technologies.","name":"copyright","label":"Copyright"}],"article-number":"132359"}}