{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,11]],"date-time":"2026-07-11T11:12:06Z","timestamp":1783768326889,"version":"3.55.0"},"reference-count":38,"publisher":"Springer Science and Business Media LLC","issue":"4","license":[{"start":{"date-parts":[[2026,5,6]],"date-time":"2026-05-06T00:00:00Z","timestamp":1778025600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,5,6]],"date-time":"2026-05-06T00:00:00Z","timestamp":1778025600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimedia Systems"],"published-print":{"date-parts":[[2026,6]]},"DOI":"10.1007\/s00530-026-02337-3","type":"journal-article","created":{"date-parts":[[2026,5,6]],"date-time":"2026-05-06T05:38:12Z","timestamp":1778045892000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["AAMN: cross-modal fusion network with association alignment matrix for radiological report generation"],"prefix":"10.1007","volume":"32","author":[{"given":"Zonglin","family":"Liang","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xiaodi","family":"Hou","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Fei","family":"Chen","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yijia","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,5,6]]},"reference":[{"key":"2337_CR1","unstructured":"Jing, B., Xie, P., Xing, E.: On the automatic generation of medical imaging reports. arXiv preprint arXiv:1711.08195 (2017)"},{"issue":"1","key":"2337_CR2","doi-asserted-by":"publisher","first-page":"253","DOI":"10.1007\/s11280-022-01013-6","volume":"26","author":"M Li","year":"2023","unstructured":"Li, M., Liu, R., Wang, F., Chang, X., Liang, X.: Auxiliary signal-guided knowledge encoder-decoder for medical report generation. World Wide Web 26(1), 253\u2013270 (2023)","journal-title":"World Wide Web"},{"key":"2337_CR3","doi-asserted-by":"publisher","DOI":"10.1016\/j.jbi.2023.104496","volume":"146","author":"X Hou","year":"2023","unstructured":"Hou, X., Liu, Z., Li, X., Li, X., Sang, S., Zhang, Y.: Mkcl: Medical knowledge with contrastive learning model for radiology report generation. J. Biomed. Inform. 146, 104496 (2023)","journal-title":"J. Biomed. Inform."},{"key":"2337_CR4","doi-asserted-by":"crossref","unstructured":"Chen, Z., Song, Y., Chang, T.-H., Wan, X.: Generating radiology reports via memory-driven transformer. arXiv preprint arXiv:2010.16056 (2020)","DOI":"10.18653\/v1\/2020.emnlp-main.112"},{"key":"2337_CR5","unstructured":"Chen, Z., Shen, Y., Song, Y., Wan, X.: Cross-modal memory networks for radiology report generation. arXiv preprint arXiv:2204.13258 (2022)"},{"key":"2337_CR6","first-page":"277","volume":"37","author":"Y Cao","year":"2023","unstructured":"Cao, Y., Cui, L., Zhang, L., Yu, F., Li, Z., Xu, Y.: MMTN: multi-modal memory transformer network for image-report consistent medical report generation. Proc. AAAI Conf. Artif. Intell. 37, 277\u2013285 (2023)","journal-title":"Proc. AAAI Conf. Artif. Intell."},{"key":"2337_CR7","first-page":"563","volume-title":"European conference on computer vision","author":"J Wang","year":"2022","unstructured":"Wang, J., Bhalerao, A., He, Y.: Cross-modal prototype driven network for radiology report generation. In: European conference on computer vision, pp. 563\u2013579. Springer, Cham (2022)"},{"key":"2337_CR8","doi-asserted-by":"crossref","unstructured":"Jing, B., Wang, Z., Xing, E.: Show, describe and conclude: On exploiting the structure information of chest x-ray reports. arXiv preprint arXiv:2004.12274 (2020)","DOI":"10.18653\/v1\/P19-1657"},{"key":"2337_CR9","unstructured":"Ji, S., Sun, W., Dong, H., Wu, H., Marttinen, P.: A unified review of deep learning for automated medical coding. arXiv preprint arXiv:2201.02797 (2022)"},{"key":"2337_CR10","first-page":"10578","volume-title":"Proceedings of the IEEE\/CVF conference on computer vision and pattern Recognition","author":"M Cornia","year":"2020","unstructured":"Cornia, M., Stefanini, M., Baraldi, L., Cucchiara, R.: Meshed-memory transformer for image captioning. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern Recognition, pp. 10578\u201310587. IEEE, Geneva (2020)"},{"key":"2337_CR11","first-page":"3156","volume-title":"Proceedings of the IEEE conference on computer vision and pattern Recognition","author":"O Vinyals","year":"2015","unstructured":"Vinyals, O., Toshev, A., Bengio, S., Erhan, D.: Show and tell: a neural image caption generator. In: Proceedings of the IEEE conference on computer vision and pattern Recognition, pp. 3156\u20133164. IEEE, Geneva (2015)"},{"key":"2337_CR12","unstructured":"Xu, K.: Show, attend and tell: Neural image caption generation with visual attention. arXiv preprint arXiv:1502.03044 (2015)"},{"key":"2337_CR13","first-page":"6077","volume-title":"Proceedings of the IEEE conference on computer vision and pattern recognition","author":"P Anderson","year":"2018","unstructured":"Anderson, P., He, X., Buehler, C., Teney, D., Johnson, M., Gould, S., Zhang, L.: Bottom-up and top-down attention for image captioning and visual question answering. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp. 6077\u20136086. IEEE, Geneva (2018)"},{"key":"2337_CR14","first-page":"8957","volume":"33","author":"W Wang","year":"2019","unstructured":"Wang, W., Chen, Z., Hu, H.: Hierarchical attention network for image captioning. Proc. AAAI Conf. Artifi. Intell. 33, 8957\u20138964 (2019)","journal-title":"Proc. AAAI Conf. Artifi. Intell."},{"key":"2337_CR15","first-page":"13753","volume-title":"Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition","author":"F Liu","year":"2021","unstructured":"Liu, F., Wu, X., Ge, S., Fan, W., Zou, Y.: Exploring and distilling posterior and prior knowledge for radiology report generation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp. 13753\u201313762. IEEE, Geneva (2021)"},{"key":"2337_CR16","doi-asserted-by":"crossref","unstructured":"You, D., Liu, F., Ge, S., Xie, X., Zhang, J., Wu, X.: Aligntransformer: Hierarchical Alignment of Visual Regions and Disease Tags for Medical Report Generation. In: Medical Image Computing and Computer Assisted Intervention\u2013MICCAI 2021: 24th International Conference, Strasbourg, France, September 27\u2013October 1, 2021, Proceedings, Part III 24, pp. 72\u201382. Springer. (2021)","DOI":"10.1007\/978-3-030-87199-4_7"},{"key":"2337_CR17","first-page":"19809","volume-title":"Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition","author":"Z Huang","year":"2023","unstructured":"Huang, Z., Zhang, X., Zhang, S.: Kiut: knowledge-injected U-transformer for radiology report generation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp. 19809\u201319818. IEEE, Geneva (2023)"},{"key":"2337_CR18","doi-asserted-by":"publisher","DOI":"10.1016\/j.eswa.2023.121260","volume":"237","author":"Y Xue","year":"2024","unstructured":"Xue, Y., Tan, Y., Tan, L., Qin, J., Xiang, X.: Generating radiology reports via auxiliary signal guidance and a memory-driven network. Expert Syst. App. (Mar. Pt.B) 237, 121260 (2024)","journal-title":"Expert Syst. App. (Mar. Pt.B)"},{"issue":"2","key":"2337_CR19","doi-asserted-by":"publisher","first-page":"359","DOI":"10.3390\/s25020359","volume":"25","author":"Y Qu","year":"2025","unstructured":"Qu, Y., Kim, J.: Efficient multi-task training with adaptive feature alignment for universal image segmentation. Sensors (Basel, Switzerland) 25(2), 359 (2025)","journal-title":"Sensors (Basel, Switzerland)"},{"key":"2337_CR20","first-page":"4725","volume":"39","author":"H Li","year":"2025","unstructured":"Li, H., Su, D., Cai, Q., Zhang, Y.: Bsafusion: a bidirectional stepwise feature alignment network for unaligned medical image fusion. Proc. AAAI Conf. Artif. Intell. 39, 4725\u20134733 (2025)","journal-title":"Proc. AAAI Conf. Artif. Intell."},{"key":"2337_CR21","doi-asserted-by":"crossref","unstructured":"Lee, K., Lee, S., Lee, K.M.: Auto-regressive transformation for image alignment. arXiv preprint arXiv:2505.04864 (2025)","DOI":"10.1109\/ICCV51701.2025.01260"},{"key":"2337_CR22","first-page":"770","volume-title":"Proceedings of the IEEE conference on computer vision and pattern recognition","author":"K He","year":"2016","unstructured":"He, K., Zhang, X., Ren, S., Sun, J.: Deep residual learning for image recognition. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp. 770\u2013778. IEEE, Geneva (2016)"},{"key":"2337_CR23","unstructured":"Simonyan, K., Zisserman, A.: Very deep convolutional networks for large-scale image recognition. arXiv preprint arXiv:1409.1556 (2014)"},{"key":"2337_CR24","first-page":"4700","volume-title":"Proceedings of the IEEE conference on computer vision and pattern recognition","author":"G Huang","year":"2017","unstructured":"Huang, G., Liu, Z., Van Der Maaten, L., Weinberger, K.Q.: Densely connected convolutional networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp. 4700\u20134708. IEEE, Geneva (2017)"},{"key":"2337_CR25","unstructured":"Howard, A.G., Zhu, M., Chen, B., Kalenichenko, D., Wang, W., Weyand, T., Andreetto, M., Adam, H.: Mobilenets: efficient convolutional neural networks for mobile vision applications (2017). arXiv preprint arXiv:1704.04861 (2017)"},{"key":"2337_CR26","doi-asserted-by":"publisher","first-page":"365","DOI":"10.1007\/BF01110298","volume":"105","author":"CS Ballantine","year":"1968","unstructured":"Ballantine, C.S.: On the Hadamard product. Math. Z. 105, 365\u2013366 (1968)","journal-title":"Math. Z."},{"key":"2337_CR27","unstructured":"Touvron, H., Lavril, T., Izacard, G., Martinet, X., Lachaux, M.-A., Lacroix, T., Rozi\u00e8re, B., Goyal, N., Hambro, E., Azhar, F., et al.: Llama: Open and efficient foundation language models. arXiv preprint arXiv:2302.13971 (2023)"},{"key":"2337_CR28","unstructured":"Peng, B., Li, C., He, P., Galley, M., Gao, J.: Instruction tuning with gpt-4. arXiv preprint arXiv:2304.03277 (2023)"},{"key":"2337_CR29","doi-asserted-by":"publisher","DOI":"10.1016\/j.media.2023.102918","volume":"89","author":"MA Mazurowski","year":"2023","unstructured":"Mazurowski, M.A., Dong, H., Gu, H., Yang, J., Konz, N., Zhang, Y.: Segment anything model for medical image analysis: an experimental study. Med. Image Anal. 89, 102918 (2023)","journal-title":"Med. Image Anal."},{"key":"2337_CR30","doi-asserted-by":"publisher","DOI":"10.48550\/arXiv.1706.03762","author":"A Vaswani","year":"2017","unstructured":"Vaswani, A., Shazeer, N., Parmar, N., Uszkoreit, J., Jones, L., Gomez, A.N., Kaiser, \u0141, Polosukhin, I.: Attention is all you need. Adv. Neural Info. Process. Syst. (2017). https:\/\/doi.org\/10.48550\/arXiv.1706.03762","journal-title":"Adv. Neural Info. Process. Syst."},{"key":"2337_CR31","first-page":"2097","volume-title":"Proceedings of the IEEE conference on computer vision and pattern recognition","author":"X Wang","year":"2017","unstructured":"Wang, X., Peng, Y., Lu, L., Lu, Z., Bagheri, M., Summers, R.M.: Chestx-ray8: hospital-scale chest X-ray database and benchmarks on weakly-supervised classification and localization of common thorax diseases. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp. 2097\u20132106. IEEE, Geneva (2017)"},{"key":"2337_CR32","doi-asserted-by":"publisher","DOI":"10.48550\/arXiv.1805.08298","author":"Y Li","year":"2018","unstructured":"Li, Y., Liang, X., Hu, Z., Xing, E.P.: Hybrid retrieval-generation reinforced agent for medical image report generation. Adv. Neural. Inf. Process. Syst. (2018). https:\/\/doi.org\/10.48550\/arXiv.1805.08298","journal-title":"Adv. Neural. Inf. Process. Syst."},{"issue":"1","key":"2337_CR33","doi-asserted-by":"publisher","first-page":"317","DOI":"10.1038\/s41597-019-0322-0","volume":"6","author":"AE Johnson","year":"2019","unstructured":"Johnson, A.E., Pollard, T.J., Berkowitz, S.J., Greenbaum, N.R., Lungren, M.P., Deng, C.-Y., Mark, R.G., Horng, S.: Mimic-CXR, a de-identified publicly available database of chest radiographs with free-text reports. Scientif. Data 6(1), 317 (2019)","journal-title":"Scientif. Data"},{"key":"2337_CR34","first-page":"311","volume-title":"Proceedings of the 40th annual meeting of the association for computational linguistics","author":"K Papineni","year":"2002","unstructured":"Papineni, K., Roukos, S., Ward, T., Zhu, W.-J.: Bleu: a method for automatic evaluation of machine translation. In: Proceedings of the 40th annual meeting of the association for computational linguistics, pp. 311\u2013318. ACM International, New York (2002)"},{"key":"2337_CR35","unstructured":"Lin, C.-Y.: Rouge: a package for automatic evaluation of summaries. In: Text summarization branches out. ACL Anthology. pp 74\u201381 (2004)"},{"key":"2337_CR36","unstructured":"Denkowski, M., Lavie, A.: Meteor 1.3: Automatic metric for reliable optimization and evaluation of machine translation systems. In: Proceedings of the sixth workshop on statistical machine translation. pp 85\u201391 (2011)"},{"key":"2337_CR37","unstructured":"Kingma, D.P., Ba, J.: Adam: A method for stochastic optimization. arXiv preprint arXiv:1412.6980 (2014)"},{"key":"2337_CR38","first-page":"375","volume-title":"Proceedings of the IEEE conference on computer vision and pattern recognition","author":"J Lu","year":"2017","unstructured":"Lu, J., Xiong, C., Parikh, D., Socher, R.: Knowing when to look: adaptive attention via a visual sentinel for image captioning. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp. 375\u2013383. IEEE, Geneva (2017)"}],"container-title":["Multimedia Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-026-02337-3.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00530-026-02337-3","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-026-02337-3.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,7,11]],"date-time":"2026-07-11T10:18:44Z","timestamp":1783765124000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00530-026-02337-3"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,5,6]]},"references-count":38,"journal-issue":{"issue":"4","published-print":{"date-parts":[[2026,6]]}},"alternative-id":["2337"],"URL":"https:\/\/doi.org\/10.1007\/s00530-026-02337-3","relation":{},"ISSN":["0942-4962","1432-1882"],"issn-type":[{"value":"0942-4962","type":"print"},{"value":"1432-1882","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,5,6]]},"assertion":[{"value":"18 February 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"4 March 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"6 May 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"253"}}