{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,7]],"date-time":"2026-07-07T15:47:21Z","timestamp":1783439241795,"version":"3.54.6"},"publisher-location":"Cham","reference-count":28,"publisher":"Springer Nature Switzerland","isbn-type":[{"value":"9783031723773","type":"print"},{"value":"9783031723780","type":"electronic"}],"license":[{"start":{"date-parts":[[2024,1,1]],"date-time":"2024-01-01T00:00:00Z","timestamp":1704067200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,1,1]],"date-time":"2024-01-01T00:00:00Z","timestamp":1704067200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2024]]},"DOI":"10.1007\/978-3-031-72378-0_67","type":"book-chapter","created":{"date-parts":[[2024,10,2]],"date-time":"2024-10-02T07:02:53Z","timestamp":1727852573000},"page":"722-732","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":17,"title":["MM-Retinal: Knowledge-Enhanced Foundational Pretraining with\u00a0Fundus Image-Text Expertise"],"prefix":"10.1007","author":[{"given":"Ruiqi","family":"Wu","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Chenran","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jianle","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yi","family":"Zhou","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Tao","family":"Zhou","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Huazhu","family":"Fu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2024,10,3]]},"reference":[{"key":"67_CR1","unstructured":"Alsentzer, E., et al.: Publicly available clinical bert embeddings. arXiv preprint arXiv:1904.03323 (2019)"},{"issue":"3","key":"67_CR2","doi-asserted-by":"publisher","first-page":"509","DOI":"10.1177\/193229680900300315","volume":"3","author":"J Cuadros","year":"2009","unstructured":"Cuadros, J., Bresnick, G.: EyePACS: an adaptable telemedicine system for diabetic retinopathy screening. J. Diabetes Sci. Technol. 3(3), 509\u2013516 (2009)","journal-title":"J. Diabetes Sci. Technol."},{"issue":"3","key":"67_CR3","doi-asserted-by":"publisher","first-page":"231","DOI":"10.5566\/ias.1155","volume":"33","author":"E Decenci\u00e8re","year":"2014","unstructured":"Decenci\u00e8re, E., et al.: Feedback on a publicly distributed image database: the messidor database. Image Anal. Stereol. 33(3), 231\u2013234 (2014)","journal-title":"Image Anal. Stereol."},{"key":"67_CR4","doi-asserted-by":"publisher","DOI":"10.1016\/j.bspc.2023.104810","volume":"84","author":"S Diao","year":"2023","unstructured":"Diao, S., et al.: Classification and segmentation of oct images for age-related macular degeneration based on dual guidance networks. Biomed. Signal Process. Control 84, 104810 (2023)","journal-title":"Biomed. Signal Process. Control"},{"key":"67_CR5","doi-asserted-by":"publisher","unstructured":"Fu, H., et al.: Palm: pathologic myopia challenge (2019). https:\/\/doi.org\/10.21227\/55pk-8z03","DOI":"10.21227\/55pk-8z03"},{"key":"67_CR6","doi-asserted-by":"publisher","unstructured":"Fu, H., et al.: Adam: automatic detection challenge on age-related macular degeneration (2020). https:\/\/doi.org\/10.21227\/dt4f-rt59","DOI":"10.21227\/dt4f-rt59"},{"issue":"2","key":"67_CR7","doi-asserted-by":"publisher","first-page":"581","DOI":"10.1007\/s11263-023-01891-x","volume":"132","author":"P Gao","year":"2024","unstructured":"Gao, P., et al.: Clip-adapter: better vision-language models with feature adapters. Int. J. Comput. Vis. 132(2), 581\u2013595 (2024)","journal-title":"Int. J. Comput. Vis."},{"key":"67_CR8","unstructured":"Lei, J., et\u00a0al.: Unibrain: universal brain MRI diagnosis with hierarchical knowledge-enhanced pre-training. arXiv preprint arXiv:2309.06828 (2023)"},{"key":"67_CR9","doi-asserted-by":"crossref","unstructured":"Li, L., Xu, M., Wang, X., Jiang, L., Liu, H.: Attention based glaucoma detection: a large-scale database and CNN model. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 10571\u201310580 (2019)","DOI":"10.1109\/CVPR.2019.01082"},{"key":"67_CR10","unstructured":"Li, M., et\u00a0al.: FFA-IR: towards an explainable and reliable medical report generation benchmark. In: Thirty-fifth Conference on Neural Information Processing Systems Datasets and Benchmarks Track (Round 2) (2021)"},{"key":"67_CR11","doi-asserted-by":"crossref","unstructured":"Li, X., et al.: Multi-modal multi-instance learning for retinal disease recognition. In: Proceedings of the 29th ACM International Conference on Multimedia, pp. 2474\u20132482 (2021)","DOI":"10.1145\/3474085.3475418"},{"key":"67_CR12","doi-asserted-by":"crossref","unstructured":"Liu, J., et al.: Clip-driven universal model for organ segmentation and tumor detection. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 21152\u201321164 (2023)","DOI":"10.1109\/ICCV51070.2023.01934"},{"key":"67_CR13","doi-asserted-by":"publisher","DOI":"10.1016\/j.media.2019.101570","volume":"59","author":"JI Orlando","year":"2020","unstructured":"Orlando, J.I., et al.: Refuge challenge: a unified framework for evaluating automated methods for glaucoma assessment from fundus photographs. Med. Image Anal. 59, 101570 (2020)","journal-title":"Med. Image Anal."},{"key":"67_CR14","doi-asserted-by":"crossref","unstructured":"Pellegrini, C., Keicher, M., \u00d6zsoy, E., Jiraskova, P., Braren, R., Navab, N.: Xplainer: from x-ray observations to explainable zero-shot diagnosis. arXiv preprint arXiv:2303.13391 (2023)","DOI":"10.1007\/978-3-031-43904-9_41"},{"key":"67_CR15","unstructured":"Radford, A., et\u00a0al.: Learning transferable visual models from natural language supervision. In: International Conference on Machine Learning, pp. 8748\u20138763. PMLR (2021)"},{"key":"67_CR16","unstructured":"Shang, F., et al.: Synfundus: a synthetic fundus images dataset with millions of samples and multi-disease annotations. arXiv preprint arXiv:2312.00377 (2023)"},{"key":"67_CR17","unstructured":"Silva-Rodriguez, J., Chakor, H., Kobbi, R., Dolz, J., Ayed, I.B.: A foundation language-image model of the retina (flair): encoding expert knowledge in text supervision. arXiv preprint arXiv:2308.07898 (2023)"},{"issue":"12","key":"67_CR18","doi-asserted-by":"publisher","first-page":"1399","DOI":"10.1038\/s41551-022-00936-9","volume":"6","author":"E Tiu","year":"2022","unstructured":"Tiu, E., Talius, E., Patel, P., Langlotz, C.P., Ng, A.Y., Rajpurkar, P.: Expert-level detection of pathologies from unannotated chest x-ray images via self-supervised learning. Nat. Biomed. Eng. 6(12), 1399\u20131406 (2022)","journal-title":"Nat. Biomed. Eng."},{"key":"67_CR19","unstructured":"Vaswani, A., et al.: Attention is all you need. Adv. Neural Inf. Process. Syst. 30 (2017)"},{"key":"67_CR20","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"30","DOI":"10.1007\/978-3-030-32239-7_4","volume-title":"Medical Image Computing and Computer Assisted Intervention \u2013 MICCAI 2019","author":"X Wang","year":"2019","unstructured":"Wang, X., Ju, L., Zhao, X., Ge, Z.: Retinal abnormalities recognition using regional multitask learning. In: Shen, D., et al. (eds.) MICCAI 2019. LNCS, vol. 11764, pp. 30\u201338. Springer, Cham (2019). https:\/\/doi.org\/10.1007\/978-3-030-32239-7_4"},{"key":"67_CR21","unstructured":"Wu, C., Yin, S., Qi, W., Wang, X., Tang, Z., Duan, N.: Visual chatgpt: talking, drawing and editing with visual foundation models. arXiv preprint arXiv:2303.04671 (2023)"},{"key":"67_CR22","unstructured":"Wu, J., et al.: Medsegdiff: medical image segmentation with diffusion probabilistic model. In: Medical Imaging with Deep Learning, pp. 1623\u20131639. PMLR (2024)"},{"key":"67_CR23","doi-asserted-by":"publisher","unstructured":"Zhang, R., et al.: Tip-adapter: training-free adaption of CLIP for few-shot classification. In: Avidan, S., Brostow, G., Ciss\u00e9, M., Farinella, G.M., Hassner, T. (eds.) Computer Vision \u2013 ECCV 2022. ECCV 2022. LNCS, vol. 13695, pp. 493\u2013510. Springer, Cham (2022). https:\/\/doi.org\/10.1007\/978-3-031-19833-5_29","DOI":"10.1007\/978-3-031-19833-5_29"},{"key":"67_CR24","doi-asserted-by":"crossref","unstructured":"Zhao, Z., et al.: Bira-net: bilinear attention net for diabetic retinopathy grading. In: 2019 IEEE International Conference on Image Processing (ICIP), pp. 1385\u20131389. IEEE (2019)","DOI":"10.1109\/ICIP.2019.8803074"},{"key":"67_CR25","doi-asserted-by":"crossref","unstructured":"Zhou, Y., et al.: Collaborative learning of semi-supervised segmentation and classification for medical images. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 2079\u20132088 (2019)","DOI":"10.1109\/CVPR.2019.00218"},{"key":"67_CR26","doi-asserted-by":"publisher","unstructured":"Zhou, Y., Yang, G., Zhou, Y., Ding, D., Zhao, J.: Representation, alignment, fusion: a generic transformer-based framework for multi-modal glaucoma recognition. In: Greenspan, H., et al. (eds.) Medical Image Computing and Computer Assisted Intervention \u2013 MICCAI 2023. MICCAI 2023. LNCS, vol. 14226, pp. 704\u2013713. Springer, Cham (2023). https:\/\/doi.org\/10.1007\/978-3-031-43990-2_66","DOI":"10.1007\/978-3-031-43990-2_66"},{"issue":"7981","key":"67_CR27","doi-asserted-by":"publisher","first-page":"156","DOI":"10.1038\/s41586-023-06555-x","volume":"622","author":"Y Zhou","year":"2023","unstructured":"Zhou, Y., et al.: A foundation model for generalizable disease detection from retinal images. Nature 622(7981), 156\u2013163 (2023)","journal-title":"Nature"},{"key":"67_CR28","unstructured":"Zhu, D., Chen, J., Shen, X., Li, X., Elhoseiny, M.: Minigpt-4: enhancing vision-language understanding with advanced large language models. arXiv preprint arXiv:2304.10592 (2023)"}],"container-title":["Lecture Notes in Computer Science","Medical Image Computing and Computer Assisted Intervention \u2013 MICCAI 2024"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/978-3-031-72378-0_67","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,2,14]],"date-time":"2026-02-14T07:30:25Z","timestamp":1771054225000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/978-3-031-72378-0_67"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024]]},"ISBN":["9783031723773","9783031723780"],"references-count":28,"URL":"https:\/\/doi.org\/10.1007\/978-3-031-72378-0_67","relation":{},"ISSN":["0302-9743","1611-3349"],"issn-type":[{"value":"0302-9743","type":"print"},{"value":"1611-3349","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024]]},"assertion":[{"value":"3 October 2024","order":1,"name":"first_online","label":"First Online","group":{"name":"ChapterHistory","label":"Chapter History"}},{"value":"The authors\u00a0have no competing interests to declare that are relevant to the content of this article.","order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Disclosure of Interests"}},{"value":"MICCAI","order":1,"name":"conference_acronym","label":"Conference Acronym","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"International Conference on Medical Image Computing and Computer-Assisted Intervention","order":2,"name":"conference_name","label":"Conference Name","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Marrakesh","order":3,"name":"conference_city","label":"Conference City","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Morocco","order":4,"name":"conference_country","label":"Conference Country","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"2024","order":5,"name":"conference_year","label":"Conference Year","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"7 October 2024","order":7,"name":"conference_start_date","label":"Conference Start Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"11 October 2024","order":8,"name":"conference_end_date","label":"Conference End Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"27","order":9,"name":"conference_number","label":"Conference Number","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"miccai2024","order":10,"name":"conference_id","label":"Conference ID","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"https:\/\/conferences.miccai.org\/2024\/en\/","order":11,"name":"conference_url","label":"Conference URL","group":{"name":"ConferenceInfo","label":"Conference Information"}}]}}