{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,20]],"date-time":"2026-05-20T16:36:03Z","timestamp":1779294963838,"version":"3.51.4"},"reference-count":162,"publisher":"Springer Science and Business Media LLC","issue":"12","license":[{"start":{"date-parts":[[2025,5,24]],"date-time":"2025-05-24T00:00:00Z","timestamp":1748044800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,5,24]],"date-time":"2025-05-24T00:00:00Z","timestamp":1748044800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Vis Comput"],"published-print":{"date-parts":[[2025,9]]},"DOI":"10.1007\/s00371-025-03988-5","type":"journal-article","created":{"date-parts":[[2025,5,24]],"date-time":"2025-05-24T02:24:25Z","timestamp":1748053465000},"page":"9601-9627","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":3,"title":["Generative artificial intelligence for ophthalmic images: developments, applications and challenges"],"prefix":"10.1007","volume":"41","author":[{"given":"Tingyao","family":"Li","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Zheyuan","family":"Wang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Zehua","family":"Jiang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Huaiqin","family":"Zhong","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yiming","family":"Qin","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,5,24]]},"reference":[{"issue":"1","key":"3988_CR1","doi-asserted-by":"crossref","first-page":"355","DOI":"10.1186\/s12886-022-02577-7","volume":"22","author":"H Abdelmotaal","year":"2022","unstructured":"Abdelmotaal, H., Sharaf, M., Soliman, W., Wasfi, E., Kedwany, S.M.: Bridging the resources gap: deep learning for fluorescein angiography and optical coherence tomography macular thickness map image translation. BMC Ophthalmol. 22(1), 355 (2022)","journal-title":"BMC Ophthalmol."},{"key":"3988_CR2","unstructured":"Agrawal, P., Antoniak, S., Hanna, E.B., Bout, B., Chaplot, D., Chudnovsky, J., Costa, D., De\u00a0Monicault, B., Garg, S., Gervet, T., et\u00a0al.: Pixtral 12b. arXiv preprint arXiv:2410.07073 (2024)"},{"key":"3988_CR3","unstructured":"Alaa, A., Van\u00a0Breugel, B., Saveliev, E.S., Van Der\u00a0Schaar, M.: How faithful is your synthetic data? Sample-level metrics for evaluating and auditing generative models. In: International Conference on Machine Learning, pp. 290\u2013306. PMLR (2022)"},{"key":"3988_CR4","first-page":"23716","volume":"35","author":"JB Alayrac","year":"2022","unstructured":"Alayrac, J.B., Donahue, J., Luc, P., Miech, A., Barr, I., Hasson, Y., Lenc, K., Mensch, A., Millican, K., Reynolds, M., et al.: Flamingo: a visual language model for few-shot learning. Adv. Neural Inf. Process. Syst. 35, 23716\u201323736 (2022)","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"3988_CR5","doi-asserted-by":"crossref","unstructured":"Alimanov, A., Islam, M.B.: Denoising diffusion probabilistic model for retinal image generation and segmentation. In: 2023 IEEE International Conference on Computational Photography (ICCP), pp. 1\u201312. IEEE (2023)","DOI":"10.1109\/ICCP56744.2023.10233841"},{"key":"3988_CR6","doi-asserted-by":"crossref","unstructured":"Alimanov, A., Islam, M.B.: Advancing retinal image segmentation: a denoising diffusion probabilistic model perspective. In: 2024 IEEE 7th International Conference on Multimedia Information Processing and Retrieval (MIPR), pp. 572\u2013578. IEEE (2024)","DOI":"10.1109\/MIPR62202.2024.00098"},{"issue":"1","key":"3988_CR7","doi-asserted-by":"crossref","first-page":"60","DOI":"10.3390\/electronics11010060","volume":"11","author":"P Andreini","year":"2021","unstructured":"Andreini, P., Ciano, G., Bonechi, S., Graziani, C., Lachi, V., Mecocci, A., Sodi, A., Scarselli, F., Bianchini, M.: A two-stage gan for high-resolution retinal image generation and segmentation. Electronics 11(1), 60 (2021)","journal-title":"Electronics"},{"key":"3988_CR8","doi-asserted-by":"crossref","unstructured":"Antol, S., Agrawal, A., Lu, J., Mitchell, M., Batra, D., Zitnick, C.L., Parikh, D.: Vqa: visual question answering. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 2425\u20132433 (2015)","DOI":"10.1109\/ICCV.2015.279"},{"key":"3988_CR9","unstructured":"Arjovsky, M., Chintala, S., Bottou, L.: Wasserstein generative adversarial networks. In: D.\u00a0Precup, Y.W. Teh (eds.) Proceedings of the 34th International Conference on Machine Learning, Proceedings of Machine Learning Research, vol.\u00a070, pp. 214\u2013223. PMLR (2017). https:\/\/proceedings.mlr.press\/v70\/arjovsky17a.html"},{"issue":"7","key":"3988_CR10","first-page":"2834","volume":"65","author":"J Baek","year":"2024","unstructured":"Baek, J., He, Y., Emamverdi, M., Mahmoudi, A., Nittala, M.G., Corradetti, G., Ip, M.S., Sadda, S.R.: Prediction of long-term anatomic treatment outcomes for diabetic macular edema using a generative adversarial network. Investig. Ophthalmol. Vis. Sci. 65(7), 2834\u20132834 (2024)","journal-title":"Investig. Ophthalmol. Vis. Sci."},{"key":"3988_CR11","unstructured":"Bordes, F., Pang, R.Y., Ajay, A., Li, A.C., Bardes, A., Petryk, S., Ma\u00f1as, O., Lin, Z., Mahmoud, A., Jayaraman, B., et\u00a0al.: An introduction to vision-language modeling. arXiv preprint arXiv:2405.17247 (2024)"},{"key":"3988_CR12","doi-asserted-by":"crossref","unstructured":"Boreiko, V., Ilanchezian, I., Ayhan, M.S., M\u00fcller, S., Koch, L.M., Faber, H., Berens, P., Hein, M.: Visual explanations for the detection of diabetic retinopathy from retinal fundus images. In: International Conference on Medical Image Computing and Computer-assisted Intervention, pp. 539\u2013549. Springer, Berlin (2022)","DOI":"10.1007\/978-3-031-16434-7_52"},{"key":"3988_CR13","unstructured":"Carlini, N., Hayes, J., Nasr, M., Jagielski, M., Sehwag, V., Tramer, F., Balle, B., Ippolito, D., Wallace, E.: Extracting Training Data from Diffusion Models. In: 32nd USENIX Security Symposium (USENIX Security 23), pp. 5253\u20135270 (2023)"},{"issue":"4","key":"3988_CR14","doi-asserted-by":"crossref","first-page":"2195","DOI":"10.3390\/app13042195","volume":"13","author":"CW Chang","year":"2023","unstructured":"Chang, C.W., Chang, C.Y., Lin, Y.Y., Su, W.W., Chen, H.S.L.: A glaucoma detection system based on generative adversarial network and incremental learning. Appl. Sci. 13(4), 2195 (2023)","journal-title":"Appl. Sci."},{"key":"3988_CR15","unstructured":"Chen, R., Xu, K., Zheng, K., Zhang, W., Lu, Y., Shi, D., He, M.: Uwf-ri2fa: generating multi-frame ultrawide-field fluorescein angiography from ultrawide-field retinal imaging improves diabetic retinopathy stratification. arXiv preprint arXiv:2408.10636 (2024)"},{"issue":"1","key":"3988_CR16","doi-asserted-by":"crossref","first-page":"34","DOI":"10.1038\/s41746-024-01018-7","volume":"7","author":"R Chen","year":"2024","unstructured":"Chen, R., Zhang, W., Song, F., Yu, H., Cao, D., Zheng, Y., He, M., Shi, D.: Translating color fundus photography to indocyanine green angiography using deep-learning for age-related macular degeneration screening. NPJ Digit. Med. 7(1), 34 (2024)","journal-title":"NPJ Digit. Med."},{"issue":"1","key":"3988_CR17","doi-asserted-by":"crossref","first-page":"111","DOI":"10.1038\/s41746-024-01101-z","volume":"7","author":"X Chen","year":"2024","unstructured":"Chen, X., Zhang, W., Xu, P., Zhao, Z., Zheng, Y., Shi, D., He, M.: FFA-GPT: an automated pipeline for fundus fluorescein angiography interpretation and question-answer. NPJ Dig. Med. 7(1), 111 (2024)","journal-title":"NPJ Dig. Med."},{"issue":"10","key":"3988_CR18","doi-asserted-by":"crossref","first-page":"1450","DOI":"10.1136\/bjo-2023-324446","volume":"108","author":"X Chen","year":"2024","unstructured":"Chen, X., Zhang, W., Zhao, Z., Xu, P., Zheng, Y., Shi, D., He, M.: ICGA-GPT: report generation and question answering for indocyanine green angiography images. Br. J. Ophthalmol. 108(10), 1450\u20131456 (2024)","journal-title":"Br. J. Ophthalmol."},{"key":"3988_CR19","volume":"55","author":"Z Chen","year":"2020","unstructured":"Chen, Z., Zeng, Z., Shen, H., Zheng, X., Dai, P., Ouyang, P.: DN-GAN: denoising generative adversarial networks for speckle noise reduction in optical coherence tomography images. Biomed. Signal Process. Control 55, 101632 (2020)","journal-title":"Biomed. Signal Process. Control"},{"key":"3988_CR20","unstructured":"Cheng, P., Lin, L., Huang, Y., He, H., Luo, W., Tang, X.: Learning enhancement from degradation: a diffusion model for fundus image enhancement. arXiv preprint arXiv:2303.04603 (2023)"},{"issue":"2","key":"3988_CR21","doi-asserted-by":"crossref","first-page":"23","DOI":"10.1167\/tvst.9.2.23","volume":"9","author":"H Cheong","year":"2020","unstructured":"Cheong, H., Devalla, S.K., Pham, T.H., Zhang, L., Tun, T.A., Wang, X., Perera, S., Schmetterer, L., Aung, T., Boote, C., et al.: Deshadowgan: a deep learning approach to remove shadows from optical coherence tomography images. Transl. Vis. Sci. Technol. 9(2), 23\u201323 (2020)","journal-title":"Transl. Vis. Sci. Technol."},{"key":"3988_CR22","doi-asserted-by":"crossref","DOI":"10.1016\/j.media.2022.102479","volume":"80","author":"H Chung","year":"2022","unstructured":"Chung, H., Ye, J.C.: Score-based diffusion models for accelerated MRI. Med. Image Anal. 80, 102479 (2022)","journal-title":"Med. Image Anal."},{"issue":"3","key":"3988_CR23","doi-asserted-by":"crossref","first-page":"781","DOI":"10.1109\/TMI.2017.2759102","volume":"37","author":"P Costa","year":"2017","unstructured":"Costa, P., Galdran, A., Meyer, M.I., Niemeijer, M., Abr\u00e0moff, M., Mendon\u00e7a, A.M., Campilho, A.: End-to-end adversarial retinal image synthesis. IEEE Trans. Med. Imaging 37(3), 781\u2013791 (2017)","journal-title":"IEEE Trans. Med. Imaging"},{"issue":"2","key":"3988_CR24","doi-asserted-by":"crossref","first-page":"584","DOI":"10.1038\/s41591-023-02702-z","volume":"30","author":"L Dai","year":"2024","unstructured":"Dai, L., Sheng, B., Chen, T., Wu, Q., Liu, R., Cai, C., Wu, L., Yang, D., Hamzah, H., Liu, Y., et al.: A deep learning system for predicting time to progression of diabetic retinopathy. Nat. Med. 30(2), 584\u2013594 (2024)","journal-title":"Nat. Med."},{"issue":"1","key":"3988_CR25","doi-asserted-by":"crossref","first-page":"3242","DOI":"10.1038\/s41467-021-23458-5","volume":"12","author":"L Dai","year":"2021","unstructured":"Dai, L., Wu, L., Li, H., Cai, C., Wu, Q., Kong, H., Liu, R., Wang, X., Hou, X., Liu, Y., et al.: A deep learning system for detecting diabetic retinopathy across the disease spectrum. Nat. Commun. 12(1), 3242 (2021)","journal-title":"Nat. Commun."},{"issue":"15","key":"3988_CR26","doi-asserted-by":"crossref","first-page":"8746","DOI":"10.1109\/JSEN.2020.2985131","volume":"20","author":"V Das","year":"2020","unstructured":"Das, V., Dandapat, S., Bora, P.K.: Unsupervised super-resolution of oct images using generative adversarial network for improved age-related macular degeneration diagnosis. IEEE Sens. J. 20(15), 8746\u20138756 (2020)","journal-title":"IEEE Sens. J."},{"key":"3988_CR27","unstructured":"de\u00a0Vente, C., Islam, M.M., Valmaggia, P., Hoyng, C., Tufail, A., S\u00e1nchez, C.I.: Conditioning 3d diffusion models with 2d images: Towards standardized oct volumes through en face-informed super-resolution. arXiv preprint arXiv:2410.09862 (2024)"},{"key":"3988_CR28","doi-asserted-by":"crossref","first-page":"279","DOI":"10.1016\/j.patrec.2024.10.012","volume":"186","author":"J Dong","year":"2024","unstructured":"Dong, J., Qian, T., Jiang, Y., Bi, L., Kim, J., Wang, L., Xu, X.: Claritydiffusenet: enhancing fundus image quality under black shadows with diffusion model-based research. Pattern Recogn. Lett. 186, 279\u2013285 (2024)","journal-title":"Pattern Recogn. Lett."},{"issue":"9","key":"3988_CR29","first-page":"1555","volume":"11","author":"XL Du","year":"2018","unstructured":"Du, X.L., Li, W.B., Hu, B.J.: Application of artificial intelligence in ophthalmology. Int. J. Ophthalmol. 11(9), 1555 (2018)","journal-title":"Int. J. Ophthalmol."},{"key":"3988_CR30","doi-asserted-by":"crossref","unstructured":"Elbatel, M., Kamnitsas, K., Li, X.: An organism starts with a single pix-cell: a neural cellular diffusion for high-resolution image synthesis. In: International Conference on Medical Image Computing and Computer-Assisted Intervention, pp. 656\u2013666. Springer, Berlin (2024)","DOI":"10.1007\/978-3-031-72378-0_61"},{"key":"3988_CR31","unstructured":"Esser, P., Kulal, S., Blattmann, A., Entezari, R., M\u00fcller, J., Saini, H., Levi, Y., Lorenz, D., Sauer, A., Boesel, F., Podell, D., Dockhorn, T., English, Z., Rombach, R.: Scaling rectified flow transformers for high-resolution image synthesis. In: Proceedings of the 41st International Conference on Machine Learning, ICML\u201924. JMLR.org (2024)"},{"key":"3988_CR32","doi-asserted-by":"crossref","unstructured":"Esser, P., Rombach, R., Ommer, B.: Taming transformers for high-resolution image synthesis. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 12873\u201312883 (2021)","DOI":"10.1109\/CVPR46437.2021.01268"},{"key":"3988_CR33","doi-asserted-by":"crossref","unstructured":"Fang, Z., Chen, Z., Wei, P., Li, W., Zhang, S., Elazab, A., Jia, G., Ge, R., Wang, C.: Uwat-gan: fundus fluorescein angiography synthesis via ultra-wide-angle transformation multi-scale gan. In: International Conference on Medical Image Computing and Computer-Assisted Intervention, pp. 745\u2013755. Springer, Berlin (2023)","DOI":"10.1007\/978-3-031-43990-2_70"},{"key":"3988_CR34","doi-asserted-by":"crossref","DOI":"10.1016\/j.eswa.2025.126471","volume":"270","author":"Z Fang","year":"2025","unstructured":"Fang, Z., Yu, X., Zhou, G., Zhuang, K., Chen, Y., Ge, R., Wang, C., Jia, G., Wu, Q., Ye, J., et al.: LPUWF-LDM: enhanced latent diffusion model for precise late-phase UWF-fa generation on limited dataset. Expert Syst.Appl. 270, 126471 (2025)","journal-title":"Expert Syst.Appl."},{"key":"3988_CR35","doi-asserted-by":"crossref","unstructured":"Ge, R., Fang, Z., Wei, P., Chen, Z., Jiang, H., Elazab, A., Li, W., Wan, X., Zhang, S., Wang, C.: UWAFA-GAN: ultra-wide-angle fluorescein angiography transformation via multi-scale generation and registration enhancement. IEEE J. Biomed. Health Inform. (2024)","DOI":"10.1109\/JBHI.2024.3394597"},{"key":"3988_CR36","doi-asserted-by":"crossref","unstructured":"Go, S., Ji, Y., Park, S.J., Lee, S.: Generation of structurally realistic retinal fundus images with diffusion models. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 2335\u20132344 (2024)","DOI":"10.1109\/CVPRW63382.2024.00239"},{"key":"3988_CR37","unstructured":"Goodfellow, I.J., Pouget-Abadie, J., Mirza, M., Xu, B., Warde-Farley, D., Ozair, S., Courville, A.C., Bengio, Y.: Generative adversarial nets. In: Z.\u00a0Ghahramani, M.\u00a0Welling, C.\u00a0Cortes, N.D. Lawrence, K.Q. Weinberger (eds.) Advances in Neural Information Processing Systems 27: Annual Conference on Neural Information Processing Systems 2014, December 8-13 2014, Montreal, Quebec, Canada, pp. 2672\u20132680 (2014)"},{"key":"3988_CR38","unstructured":"Guo, D., Yang, D., Zhang, H., Song, J., Zhang, R., Xu, R., Zhu, Q., Ma, S., Wang, P., Bi, X., et\u00a0al.: Deepseek-r1: incentivizing reasoning capability in llms via reinforcement learning. arXiv preprint arXiv:2501.12948 (2025)"},{"issue":"12","key":"3988_CR39","doi-asserted-by":"crossref","first-page":"6205","DOI":"10.1364\/BOE.9.006205","volume":"9","author":"KJ Halupka","year":"2018","unstructured":"Halupka, K.J., Antony, B.J., Lee, M.H., Lucy, K.A., Rai, R.S., Ishikawa, H., Wollstein, G., Schuman, J.S., Garnavi, R.: Retinal optical coherence tomography image enhancement via deep learning. Biomed. Opt. Express 9(12), 6205\u20136221 (2018)","journal-title":"Biomed. Opt. Express"},{"key":"3988_CR40","doi-asserted-by":"crossref","first-page":"1430984","DOI":"10.3389\/frai.2024.1430984","volume":"7","author":"I Hartsock","year":"2024","unstructured":"Hartsock, I., Rasool, G.: Vision-language models for medical report generation and visual question answering: a review. Front. Artif. Intell. 7, 1430984 (2024)","journal-title":"Front. Artif. Intell."},{"issue":"12","key":"3988_CR41","doi-asserted-by":"crossref","first-page":"20","DOI":"10.1167\/tvst.12.12.20","volume":"12","author":"S He","year":"2023","unstructured":"He, S., Joseph, S., Bulloch, G., Jiang, F., Kasturibai, H., Kim, R., Ravilla, T.D., Wang, Y., Shi, D., He, M.: Bridging the camera domain gap with image-to-image translation improves glaucoma diagnosis. Transl. Vis. Sci. Technol. 12(12), 20\u201320 (2023)","journal-title":"Transl. Vis. Sci. Technol."},{"key":"3988_CR42","first-page":"6840","volume":"33","author":"J Ho","year":"2020","unstructured":"Ho, J., Jain, A., Abbeel, P.: Denoising diffusion probabilistic models. Adv. Neural. Inf. Process. Syst. 33, 6840\u20136851 (2020)","journal-title":"Adv. Neural. Inf. Process. Syst."},{"key":"3988_CR43","unstructured":"Ho, J., Salimans, T.: Classifier-free diffusion guidance. In: NeurIPS 2021 Workshop on Deep Generative Models and Downstream Applications (2021)"},{"key":"3988_CR44","doi-asserted-by":"crossref","unstructured":"Hu, D., Li, H., Liu, H., Wang, J., Yao, X., Lu, D., Oguz, I.: Adaptdiff: cross-modality domain adaptation via weak conditional semantic diffusion for retinal vessel segmentation. In: International Workshop on Simulation and Synthesis in Medical Imaging, pp. 13\u201323. Springer, Berlin (2024)","DOI":"10.1007\/978-3-031-73281-2_2"},{"key":"3988_CR45","doi-asserted-by":"crossref","unstructured":"Hu, D., Tao, Y.K., Oguz, I.: Unsupervised denoising of retinal oct with diffusion probabilistic model. In: Medical Imaging 2022: Image Processing, vol. 12032, pp. 25\u201334. SPIE (2022)","DOI":"10.1117\/12.2612235"},{"key":"3988_CR46","volume":"229","author":"K Huang","year":"2023","unstructured":"Huang, K., Li, M., Yu, J., Miao, J., Hu, Z., Yuan, S., Chen, Q.: Lesion-aware generative adversarial networks for color fundus image to fundus fluorescein angiography translation. Comput. Methods Progr. Biomed. 229, 107306 (2023)","journal-title":"Comput. Methods Progr. Biomed."},{"key":"3988_CR47","doi-asserted-by":"crossref","unstructured":"Huang, K., Ma, X., Zhang, Y., Su, N., Yuan, S., Liu, Y., Chen, Q., Fu, H.: Memory-efficient high-resolution oct volume synthesis with cascaded amortized latent diffusion models. In: International Conference on Medical Image Computing and Computer-Assisted Intervention, pp. 478\u2013487. Springer, Berlin (2024)","DOI":"10.1007\/978-3-031-72104-5_46"},{"key":"3988_CR48","volume":"253","author":"W Huang","year":"2024","unstructured":"Huang, W., Liu, F.: HiDiffSeg: a hierarchical diffusion model for blood vessel segmentation in retinal fundus images. Expert Syst. Appl. 253, 124249 (2024)","journal-title":"Expert Syst. Appl."},{"issue":"9","key":"3988_CR49","doi-asserted-by":"crossref","first-page":"12289","DOI":"10.1364\/OE.27.012289","volume":"27","author":"Y Huang","year":"2019","unstructured":"Huang, Y., Lu, Z., Shao, Z., Ran, M., Zhou, J., Fang, L., Zhang, Y.: Simultaneous denoising and super-resolution of optical coherence tomography images based on generative adversarial network. Opt. Express 27(9), 12289\u201312307 (2019)","journal-title":"Opt. Express"},{"key":"3988_CR50","doi-asserted-by":"crossref","DOI":"10.1016\/j.compbiomed.2025.109834","volume":"189","author":"M Ibrahim","year":"2025","unstructured":"Ibrahim, M., Al Khalil, Y., Amirrajab, S., Sun, C., Breeuwer, M., Pluim, J., Elen, B., Ertaylan, G., Dumontier, M.: Generative ai for synthetic data across multiple medical modalities: a systematic review of recent developments and challenges. Comput. Biol. Med. 189, 109834 (2025)","journal-title":"Comput. Biol. Med."},{"key":"3988_CR51","unstructured":"Ilanchezian, I., Boreiko, V., K\u00fchlewein, L., Huang, Z., Ayhan, M.S., Hein, M., Koch, L., Berens, P.: Generating realistic counterfactuals for retinal fundus and oct images using diffusion models. arXiv preprint arXiv:2311.11629 (2023)"},{"key":"3988_CR52","doi-asserted-by":"crossref","unstructured":"Isola, P., Zhu, J.Y., Zhou, T., Efros, A.A.: Image-to-image translation with conditional adversarial networks. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 1125\u20131134 (2017)","DOI":"10.1109\/CVPR.2017.632"},{"key":"3988_CR53","first-page":"14938","volume":"34","author":"A Jalal","year":"2021","unstructured":"Jalal, A., Arvinte, M., Daras, G., Price, E., Dimakis, A.G., Tamir, J.: Robust compressed sensing MRI with deep generative priors. Adv. Neural Inf. Process. Syst. 34, 14938\u201314954 (2021)","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"3988_CR54","unstructured":"Jia, C., Yang, Y., Xia, Y., Chen, Y.T., Parekh, Z., Pham, H., Le, Q., Sung, Y.H., Li, Z., Duerig, T.: Scaling up visual and vision-language representation learning with noisy text supervision. In: International Conference on Machine Learning, pp. 4904\u20134916. PMLR (2021)"},{"issue":"10","key":"3988_CR55","doi-asserted-by":"crossref","first-page":"2911","DOI":"10.1109\/TMI.2021.3056395","volume":"40","author":"L Ju","year":"2021","unstructured":"Ju, L., Wang, X., Zhao, X., Bonnington, P., Drummond, T., Ge, Z.: Leveraging regular fundus images for training UWF fundus diagnosis models via adversarial learning and pseudo-labeling. IEEE Trans. Med. Imaging 40(10), 2911\u20132925 (2021)","journal-title":"IEEE Trans. Med. Imaging"},{"issue":"4","key":"3988_CR56","doi-asserted-by":"crossref","DOI":"10.1016\/j.xops.2024.100493","volume":"4","author":"SA Kamran","year":"2024","unstructured":"Kamran, S.A., Hossain, K.F., Ong, J., Waisberg, E., Zaman, N., Baker, S.A., Lee, A.G., Tavakkoli, A.: FA4SANS-GAN: a novel machine learning generative adversarial network to further understand ophthalmic changes in spaceflight associated neuro-ocular syndrome (sans). Ophthalmol. Sci. 4(4), 100493 (2024)","journal-title":"Ophthalmol. Sci."},{"key":"3988_CR57","doi-asserted-by":"crossref","unstructured":"Kamran, S.A., Hossain, K.F., Tavakkoli, A., Zuckerbrod, S.L., Baker, S.A.: Vtgan: semi-supervised retinal image synthesis and disease prediction using vision transformers. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 3235\u20133245 (2021)","DOI":"10.1109\/ICCVW54120.2021.00362"},{"issue":"2","key":"3988_CR58","doi-asserted-by":"crossref","first-page":"233","DOI":"10.1016\/j.survophthal.2018.09.002","volume":"64","author":"R Kapoor","year":"2019","unstructured":"Kapoor, R., Walters, S.P., Al-Aswad, L.A.: The current state of artificial intelligence in ophthalmology. Surv. Ophthalmol. 64(2), 233\u2013240 (2019)","journal-title":"Surv. Ophthalmol."},{"key":"3988_CR59","unstructured":"Karras, T., Aila, T., Laine, S., Lehtinen, J.: Progressive growing of gans for improved quality, stability, and variation. In: International Conference on Learning Representations (2018)"},{"key":"3988_CR60","doi-asserted-by":"crossref","unstructured":"Karras, T., Laine, S., Aila, T.: A style-based generator architecture for generative adversarial networks. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 4401\u20134410 (2019)","DOI":"10.1109\/CVPR.2019.00453"},{"key":"3988_CR61","doi-asserted-by":"crossref","DOI":"10.1016\/j.media.2023.102846","volume":"88","author":"A Kazerouni","year":"2023","unstructured":"Kazerouni, A., Aghdam, E.K., Heidari, M., Azad, R., Fayyaz, M., Hacihaliloglu, I., Merhof, D.: Diffusion models in medical imaging: a comprehensive survey. Med. Image Anal. 88, 102846 (2023)","journal-title":"Med. Image Anal."},{"issue":"4","key":"3988_CR62","doi-asserted-by":"crossref","first-page":"188","DOI":"10.1007\/s42452-024-05871-9","volume":"6","author":"HK Kim","year":"2024","unstructured":"Kim, H.K., Ryu, I.H., Choi, J.Y., Yoo, T.K.: A feasibility study on the adoption of a generative denoising diffusion model for the synthesis of fundus photographs using a small dataset. Discov. Appl. Sci. 6(4), 188 (2024)","journal-title":"Discov. Appl. Sci."},{"issue":"1","key":"3988_CR63","doi-asserted-by":"crossref","first-page":"17307","DOI":"10.1038\/s41598-022-20698-3","volume":"12","author":"M Kim","year":"2022","unstructured":"Kim, M., Kim, Y.N., Jang, M., Hwang, J., Kim, H.K., Yoon, S.C., Kim, Y.J., Kim, N.: Synthesizing realistic high-resolution retina image by style-based generative adversarial network and its utilization. Sci. Rep. 12(1), 17307 (2022)","journal-title":"Sci. Rep."},{"key":"3988_CR64","doi-asserted-by":"crossref","unstructured":"Kim, S., Chung, H., Park, S.H., Chung, E.S., Yi, K., Ye, J.C.: Fundus image enhancement through direct diffusion bridges. IEEE J. Biomed. Health Inform. (2024)","DOI":"10.1109\/JBHI.2024.3446866"},{"key":"3988_CR65","unstructured":"Kingma, D.P., Welling, M.: Auto-encoding variational bayes (2022). arxiv: 1312.6114"},{"issue":"4","key":"3988_CR66","doi-asserted-by":"crossref","first-page":"1166","DOI":"10.1038\/s41591-024-02838-6","volume":"30","author":"I Ktena","year":"2024","unstructured":"Ktena, I., Wiles, O., Albuquerque, I., Rebuffi, S.A., Tanno, R., Roy, A.G., Azizi, S., Belgrave, D., Kohli, P., Cemgil, T., et al.: Generative models improve fairness of medical classifiers under distribution shifts. Nat. Med. 30(4), 1166\u20131173 (2024)","journal-title":"Nat. Med."},{"key":"3988_CR67","doi-asserted-by":"crossref","unstructured":"Lang, O., Gandelsman, Y., Yarom, M., Wald, Y., Elidan, G., Hassidim, A., Freeman, W.T., Isola, P., Globerson, A., Irani, M., et\u00a0al.: Explaining in style: training a gan to explain a classifier in stylespace. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 693\u2013702 (2021)","DOI":"10.1109\/ICCV48922.2021.00073"},{"key":"3988_CR68","doi-asserted-by":"crossref","DOI":"10.1016\/j.media.2020.101906","volume":"68","author":"G Lazaridis","year":"2021","unstructured":"Lazaridis, G., Lorenzi, M., Ourselin, S., Garway-Heath, D.: Improving statistical power of glaucoma clinical trials using an ensemble of cyclical generative adversarial networks. Med. Image Anal. 68, 101906 (2021)","journal-title":"Med. Image Anal."},{"issue":"3","key":"3988_CR69","doi-asserted-by":"crossref","first-page":"572","DOI":"10.1097\/IAE.0000000000002898","volume":"41","author":"H Lee","year":"2021","unstructured":"Lee, H., Kim, S., Kim, M.A., Chung, H., Kim, H.C.: Post-treatment prediction of optical coherence tomography using a conditional generative adversarial network in age-related macular degeneration. Retina 41(3), 572\u2013580 (2021)","journal-title":"Retina"},{"issue":"1","key":"3988_CR70","doi-asserted-by":"crossref","first-page":"464","DOI":"10.1038\/s42003-023-04846-7","volume":"6","author":"W Lee","year":"2023","unstructured":"Lee, W., Nam, H.S., Seok, J.Y., Oh, W.Y., Kim, J.W., Yoo, H.: Deep learning-based image enhancement in optical coherence tomography by exploiting interference fringe. Commun. Biol. 6(1), 464 (2023)","journal-title":"Commun. Biol."},{"key":"3988_CR71","first-page":"28541","volume":"36","author":"C Li","year":"2023","unstructured":"Li, C., Wong, C., Zhang, S., Usuyama, N., Liu, H., Yang, J., Naumann, T., Poon, H., Gao, J.: Llava-med: training a large language-and-vision assistant for biomedicine in one day. Adv. Neural Inf. Process. Syst. 36, 28541\u201328564 (2023)","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"3988_CR72","doi-asserted-by":"crossref","unstructured":"Li, H., Liu, H., Fu, H., Shu, H., Zhao, Y., Luo, X., Hu, Y., Liu, J.: Structure-consistent restoration network for cataract fundus image enhancement. In: Wang, L., Dou, Q., Fletcher, P.T., Speidel, S., Li, S. (eds.) Medical Image Computing and Computer Assisted Intervention\u2014MICCAI 2022, pp. 487\u2013496. Springer, Cham (2022)","DOI":"10.1007\/978-3-031-16434-7_47"},{"key":"3988_CR73","volume":"90","author":"H Li","year":"2023","unstructured":"Li, H., Liu, H., Fu, H., Xu, Y., Shu, H., Niu, K., Hu, Y., Liu, J.: A generic fundus image enhancement network boosted by frequency self-supervised representation learning. Med. Image Anal. 90, 102945 (2023)","journal-title":"Med. Image Anal."},{"issue":"7","key":"3988_CR74","doi-asserted-by":"crossref","first-page":"1699","DOI":"10.1109\/TMI.2022.3147854","volume":"41","author":"H Li","year":"2022","unstructured":"Li, H., Liu, H., Hu, Y., Fu, H., Zhao, Y., Miao, H., Liu, J.: An annotation-free restoration network for cataractous fundus images. IEEE Trans. Med. Imaging 41(7), 1699\u20131710 (2022)","journal-title":"IEEE Trans. Med. Imaging"},{"issue":"10","key":"3988_CR75","doi-asserted-by":"crossref","first-page":"2886","DOI":"10.1038\/s41591-024-03139-8","volume":"30","author":"J Li","year":"2024","unstructured":"Li, J., Guan, Z., Wang, J., Cheung, C.Y., Zheng, Y., Lim, L.L., Lim, C.C., Ruamviboonsuk, P., Raman, R., Corsino, L., et al.: Integrated image-based deep learning and language models for primary diabetes care. Nat. Med. 30(10), 2886\u20132896 (2024)","journal-title":"Nat. Med."},{"key":"3988_CR76","unstructured":"Li, J., Li, D., Xiong, C., Hoi, S.: Blip: Bootstrapping language-image pre-training for unified vision-language understanding and generation. In: International Conference on Machine Learning, pp. 12888\u201312900. PMLR (2022)"},{"key":"3988_CR77","doi-asserted-by":"crossref","unstructured":"Li, S., Higashita, R., Fu, H., Li, H., Niu, J., Liu, J.: Content-preserving diffusion model for unsupervised as-oct image despeckling. In: International Conference on Medical Image Computing and Computer-Assisted Intervention, pp. 660\u2013670. Springer, Berlin (2023)","DOI":"10.1007\/978-3-031-43990-2_62"},{"key":"3988_CR78","unstructured":"Li, Y., Sun, H., Lin, M., Li, T., Dong, G., Zhang, T., Ding, B., Song, W., Cheng, Z., Huo, Y., et\u00a0al.: Ocean-omni: To understand the world with omni-modality. arXiv preprint arXiv:2410.08565 (2024)"},{"key":"3988_CR79","doi-asserted-by":"crossref","unstructured":"Li, Z., Song, D., Yang, Z., Wang, D., Li, F., Zhang, X., Kinahan, P.E., Qiao, Y.: Visionunite: A vision-language foundation model for ophthalmology enhanced with clinical knowledge. arXiv preprint arXiv:2408.02865 (2024)","DOI":"10.1109\/TPAMI.2025.3598734"},{"key":"3988_CR80","doi-asserted-by":"crossref","unstructured":"Li, Z., Wang, L., Wu, X., Jiang, J., Qiang, W., Xie, H., Zhou, H., Wu, S., Shao, Y., Chen, W.: Artificial intelligence in ophthalmology: the path to the real-world clinic. Cell Rep. Med. 4(7) (2023)","DOI":"10.1016\/j.xcrm.2023.101095"},{"issue":"12","key":"3988_CR81","doi-asserted-by":"crossref","first-page":"6563","DOI":"10.1364\/BOE.506205","volume":"14","author":"F Liu","year":"2023","unstructured":"Liu, F., Huang, W.: ESDiff: a joint model for low-quality retinal image enhancement and vessel segmentation using a diffusion model. Biomed. Opt. Express 14(12), 6563\u20136578 (2023)","journal-title":"Biomed. Opt. Express"},{"key":"3988_CR82","doi-asserted-by":"crossref","unstructured":"Liu, F., Shareghi, E., Meng, Z., Basaldella, M., Collier, N.: Self-alignment pretraining for biomedical entity representations. In: Proceedings of the 2021 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies, pp. 4228\u20134238 (2021)","DOI":"10.18653\/v1\/2021.naacl-main.334"},{"key":"3988_CR83","first-page":"34892","volume":"36","author":"H Liu","year":"2023","unstructured":"Liu, H., Li, C., Wu, Q., Lee, Y.J.: Visual instruction tuning. Adv. Neural Inf. Process. Syst. 36, 34892\u201334916 (2023)","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"3988_CR84","unstructured":"Liu, L., Ren, Y., Lin, Z., Zhao, Z.: Pseudo numerical methods for diffusion models on manifolds. In: International Conference on Learning Representations (2022). https:\/\/openreview.net\/forum?id=PlKWVd2yBkY"},{"key":"3988_CR85","unstructured":"Liu, M.Y., Breuel, T., Kautz, J.: Unsupervised image-to-image translation networks. Adv. Neural Inf. Process. Syst. 30 (2017)"},{"key":"3988_CR86","volume":"41","author":"S Liu","year":"2023","unstructured":"Liu, S., Hu, W., Xu, F., Chen, W., Liu, J., Yu, X., Wang, Z., Li, Z., Li, Z., Yang, X., et al.: Prediction of oct images of short-term response to anti-vegf treatment for diabetic macular edema using different generative adversarial networks. Photodiagnosis Photodyn. Ther. 41, 103272 (2023)","journal-title":"Photodiagnosis Photodyn. Ther."},{"issue":"12","key":"3988_CR87","doi-asserted-by":"crossref","first-page":"1735","DOI":"10.1136\/bjophthalmol-2019-315338","volume":"104","author":"Y Liu","year":"2020","unstructured":"Liu, Y., Yang, J., Zhou, Y., Wang, W., Zhao, J., Yu, W., Zhang, D., Ding, D., Li, X., Chen, Y.: Prediction of oct images of short-term response to anti-vegf treatment for neovascular age-related macular degeneration using generative adversarial network. Br. J. Ophthalmol. 104(12), 1735\u20131740 (2020)","journal-title":"Br. J. Ophthalmol."},{"key":"3988_CR88","unstructured":"Liu, Z., Sun, Z., Zang, Y., Dong, X., Cao, Y., Duan, H., Lin, D., Wang, J.: Visual-rft: visual reinforcement fine-tuning. arXiv preprint arXiv:2503.01785 (2025)"},{"issue":"1","key":"3988_CR89","first-page":"5278196","volume":"2018","author":"W Lu","year":"2018","unstructured":"Lu, W., Tong, Y., Yu, Y., Xing, Y., Chen, C., Shen, Y.: Applications of artificial intelligence in ophthalmology: general overview. J. Ophthalmol. 2018(1), 5278196 (2018)","journal-title":"J. Ophthalmol."},{"key":"3988_CR90","doi-asserted-by":"crossref","unstructured":"Lugmayr, A., Danelljan, M., Romero, A., Yu, F., Timofte, R., Van\u00a0Gool, L.: Repaint: inpainting using denoising diffusion probabilistic models. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 11461\u201311471 (2022)","DOI":"10.1109\/CVPR52688.2022.01117"},{"key":"3988_CR91","doi-asserted-by":"crossref","unstructured":"Luo, Y., Zhang, J., Fan, S., Yang, K., Hong, M., Wu, Y., Qiao, M., Nie, Z.: Biomedgpt: An open multimodal large language model for biomedicine. IEEE J. Biomed. Health Inform. (2024)","DOI":"10.1109\/JBHI.2024.3505955"},{"issue":"3","key":"3988_CR92","doi-asserted-by":"crossref","first-page":"1831","DOI":"10.1364\/BOE.517819","volume":"15","author":"X Ma","year":"2024","unstructured":"Ma, X., Ji, Z., Chen, Q., Ge, L., Wang, X., Chen, C., Fan, W.: Controllable editing via diffusion inversion on ultra-widefield fluorescein angiography for the comprehensive analysis of diabetic retinopathy. Biomed. Opt. Express 15(3), 1831\u20131846 (2024)","journal-title":"Biomed. Opt. Express"},{"key":"3988_CR93","doi-asserted-by":"crossref","DOI":"10.1016\/j.patcog.2019.107109","volume":"100","author":"D Mahapatra","year":"2020","unstructured":"Mahapatra, D., Ge, Z.: Training data independent image registration using generative adversarial networks and domain adaptation. Pattern Recognit. 100, 107109 (2020)","journal-title":"Pattern Recognit."},{"key":"3988_CR94","unstructured":"Meng, F., Du, L., Liu, Z., Zhou, Z., Lu, Q., Fu, D., Shi, B., Wang, W., He, J., Zhang, K., et\u00a0al.: MM-Eureka: exploring visual aha moment with rule-based large-scale reinforcement learning. arXiv preprint arXiv:2503.07365 (2025)"},{"issue":"3","key":"3988_CR95","doi-asserted-by":"crossref","DOI":"10.1016\/j.xops.2023.100294","volume":"3","author":"MJ Menten","year":"2023","unstructured":"Menten, M.J., Holland, R., Leingang, O., Bogunovi\u0107, H., Hagag, A.M., Kaye, R., Riedl, S., Traber, G.L., Hassan, O.N., Pawlowski, N., et al.: Exploring healthy retinal aging with deep learning. Ophthalmol. Sci. 3(3), 100294 (2023)","journal-title":"Ophthalmol. Sci."},{"issue":"1","key":"3988_CR96","doi-asserted-by":"crossref","first-page":"5639","DOI":"10.1038\/s41598-023-32398-7","volume":"13","author":"S Moon","year":"2023","unstructured":"Moon, S., Lee, Y., Hwang, J., Kim, C.G., Kim, J.W., Yoon, W.T., Kim, J.H.: Prediction of anti-vascular endothelial growth factor agent-specific treatment outcomes in neovascular age-related macular degeneration using a generative adversarial network. Sci. Rep. 13(1), 5639 (2023)","journal-title":"Sci. Rep."},{"key":"3988_CR97","unstructured":"Moor, M., Huang, Q., Wu, S., Yasunaga, M., Dalmia, Y., Leskovec, J., Zakka, C., Reis, E.P., Rajpurkar, P.: Med-flamingo: a multimodal medical few-shot learner. In: Machine Learning for Health (ML4H), pp. 353\u2013367. PMLR (2023)"},{"issue":"10","key":"3988_CR98","doi-asserted-by":"crossref","first-page":"5291","DOI":"10.1364\/BOE.10.005291","volume":"10","author":"J Ouyang","year":"2019","unstructured":"Ouyang, J., Mathai, T.S., Lathrop, K., Galeotti, J.: Accurate tissue interface segmentation via adversarial pre-segmentation of anterior segment oct images. Biomed. Opt. Express 10(10), 5291\u20135324 (2019)","journal-title":"Biomed. Opt. Express"},{"key":"3988_CR99","doi-asserted-by":"crossref","unstructured":"Peebles, W., Xie, S.: Scalable diffusion models with transformers. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 4195\u20134205 (2023)","DOI":"10.1109\/ICCV51070.2023.00387"},{"key":"3988_CR100","unstructured":"Podell, D., English, Z., Lacey, K., Blattmann, A., Dockhorn, T., M\u00fcller, J., Penna, J., Rombach, R.: Sdxl: Improving latent diffusion models for high-resolution image synthesis. arXiv preprint arXiv:2307.01952 (2023)"},{"key":"3988_CR101","doi-asserted-by":"crossref","unstructured":"Qin, Z., Yin, Y., Campbell, D., Wu, X., Zou, K., Tham, Y.C., Liu, N., Zhang, X., Chen, Q.: Lmod: A large multimodal ophthalmology dataset and benchmark for large vision-language models. arXiv preprint arXiv:2410.01620 (2024)","DOI":"10.18653\/v1\/2025.findings-naacl.135"},{"issue":"2","key":"3988_CR102","doi-asserted-by":"crossref","first-page":"275","DOI":"10.1007\/s12626-023-00145-z","volume":"17","author":"M Quaranta","year":"2023","unstructured":"Quaranta, M., Amantea, I.A., Grosso, M.: Obligation for ai systems in healthcare: prepare for trouble and make it double? Rev. Socionetwork Strateg. 17(2), 275\u2013295 (2023)","journal-title":"Rev. Socionetwork Strateg."},{"key":"3988_CR103","unstructured":"Radford, A., Kim, J.W., Hallacy, C., Ramesh, A., Goh, G., Agarwal, S., Sastry, G., Askell, A., Mishkin, P., Clark, J., et\u00a0al.: Learning transferable visual models from natural language supervision. In: International Conference on Machine Learning, pp. 8748\u20138763. PMLR (2021)"},{"issue":"1","key":"3988_CR104","doi-asserted-by":"crossref","first-page":"86","DOI":"10.1038\/s41746-021-00455-y","volume":"4","author":"L Rasmy","year":"2021","unstructured":"Rasmy, L., Xiang, Y., Xie, Z., Tao, C., Zhi, D.: Med-BERT: pretrained contextualized embeddings on large-scale structured electronic health records for disease prediction. NPJ Dig. Med. 4(1), 86 (2021)","journal-title":"NPJ Dig. Med."},{"key":"3988_CR105","doi-asserted-by":"crossref","unstructured":"Rezagholiradeh, M., Haidar, M.A.: Reg-gan: Semi-supervised learning based on generative adversarial networks for regression. In: 2018 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), pp. 2806\u20132810. IEEE (2018)","DOI":"10.1109\/ICASSP.2018.8462534"},{"issue":"377","key":"3988_CR106","first-page":"1","volume":"24","author":"A Roberts","year":"2023","unstructured":"Roberts, A., Chung, H.W., Mishra, G., Levskaya, A., Bradbury, J., Andor, D., Narang, S., Lester, B., Gaffney, C., Mohiuddin, A., Hawthorne, C., Lewkowycz, A., Salcianu, A., van Zee, M., Austin, J., Goodman, S., Soares, L.B., Hu, H., Tsvyashchenko, S., Chowdhery, A., Bastings, J., Bulian, J., Garcia, X., Ni, J., Chen, A., Kenealy, K., Han, K., Casbon, M., Clark, J.H., Lee, S., Garrette, D., Lee-Thorp, J., Raffel, C., Shazeer, N., Ritter, M., Bosma, M., Passos, A., Maitin-Shepard, J., Fiedel, N., Omernick, M., Saeta, B., Sepassi, R., Spiridonov, A., Newlan, J., Gesmundo, A.: Scaling up models and data with t5x and seqio. J. Mach. Learn. Res. 24(377), 1\u20138 (2023)","journal-title":"J. Mach. Learn. Res."},{"key":"3988_CR107","doi-asserted-by":"crossref","unstructured":"Rombach, R., Blattmann, A., Lorenz, D., Esser, P., Ommer, B.: High-resolution image synthesis with latent diffusion models. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 10684\u201310695 (2022)","DOI":"10.1109\/CVPR52688.2022.01042"},{"key":"3988_CR108","doi-asserted-by":"crossref","unstructured":"Ronneberger, O., Fischer, P., Brox, T.: U-net: Convolutional networks for biomedical image segmentation. In: Medical Image Computing and Computer-assisted Intervention\u2013MICCAI 2015: 18th International Conference, Munich, Germany, October 5-9, 2015, Proceedings, part III 18, pp. 234\u2013241. Springer, Berlin (2015)","DOI":"10.1007\/978-3-319-24574-4_28"},{"key":"3988_CR109","unstructured":"Shang, F., Fu, J., Yang, Y., Huang, H., Liu, J., Ma, L.: Synfundus-1m: a high-quality million-scale synthetic fundus images dataset with fifteen types of annotation. arXiv preprint arXiv:2312.00377 (2023)"},{"issue":"3","key":"3988_CR110","doi-asserted-by":"publisher","first-page":"996","DOI":"10.1109\/TMI.2020.3043495","volume":"40","author":"Z Shen","year":"2021","unstructured":"Shen, Z., Fu, H., Shen, J., Shao, L.: Modeling and enhancing low-quality retinal fundus images. IEEE Trans. Med. Imaging 40(3), 996\u20131006 (2021). https:\/\/doi.org\/10.1109\/TMI.2020.3043495","journal-title":"IEEE Trans. Med. Imaging"},{"key":"3988_CR111","volume":"10","author":"B Sheng","year":"2022","unstructured":"Sheng, B., Chen, X., Li, T., Ma, T., Yang, Y., Bi, L., Zhang, X.: An overview of artificial intelligence in diabetic retinopathy and other ocular diseases. Front. Public Health 10, 971943 (2022)","journal-title":"Front. Public Health"},{"key":"3988_CR112","unstructured":"Shi, D., Zhang, W., Chen, X., Liu, Y., Yang, J., Huang, S., Tham, Y.C., Zheng, Y., He, M.: Eyefound: a multimodal generalist foundation model for ophthalmic imaging. arXiv preprint arXiv:2405.11338 (2024)"},{"issue":"4","key":"3988_CR113","volume":"3","author":"D Shi","year":"2023","unstructured":"Shi, D., Zhang, W., He, S., Chen, Y., Song, F., Liu, S., Wang, R., Zheng, Y., He, M.: Translation of color fundus photography into fluorescein angiography using deep learning for enhanced diabetic retinopathy screening. Ophthalmol. Sci. 3(4), 100401 (2023)","journal-title":"Ophthalmol. Sci."},{"issue":"8022","key":"3988_CR114","doi-asserted-by":"crossref","first-page":"755","DOI":"10.1038\/s41586-024-07566-y","volume":"631","author":"I Shumailov","year":"2024","unstructured":"Shumailov, I., Shumaylov, Z., Zhao, Y., Papernot, N., Anderson, R., Gal, Y.: Ai models collapse when trained on recursively generated data. Nature 631(8022), 755\u2013759 (2024)","journal-title":"Nature"},{"key":"3988_CR115","doi-asserted-by":"crossref","DOI":"10.1016\/j.media.2024.103357","volume":"99","author":"J Silva-Rodriguez","year":"2025","unstructured":"Silva-Rodriguez, J., Chakor, H., Kobbi, R., Dolz, J., Ayed, I.B.: A foundation language-image model of the retina (flair): encoding expert knowledge in text supervision. Med. Image Anal. 99, 103357 (2025)","journal-title":"Med. Image Anal."},{"key":"3988_CR116","doi-asserted-by":"crossref","DOI":"10.1016\/j.compbiomed.2024.109645","volume":"186","author":"VB Sivaraman","year":"2025","unstructured":"Sivaraman, V.B., Imran, M., Wei, Q., Muralidharan, P., Tamplin, M.R., Grumbach, I.M., Kardon, R.H., Wang, J.K., Zhou, Y., Shao, W.: Retinaregnet: a zero-shot approach for retinal image registration. Comput. Biol. Med. 186, 109645 (2025)","journal-title":"Comput. Biol. Med."},{"key":"3988_CR117","unstructured":"Sohl-Dickstein, J., Weiss, E., Maheswaranathan, N., Ganguli, S.: Deep unsupervised learning using nonequilibrium thermodynamics. In: International Conference on Machine Learning, pp. 2256\u20132265. PMLR (2015)"},{"issue":"4","key":"3988_CR118","doi-asserted-by":"crossref","first-page":"192","DOI":"10.1016\/j.aopr.2023.11.001","volume":"3","author":"F Song","year":"2023","unstructured":"Song, F., Zhang, W., Zheng, Y., Shi, D., He, M.: A deep learning model for generating fundus autofluorescence images from color fundus photography. Adv. Ophthalmol. Pract. Res. 3(4), 192\u2013198 (2023)","journal-title":"Adv. Ophthalmol. Pract. Res."},{"key":"3988_CR119","unstructured":"Song, J., Meng, C., Ermon, S.: Denoising diffusion implicit models. In: International Conference on Learning Representations (2021). https:\/\/openreview.net\/forum?id=St1giarCHLP"},{"issue":"10","key":"3988_CR120","doi-asserted-by":"crossref","first-page":"1335","DOI":"10.1136\/bjo-2024-325458","volume":"108","author":"SC Sonmez","year":"2024","unstructured":"Sonmez, S.C., Sevgi, M., Antaki, F., Huemer, J., Keane, P.A.: Generative artificial intelligence in ophthalmology: current innovations, future applications and challenges. Br. J. Ophthalmol. 108(10), 1335\u20131340 (2024)","journal-title":"Br. J. Ophthalmol."},{"key":"3988_CR121","doi-asserted-by":"crossref","unstructured":"Sun, Y., Tan, W., Gu, Z., He, R., Chen, S., Pang, M., Yan, B.: A data-efficient strategy for building high-performing medical foundation models. Nat. Biomed. Eng. 1\u201313 (2025)","DOI":"10.1038\/s41551-025-01365-0"},{"issue":"1","key":"3988_CR122","doi-asserted-by":"crossref","first-page":"21580","DOI":"10.1038\/s41598-020-78696-2","volume":"10","author":"A Tavakkoli","year":"2020","unstructured":"Tavakkoli, A., Kamran, S.A., Hossain, K.F., Zuckerbrod, S.L.: A novel deep learning conditional generative adversarial network for producing angiography images from retinal fundus photographs. Sci. Rep. 10(1), 21580 (2020)","journal-title":"Sci. Rep."},{"key":"3988_CR123","unstructured":"Team, G., Anil, R., Borgeaud, S., Alayrac, J.B., Yu, J., Soricut, R., Schalkwyk, J., Dai, A.M., Hauth, A., Millican, K., et\u00a0al.: Gemini: a family of highly capable multimodal models. arXiv preprint arXiv:2312.11805 (2023)"},{"key":"3988_CR124","unstructured":"Tian, Y., Shi, M., Luo, Y., Kouhana, A., Elze, T., Wang, M.: Fairseg: a large-scale medical image segmentation dataset for fairness learning using segment anything model with fair error-bound scaling. arXiv preprint arXiv:2311.02189 (2023)"},{"issue":"2","key":"3988_CR125","doi-asserted-by":"crossref","first-page":"167","DOI":"10.1136\/bjophthalmol-2018-313173","volume":"103","author":"DSW Ting","year":"2019","unstructured":"Ting, D.S.W., Pasquale, L.R., Peng, L., Campbell, J.P., Lee, A.Y., Raman, R., Tan, G.S.W., Schmetterer, L., Keane, P.A., Wong, T.Y.: Artificial intelligence and deep learning in ophthalmology. Br. J. Ophthalmol. 103(2), 167\u2013175 (2019)","journal-title":"Br. J. Ophthalmol."},{"key":"3988_CR126","unstructured":"van\u00a0den Oord, A., Vinyals, O., Kavukcuoglu, K.: Neural discrete representation learning. In: Proceedings of the 31st International Conference on Neural Information Processing Systems, NIPS\u201917, pp. 6309\u20136318. Curran Associates Inc., Red Hook, NY, USA (2017)"},{"key":"3988_CR127","doi-asserted-by":"crossref","unstructured":"Waisberg, E., Ong, J., Kamran, S.A., Masalkhi, M., Paladugu, P., Zaman, N., Lee, A.G., Tavakkoli, A.: Generative artificial intelligence in ophthalmology. Surv. Ophthalmol. (2024)","DOI":"10.1016\/j.survophthal.2024.04.009"},{"key":"3988_CR128","doi-asserted-by":"crossref","unstructured":"Wang, H., Xing, Z., Wu, W., Yang, Y., Tang, Q., Zhang, M., Xu, Y., Zhu, L.: Non-invasive to invasive: Enhancing FFA synthesis from CFP with a benchmark dataset and a novel network. In: Proceedings of the 1st International Workshop on Multimedia Computing for Health and Medicine, pp. 7\u201315 (2024)","DOI":"10.1145\/3688868.3689194"},{"issue":"2","key":"3988_CR129","doi-asserted-by":"crossref","first-page":"609","DOI":"10.1038\/s41591-024-03359-y","volume":"31","author":"J Wang","year":"2025","unstructured":"Wang, J., Wang, K., Yu, Y., Lu, Y., Xiao, W., Sun, Z., Liu, F., Zou, Z., Gao, Y., Yang, L., et al.: Self-improving generative foundation model for synthetic medical image generation and clinical applications. Nat. Med. 31(2), 609\u2013617 (2025)","journal-title":"Nat. Med."},{"key":"3988_CR130","unstructured":"Wang, P., Bai, S., Tan, S., Wang, S., Fan, Z., Bai, J., Chen, K., Liu, X., Wang, J., Ge, W., et\u00a0al.: Qwen2-vl: enhancing vision-language model\u2019s perception of the world at any resolution. arXiv preprint arXiv:2409.12191 (2024)"},{"key":"3988_CR131","doi-asserted-by":"crossref","unstructured":"Wang, T.C., Liu, M.Y., Zhu, J.Y., Tao, A., Kautz, J., Catanzaro, B.: High-resolution image synthesis and semantic manipulation with conditional gans. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (2018)","DOI":"10.1109\/CVPR.2018.00917"},{"key":"3988_CR132","unstructured":"Wang, X., Zhang, X., Luo, Z., Sun, Q., Cui, Y., Wang, J., Zhang, F., Wang, Y., Li, Z., Yu, Q., et\u00a0al.: Emu3: next-token prediction is all you need. arXiv preprint arXiv:2409.18869 (2024)"},{"key":"3988_CR133","doi-asserted-by":"crossref","unstructured":"Wang, Z., Wu, J., Low, C.H., Jin, Y.: Medagent-pro: towards multi-modal evidence-based medical diagnosis via reasoning agentic workflow (2025). arxiv: 2503.18968","DOI":"10.20944\/preprints202503.1751.v2"},{"key":"3988_CR134","unstructured":"Wei, H., Liu, B., Zhang, M., Shi, P., Yuan, W.: Visionclip: An med-aigc based ethical language-image foundation model for generalizable retina image analysis. arXiv preprint arXiv:2403.10823 (2024)"},{"key":"3988_CR135","unstructured":"Wu, J., Fu, R., Fang, H., Zhang, Y., Yang, Y., Xiong, H., Liu, H., Xu, Y.: Medsegdiff: medical image segmentation with diffusion probabilistic model. In: Medical Imaging with Deep Learning, pp. 1623\u20131639. PMLR (2024)"},{"key":"3988_CR136","doi-asserted-by":"crossref","unstructured":"Wu, J., Ji, W., Fu, H., Xu, M., Jin, Y., Xu, Y.: Medsegdiff-v2: diffusion-based medical image segmentation with transformer. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol.\u00a038, pp. 6030\u20136038 (2024)","DOI":"10.1609\/aaai.v38i6.28418"},{"key":"3988_CR137","volume":"182","author":"M Wu","year":"2019","unstructured":"Wu, M., Cai, X., Chen, Q., Ji, Z., Niu, S., Leng, T., Rubin, D.L., Park, H.: Geographic atrophy segmentation in SD-OCT images using synthesized fundus autofluorescence imaging. Comput. Methods Prog. Biomed. 182, 105101 (2019)","journal-title":"Comput. Methods Prog. Biomed."},{"key":"3988_CR138","doi-asserted-by":"crossref","unstructured":"Wu, Y., He, W., Eschweiler, D., Dou, N., Fan, Z., Mi, S., Walter, P., Stegmaier, J.: Retinal oct synthesis with denoising diffusion probabilistic models for layer segmentation. In: 2024 IEEE International Symposium on Biomedical Imaging (ISBI), pp. 1\u20135. IEEE (2024)","DOI":"10.1109\/ISBI56570.2024.10635836"},{"key":"3988_CR139","doi-asserted-by":"crossref","unstructured":"Xie, Y., Qu, J., Xie, H., Wang, T., Lei, B.: Diffdgss: Generalizable retinal image segmentation with deterministic representation from diffusion models. In: International Conference on Medical Image Computing and Computer-Assisted Intervention, pp. 166\u2013176. Springer, Berlin (2024)","DOI":"10.1007\/978-3-031-72111-3_16"},{"key":"3988_CR140","volume":"10","author":"F Xu","year":"2022","unstructured":"Xu, F., Yu, X., Gao, Y., Ning, X., Huang, Z., Wei, M., Zhai, W., Zhang, R., Wang, S., Li, J.: Predicting oct images of short-term response to anti-vegf treatment for retinal vein occlusion using generative adversarial network. Front. Bioeng. Biotechnol. 10, 914964 (2022)","journal-title":"Front. Bioeng. Biotechnol."},{"issue":"10","key":"3988_CR141","doi-asserted-by":"crossref","first-page":"2838","DOI":"10.1038\/s41591-024-03113-4","volume":"30","author":"Y Yang","year":"2024","unstructured":"Yang, Y., Zhang, H., Gichoya, J.W., Katabi, D., Ghassemi, M.: The limits of fair medical imaging ai in real-world generalization. Nat. Med. 30(10), 2838\u20132848 (2024)","journal-title":"Nat. Med."},{"key":"3988_CR142","unstructured":"Yang, Z., Li, L., Lin, K., Wang, J., Lin, C.C., Liu, Z., Wang, L.: The dawn of lmms: preliminary explorations with gpt-4v (ision). arXiv preprint arXiv:2309.174219(1), 1 (2023)"},{"key":"3988_CR143","volume":"197","author":"TK Yoo","year":"2020","unstructured":"Yoo, T.K., Ryu, I.H., Kim, J.K., Lee, I.S., Kim, J.S., Kim, H.K., Choi, J.Y.: Deep learning can generate traditional retinal fundus photographs using ultra-widefield images via generative adversarial networks. Comput. Methods Progr. Biomed. 197, 105761 (2020)","journal-title":"Comput. Methods Progr. Biomed."},{"issue":"1","key":"3988_CR144","doi-asserted-by":"crossref","first-page":"6","DOI":"10.1186\/s40662-022-00277-3","volume":"9","author":"A You","year":"2022","unstructured":"You, A., Kim, J.K., Ryu, I.H., Yoo, T.K.: Application of generative adversarial networks (GAN) for ophthalmology image domains: a survey. Eye Vis. 9(1), 6 (2022)","journal-title":"Eye Vis."},{"key":"3988_CR145","unstructured":"Yu, C., Fang, H., Wang, H., Deng, T., Du, Q., Xu, Y., Yang, W.: Rethinking diffusion-based image generators for fundus fluorescein angiography synthesis on limited data. arXiv preprint arXiv:2412.12778 (2024)"},{"key":"3988_CR146","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1186\/s12938-018-0620-3","volume":"18","author":"Z Yu","year":"2019","unstructured":"Yu, Z., Xiang, Q., Meng, J., Kou, C., Ren, Q., Lu, Y.: Retinal image synthesis from multiple-landmarks input with generative adversarial networks. Biomed. Eng. Online 18, 1\u201315 (2019)","journal-title":"Biomed. Eng. Online"},{"key":"3988_CR147","doi-asserted-by":"crossref","unstructured":"Zhang, L., Rao, A., Agrawala, M.: Adding conditional control to text-to-image diffusion models. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 3836\u20133847 (2023)","DOI":"10.1109\/ICCV51070.2023.00355"},{"key":"3988_CR148","doi-asserted-by":"crossref","unstructured":"Zhang, L., Wu, F., Bronik, K., Papiez, B.W.: Diffuseg: domain-driven diffusion for medical image segmentation. IEEE J. Biomed. Health Inform. (2025)","DOI":"10.1109\/JBHI.2025.3526806"},{"key":"3988_CR149","doi-asserted-by":"crossref","unstructured":"Zhang, W., Huang, S., Yang, J., Chen, R., Ge, Z., Zheng, Y., Shi, D., He, M.: Fundus2video: Cross-modal angiography video generation from static fundus photography with clinical knowledge guidance. In: International Conference on Medical Image Computing and Computer-Assisted Intervention, pp. 689\u2013699. Springer, Berlin (2024)","DOI":"10.1007\/978-3-031-72378-0_64"},{"key":"3988_CR150","unstructured":"Zhang, W., Yang, J., Chen, R., Huang, S., Xu, P., Chen, X., Lu, S., Cao, H., He, M., Shi, D.: Fundus to fluorescein angiography video generation as a retinal generative foundation model. arXiv preprint arXiv:2410.13242 (2024)"},{"key":"3988_CR151","doi-asserted-by":"crossref","unstructured":"Zhang, Y., Huang, K., Yang, X., Ma, X., Wu, J., Wang, N., Wang, X., Heng, P.A.: Coarse-to-fine latent diffusion model for glaucoma forecast on sequential fundus images. In: International Conference on Medical Image Computing and Computer-Assisted Intervention, pp. 166\u2013176. Springer, Berlin (2024)","DOI":"10.1007\/978-3-031-72086-4_16"},{"key":"3988_CR152","unstructured":"Zhang, Y.F., Yu, T., Tian, H., Fu, C., Li, P., Zeng, J., Xie, W., Shi, Y., Zhang, H., Wu, J., et\u00a0al.: Mm-rlhf: the next step forward in multimodal llm alignment. arXiv preprint arXiv:2502.10391 (2025)"},{"key":"3988_CR153","doi-asserted-by":"crossref","first-page":"14","DOI":"10.1016\/j.media.2018.07.001","volume":"49","author":"H Zhao","year":"2018","unstructured":"Zhao, H., Li, H., Maurer-Stroh, S., Cheng, L.: Synthesizing retinal and neuronal images with generative adversarial nets. Med. Image Anal. 49, 14\u201326 (2018)","journal-title":"Med. Image Anal."},{"issue":"1","key":"3988_CR154","doi-asserted-by":"crossref","first-page":"557","DOI":"10.1109\/JBHI.2023.3302989","volume":"28","author":"P Zhao","year":"2023","unstructured":"Zhao, P., Song, X., Xi, X., Nie, X., Meng, X., Qu, Y., Yin, Y.: Biomarkers-aware asymmetric bibranch GAN with adaptive memory batch normalization for prediction of anti-VEGF treatment response in neovascular age-related macular degeneration. IEEE J. Biomed. Health Inform. 28(1), 557\u2013568 (2023)","journal-title":"IEEE J. Biomed. Health Inform."},{"key":"3988_CR155","unstructured":"Zhao, Y., Zhang, H., Si, S., Nan, L., Tang, X., Cohan, A.: Large language models are effective table-to-text generators, evaluators, and feedback providers. arXiv preprint arXiv:2305.14987 (2023)"},{"key":"3988_CR156","doi-asserted-by":"crossref","unstructured":"Zhao, Z., Yang, J., Faghihroohi, S., Zhao, Y., Zapp, D., Huang, K., Navab, N., Nasseri, M.A.: Extrapolating prospective glaucoma fundus images through diffusion in irregular longitudinal sequences. In: 2024 IEEE International Conference on Bioinformatics and Biomedicine (BIBM), pp. 4032\u20134035. IEEE (2024)","DOI":"10.1109\/BIBM62325.2024.10822368"},{"key":"3988_CR157","volume":"26","author":"Z Zhao","year":"2024","unstructured":"Zhao, Z., Zhang, W., Chen, X., Song, F., Gunasegaram, J., Huang, W., Shi, D., He, M., Liu, N.: Slit lamp report generation and question answering: Development and validation of a multimodal transformer model with large language model integration. J. Med. Internet Res. 26, e54047 (2024)","journal-title":"J. Med. Internet Res."},{"key":"3988_CR158","doi-asserted-by":"crossref","unstructured":"Zhao, Z., Zhao, Y., Yang, J., Huang, K., Navab, N., Nasseri, M.A.: Kldd: Kalman filter based linear deformable diffusion model in retinal image segmentation. In: 2024 IEEE International Conference on Bioinformatics and Biomedicine (BIBM), pp. 1763\u20131766. IEEE (2024)","DOI":"10.1109\/BIBM62325.2024.10822342"},{"key":"3988_CR159","unstructured":"Zhou, H., Li, X., Wang, R., Cheng, M., Zhou, T., Hsieh, C.J.: R1-zero\u2019s\" aha moment\" in visual reasoning on a 2b non-sft model. arXiv preprint arXiv:2503.05132 (2025)"},{"key":"3988_CR160","unstructured":"Zhou, H., Zhu, L., Zhou, Y.: Distribution aligned diffusion and prototype-guided network for unsupervised domain adaptive segmentation. arXiv preprint arXiv:2303.12313 (2023)"},{"issue":"7981","key":"3988_CR161","doi-asserted-by":"crossref","first-page":"156","DOI":"10.1038\/s41586-023-06555-x","volume":"622","author":"Y Zhou","year":"2023","unstructured":"Zhou, Y., Chia, M.A., Wagner, S.K., Ayhan, M.S., Williamson, D.J., Struyven, R.R., Liu, T., Xu, M., Lozano, M.G., Woodward-Court, P., et al.: A foundation model for generalizable disease detection from retinal images. Nature 622(7981), 156\u2013163 (2023)","journal-title":"Nature"},{"key":"3988_CR162","doi-asserted-by":"crossref","unstructured":"Zhu, J.Y., Park, T., Isola, P., Efros, A.A.: Unpaired image-to-image translation using cycle-consistent adversarial networks. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 2223\u20132232 (2017)","DOI":"10.1109\/ICCV.2017.244"}],"container-title":["The Visual Computer"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00371-025-03988-5.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00371-025-03988-5\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00371-025-03988-5.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,9,15]],"date-time":"2025-09-15T09:37:48Z","timestamp":1757929068000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00371-025-03988-5"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,5,24]]},"references-count":162,"journal-issue":{"issue":"12","published-print":{"date-parts":[[2025,9]]}},"alternative-id":["3988"],"URL":"https:\/\/doi.org\/10.1007\/s00371-025-03988-5","relation":{},"ISSN":["0178-2789","1432-2315"],"issn-type":[{"value":"0178-2789","type":"print"},{"value":"1432-2315","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,5,24]]},"assertion":[{"value":"6 May 2025","order":1,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"24 May 2025","order":2,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}]}}