{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,23]],"date-time":"2026-07-23T09:00:53Z","timestamp":1784797253300,"version":"3.55.0"},"reference-count":58,"publisher":"Springer Science and Business Media LLC","issue":"4","license":[{"start":{"date-parts":[[2026,6,17]],"date-time":"2026-06-17T00:00:00Z","timestamp":1781654400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,6,17]],"date-time":"2026-06-17T00:00:00Z","timestamp":1781654400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Mach. Intell. Res."],"published-print":{"date-parts":[[2026,8]]},"DOI":"10.1007\/s11633-025-1613-x","type":"journal-article","created":{"date-parts":[[2026,6,17]],"date-time":"2026-06-17T06:47:05Z","timestamp":1781678825000},"page":"916-930","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Countering Adversarial Attacks with Multimodal Image Fusion"],"prefix":"10.1007","volume":"23","author":[{"ORCID":"https:\/\/orcid.org\/0000-0001-9837-602X","authenticated-orcid":false,"given":"Dong","family":"Yu","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-1161-8995","authenticated-orcid":false,"given":"Chunjie","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-1630-6058","authenticated-orcid":false,"given":"Xiaoyu","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-2219-1317","authenticated-orcid":false,"given":"Gaopeng","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yao","family":"Zhao","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,6,17]]},"reference":[{"issue":"3","key":"1613_CR1","doi-asserted-by":"publisher","first-page":"417","DOI":"10.1007\/s11633-025-1546-4","volume":"22","author":"Y Zhao","year":"2025","unstructured":"Y. Zhao, R. Zhang, W. Li, L. Li. Assessing and understanding creativity in large language models. Machine Intelligence Research, vol. 22, no. 3, pp. 417\u2013436, 2025. DOI: https:\/\/doi.org\/10.1007\/s11633-025-1546-4.","journal-title":"Machine Intelligence Research"},{"issue":"5","key":"1613_CR2","doi-asserted-by":"publisher","first-page":"888","DOI":"10.1007\/s11633-024-1502-8","volume":"21","author":"T Sun","year":"2024","unstructured":"T. Sun, X. Zhang, Z. He, P. Li, Q. Cheng, X. Liu, H. Yan, Y. Shao, Q. Tang, S. Zhang, X. Zhao, K. Chen, Y. Zheng, Z. Zhou, R. Li, J. Zhan, Y. Zhou, L. Li, X. Yang, L. Wu, Z. Yin, X. Huang, Y. G. Jiang, X. Qiu. Moss: An open conversational large language model. Machine Intelligence Research, vol. 21, no. 5, pp. 888\u2013905, 2024. DOI: https:\/\/doi.org\/10.1007\/s11633-024-1502-8.","journal-title":"Machine Intelligence Research"},{"key":"1613_CR3","doi-asserted-by":"publisher","first-page":"5345","DOI":"10.1109\/TIFS.2024.3396064","volume":"19","author":"W Guan","year":"2024","unstructured":"W. Guan, W. Wang, J. Dong, B. Peng. Improving generalization of deepfake detectors by imposing gradient regularization. IEEE Transactions on Information Forensics and Security, vol. 19, pp. 5345\u20135356, 2024. DOI: https:\/\/doi.org\/10.1109\/TIFS.2024.3396064.","journal-title":"IEEE Transactions on Information Forensics and Security"},{"key":"1613_CR4","doi-asserted-by":"publisher","first-page":"6894","DOI":"10.1109\/CVPR52729.2023.00666","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Y Lyu","year":"2023","unstructured":"Y. Lyu, T. Lin, F. Li, D. He, J. Dong, T. Tan. Notice of removal: DeltaEdit: Exploring text-free training for text-driven image manipulation. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Vancouver, Canada, pp. 6894\u20136903, 2023. DOI: https:\/\/doi.org\/10.1109\/CVPR52729.2023.00666."},{"key":"1613_CR5","doi-asserted-by":"publisher","first-page":"1796","DOI":"10.23919\/APSIPA.2018.8659464","volume-title":"Proceedings of Asia-Pacific Signal and Information Processing Association Annual Summit and Conference","author":"K Ye","year":"2018","unstructured":"K. Ye, J. Dong, W. Wang, B. Peng, T. Tan. Feature pyramid deep matching and localization network for image forensics. In Proceedings of Asia-Pacific Signal and Information Processing Association Annual Summit and Conference, Honolulu, USA, pp. 1796\u20131802, 2018. DOI: https:\/\/doi.org\/10.23919\/APSIPA.2018.8659464."},{"key":"1613_CR6","doi-asserted-by":"publisher","first-page":"5256","DOI":"10.1109\/TIFS.2025.3573161","volume":"20","author":"W Guan","year":"2025","unstructured":"W. Guan, W. Wang, B. Peng, Z. He, J. Dong, H. Cheng. Noise-informed diffusion-generated image detection with anomaly attention. IEEE Transactions on Information Forensics and Security, vol. 20, pp. 5256\u20135268, 2025. DOI: https:\/\/doi.org\/10.1109\/TIFS.2025.3573161.","journal-title":"IEEE Transactions on Information Forensics and Security"},{"key":"1613_CR7","doi-asserted-by":"publisher","first-page":"48126","DOI":"10.1109\/ACCESS.2024.3381611","volume":"12","author":"A Golda","year":"2024","unstructured":"A. Golda, K. Mekonen, A. Pandey, A. Singh, V. Hassija, V. Chamola, B. Sikdar. Privacy and security concerns in generative AI: A comprehensive survey. IEEE Access, vol. 12, pp. 48126\u201348144, 2024. DOI: https:\/\/doi.org\/10.1109\/ACCESS.2024.3381611.","journal-title":"IEEE Access"},{"key":"1613_CR8","doi-asserted-by":"publisher","unstructured":"Y. Chen, P. Esmaeilzadeh. Generative AI in medical practice: In-depth exploration of privacy and security challenges. Journal of Medical Internet Research, vol. 26, Article number e53008, 2024. DOI: https:\/\/doi.org\/10.2196\/53008.","DOI":"10.2196\/53008"},{"issue":"6","key":"1613_CR9","doi-asserted-by":"publisher","first-page":"1214","DOI":"10.1007\/s11633-024-1507-3","volume":"21","author":"Z Zhang","year":"2024","unstructured":"Z. Zhang, G. Xiao, Y. Li, T. Lv, F. Qi, Z. Liu, Y. Wang, X. Jiang, M. Sun. Correction to: Red alarm for pretrained models: Universal vulnerability to neuron-level backdoor attacks. Machine Intelligence Research, vol. 21, no. 6, pp. 1214, 2024. DOI: https:\/\/doi.org\/10.1007\/s11633-024-1507-3.","journal-title":"Machine Intelligence Research"},{"issue":"2","key":"1613_CR10","doi-asserted-by":"publisher","first-page":"217","DOI":"10.1007\/s11633-023-1416-x","volume":"21","author":"B Cao","year":"2024","unstructured":"B. Cao, H. Lin, X. Han, L. Sun. The life cycle of knowledge in big language models: A survey. Machine Intelligence Research, vol. 21, no. 2, pp. 217\u2013238, 2024. DOI: https:\/\/doi.org\/10.1007\/s11633-023-1416-x.","journal-title":"Machine Intelligence Research"},{"issue":"6","key":"1613_CR11","doi-asserted-by":"publisher","first-page":"1011","DOI":"10.1007\/s11633-024-1510-8","volume":"21","author":"E Dai","year":"2024","unstructured":"E. Dai, T. Zhao, H. Zhu, J. Xu, Z. Guo, H. Liu, J. Tang, S. Wang. A comprehensive survey on trustworthy graph neural networks: Privacy, robustness, fairness, and explainability. Machine Intelligence Research, vol. 21, no. 6, pp. 1011\u20131061, 2024. DOI: https:\/\/doi.org\/10.1007\/s11633-024-1510-8.","journal-title":"Machine Intelligence Research"},{"issue":"9","key":"1613_CR12","doi-asserted-by":"publisher","first-page":"2805","DOI":"10.1109\/TNNLS.2018.2886017","volume":"30","author":"X Yuan","year":"2019","unstructured":"X. Yuan, P. He, Q. Zhu, X. Li. Adversarial examples: Attacks and defenses for deep learning. IEEE Transactions on Neural Networks and Learning Systems, vol. 30, no. 9, pp. 2805\u20132824, 2019. DOI: https:\/\/doi.org\/10.1109\/TNNLS.2018.2886017.","journal-title":"IEEE Transactions on Neural Networks and Learning Systems"},{"key":"1613_CR13","volume-title":"Generative adversarial patches for physical attacks on cross-modal pedestrian re-identification","author":"Y Su","year":"2024","unstructured":"Y. Su, H. Li, M. Gong. Generative adversarial patches for physical attacks on cross-modal pedestrian re-identification, [Online], Available: https:\/\/arxiv.org\/abs\/2410.20097, 2024."},{"key":"1613_CR14","doi-asserted-by":"publisher","first-page":"38","DOI":"10.1016\/j.knosys.2019.05.017","volume":"180","author":"P Hu","year":"2019","unstructured":"P. Hu, D. Peng, X. Wang, Y. Xiang. Multimodal adversarial network for cross-modal retrieval. Knowledge-Based Systems, vol. 180, pp. 38\u201350, 2019. DOI: https:\/\/doi.org\/10.1016\/j.knosys.2019.05.017.","journal-title":"Knowledge-Based Systems"},{"key":"1613_CR15","doi-asserted-by":"publisher","first-page":"1371","DOI":"10.1109\/WACV51458.2022.00144","volume-title":"Proceedings of IEEE\/CVF Winter Conference on Applications of Computer Vision","author":"S Wang","year":"2022","unstructured":"S. Wang, T. Wu, A. Chakrabarti, Y. Vorobeychik. Adversarial robustness of deep sensor fusion models. In Proceedings of IEEE\/CVF Winter Conference on Applications of Computer Vision, Waikoloa, USA, pp. 1371\u20131380, 2022. DOI: https:\/\/doi.org\/10.1109\/WACV51458.2022.00144."},{"key":"1613_CR16","doi-asserted-by":"publisher","first-page":"26866","DOI":"10.1109\/CVPR52733.2024.02538","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Z Gao","year":"2024","unstructured":"Z. Gao, X. Jiang, X. Xu, F. Shen, Y. Li, H. T. Shen. Embracing unimodal aleatoric uncertainty for robust multimodal fusion. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Seattle, USA, pp. 26866\u201326875, 2024. DOI: https:\/\/doi.org\/10.1109\/CVPR52733.2024.02538."},{"key":"1613_CR17","volume-title":"Revisiting the adversarial robustness of vision language models: A multimodal perspective","author":"W Zhou","year":"2024","unstructured":"W. Zhou, S. Bai, D. P. Mandic, Q. Zhao, B. Chen. Revisiting the adversarial robustness of vision language models: A multimodal perspective, [Online], Available: https:\/\/arxiv.org\/abs\/2404.19287, 2024."},{"key":"1613_CR18","doi-asserted-by":"publisher","first-page":"17161","DOI":"10.1109\/CVPR52688.2022.01667","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Y Li","year":"2022","unstructured":"Y. Li, A. W. Yu, T. Meng, B. Caine, J. Ngiam, D. Peng, J. Shen, Y. Lu, D. Zhou, Q. V. Le, A. Yuille, M. Tan. DeepFusion: Lidar-camera deep fusion for multi-modal 3D object detection. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, New Orleans, USA, pp. 17161\u201317170, 2022. DOI: https:\/\/doi.org\/10.1109\/CVPR52688.2022.01667."},{"key":"1613_CR19","doi-asserted-by":"publisher","first-page":"5597","DOI":"10.1109\/CVPR46437.2021.00555","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Y Tian","year":"2021","unstructured":"Y. Tian, C. Xu. Can audio-visual integration strengthen robustness under multimodal attacks? In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Nashville, USA, pp. 5597\u20135607, 2021. DOI: https:\/\/doi.org\/10.1109\/CVPR46437.2021.00555."},{"key":"1613_CR20","doi-asserted-by":"publisher","first-page":"3706","DOI":"10.1145\/3581783.3611928","volume-title":"Proceedings of the 31st ACM International Conference on Multimedia","author":"Z Liu","year":"2023","unstructured":"Z. Liu, J. Liu, B. Zhang, L. Ma, X. Fan, R. Liu. PAIF: Perception-aware infrared-visible image fusion for attack-tolerant semantic segmentation. In Proceedings of the 31st ACM International Conference on Multimedia, Ottawa, Canada, pp. 3706\u20133714, 2023. DOI: https:\/\/doi.org\/10.1145\/3581783.3611928."},{"key":"1613_CR21","doi-asserted-by":"publisher","first-page":"4770","DOI":"10.1609\/aaai.v39i5.32504","volume-title":"Proceedings of the 39th AAAI Conference on Artificial Intelligence","author":"J Li","year":"2025","unstructured":"J. Li, H. Yu, J. Chen, X. Ding, J. Wang, J. Liu, B. Zou, H. Ma. A2RNet: Adversarial attack resilient network for robust infrared and visible image fusion. In Proceedings of the 39th AAAI Conference on Artificial Intelligence, Philadelphia, USA, pp. 4770\u20134778, 2025. DOI: https:\/\/doi.org\/10.1609\/aaai.v39i5.32504."},{"key":"1613_CR22","volume-title":"Proceedings of the 2nd International Conference on Learning Representations","author":"C Szegedy","year":"2014","unstructured":"C. Szegedy, W. Zaremba, I. Sutskever, J. Bruna, D. Erhan, I. J. Goodfellow, R. Fergus. Intriguing properties of neural networks. In Proceedings of the 2nd International Conference on Learning Representations, Banff, Canada, 2014."},{"key":"1613_CR23","doi-asserted-by":"publisher","first-page":"5775","DOI":"10.1609\/aaai.v39i6.32616","volume-title":"Proceedings of the 39th AAAI Conference on Artificial Intelligence","author":"J Long","year":"2025","unstructured":"J. Long, Z. Xu, T. Jiang, W. Yao, S. Jia, C. Ma, X. Chen. Robust SAM: On the adversarial robustness of vision foundation models. In Proceedings of the 39th AAAI Conference on Artificial Intelligence, Philadelphia, USA, pp. 5775\u20135783, 2025. DOI: https:\/\/doi.org\/10.1609\/aaai.v39i6.32616."},{"issue":"9","key":"1613_CR24","doi-asserted-by":"publisher","first-page":"15706","DOI":"10.1109\/TNNLS.2025.3561225","volume":"36","author":"K N T Nguyen","year":"2025","unstructured":"K. N. T. Nguyen, W. Zhang, K. Lu, Y. H. Wu, X. Zheng, H. L. Tan, L. Zhen. A survey and evaluation of adversarial attacks in object detection. IEEE Transactions on Neural Networks and Learning Systems, vol. 36, no. 9, pp. 15706\u201315722, 2025. DOI: https:\/\/doi.org\/10.1109\/TNNLS.2025.3561225.","journal-title":"IEEE Transactions on Neural Networks and Learning Systems"},{"key":"1613_CR25","volume-title":"Proceedings of the 3rd International Conference on Learning Representations","author":"I J Goodfellow","year":"2015","unstructured":"I. J. Goodfellow, J. Shlens, C. Szegedy. Explaining and harnessing adversarial examples. In Proceedings of the 3rd International Conference on Learning Representations, San Diego, USA, 2015."},{"key":"1613_CR26","volume-title":"Proceedings of the 5th International Conference on Learning Representations","author":"A Kurakin","year":"2017","unstructured":"A. Kurakin, I. J. Goodfellow, S. Bengio. Adversarial machine learning at scale. In Proceedings of the 5th International Conference on Learning Representations, Toulon, France, 2017."},{"key":"1613_CR27","volume-title":"Proceedings of the 6th International Conference on Learning Representations","author":"F Tram\u00e8r","year":"2018","unstructured":"F. Tram\u00e8r, A. Kurakin, N. Papernot, I. J. Goodfellow, D. Boneh, P. D. McDaniel. Ensemble adversarial training: Attacks and defenses. In Proceedings of the 6th International Conference on Learning Representations, Vancouver, Canada, 2018."},{"key":"1613_CR28","volume-title":"Proceedings of the 32nd International Conference on Neural Information Processing Systems","author":"H Zhang","year":"2019","unstructured":"H. Zhang, J. Wang. Defense against adversarial attacks using feature scattering-based adversarial training. In Proceedings of the 32nd International Conference on Neural Information Processing Systems, Vancouver, Canada, 2019."},{"key":"1613_CR29","volume-title":"Proceedings of the 8th International Conference on Learning Representations","author":"T Wu","year":"2020","unstructured":"T. Wu, L. Tong, Y. Vorobeychik. Defending against physically realizable attacks on image classification. In Proceedings of the 8th International Conference on Learning Representations, Addis Ababa, Ethiopia, 2020."},{"key":"1613_CR30","volume-title":"Proceedings of the 3rd International Conference on Learning Representations","author":"S Gu","year":"2015","unstructured":"S. Gu, L. Rigazio. Towards deep neural network architectures robust to adversarial examples. In Proceedings of the 3rd International Conference on Learning Representations, San Diego, USA, 2015."},{"key":"1613_CR31","doi-asserted-by":"publisher","first-page":"99","DOI":"10.1016\/j.inffus.2019.07.011","volume":"54","author":"Y Zhang","year":"2020","unstructured":"Y. Zhang, Y. Liu, P. Sun, H. Yan, X. Zhao, L. Zhang. IFCNN: A general image fusion framework based on convolutional neural network. Information Fusion, vol. 54, pp. 99\u2013118, 2020. DOI: https:\/\/doi.org\/10.1016\/j.inffus.2019.07.011.","journal-title":"Information Fusion"},{"key":"1613_CR32","doi-asserted-by":"publisher","unstructured":"D. Yu, S. Lin, X. Lu, B. Wang, D. Li, Y. Wang. A multi-band image synchronous fusion method based on saliency. Infrared Physics & Technology, vol. 127, Article number 104466, 2022. DOI: https:\/\/doi.org\/10.1016\/j.infrared.2022.104466.","DOI":"10.1016\/j.infrared.2022.104466"},{"key":"1613_CR33","doi-asserted-by":"publisher","first-page":"2705","DOI":"10.1109\/ICPR.2018.8546006","volume-title":"Proceedings of the 24th International Conference on Pattern Recognition","author":"H Li","year":"2018","unstructured":"H. Li, X. J. Wu, J. Kittler. Infrared and visible image fusion using a deep learning framework. In Proceedings of the 24th International Conference on Pattern Recognition, Beijing, China, pp. 2705\u20132710, 2018. DOI: https:\/\/doi.org\/10.1109\/ICPR.2018.8546006."},{"key":"1613_CR34","doi-asserted-by":"publisher","first-page":"966","DOI":"10.1109\/TMM.2021.3134565","volume":"25","author":"F Zhao","year":"2023","unstructured":"F. Zhao, W. Zhao, H. Lu, Y. Liu, L. Yao, Y. Liu. Depth-distilled multi-focus image fusion. IEEE Transactions on Multimedia, vol. 25, pp. 966\u2013978, 2023. DOI: https:\/\/doi.org\/10.1109\/TMM.2021.3134565.","journal-title":"IEEE Transactions on Multimedia"},{"issue":"1","key":"1613_CR35","doi-asserted-by":"publisher","first-page":"502","DOI":"10.1109\/TPAMI.2020.3012548","volume":"44","author":"H Xu","year":"2022","unstructured":"H. Xu, J. Ma, J. Jiang, X. Guo, H. Ling. U2Fusion: A unified unsupervised image fusion network. IEEE Transactions on Pattern Analysis and Machine Intelligence, vol. 44, no. 1, pp. 502\u2013518, 2022. DOI: https:\/\/doi.org\/10.1109\/TPAMI.2020.3012548.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"issue":"3","key":"1613_CR36","doi-asserted-by":"publisher","first-page":"986","DOI":"10.1109\/TCSVT.2020.2998696","volume":"31","author":"R Nie","year":"2021","unstructured":"R. Nie, J. Cao, D. Zhou, W. Qian. Multi-source information exchange encoding with PCNN for medical image fusion. IEEE Transactions on Circuits and Systems for Video Technology, vol. 31, no. 3, pp. 986\u20131000, 2021. DOI: https:\/\/doi.org\/10.1109\/TCSVT.2020.2998696.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"issue":"5","key":"1613_CR37","doi-asserted-by":"publisher","first-page":"2614","DOI":"10.1109\/TIP.2018.2887342","volume":"28","author":"H Li","year":"2019","unstructured":"H. Li, X. J. Wu. DenseFuse: A fusion approach to infrared and visible images. IEEE Transactions on Image Processing, vol. 28, no. 5, pp. 2614\u20132623, 2019. DOI: https:\/\/doi.org\/10.1109\/TIP.2018.2887342.","journal-title":"IEEE Transactions on Image Processing"},{"key":"1613_CR38","doi-asserted-by":"publisher","unstructured":"H. Xu, X. Wang, J. Ma. DRF: Disentangled representation for visible and infrared image fusion. IEEE Transactions on Instrumentation and Measurement, vol. 70, Article number 5006713, 2021. DOI: https:\/\/doi.org\/10.1109\/TIM.2021.3056645.","DOI":"10.1109\/TIM.2021.3056645"},{"key":"1613_CR39","doi-asserted-by":"publisher","unstructured":"L. Jian, X. Yang, Z. Liu, G. Jeon, M. Gao, D. Chisholm. SEDRFuse: A symmetric encoder-decoder with residual block network for infrared and visible image fusion. IEEE Transactions on Instrumentation and Measurement, vol. 70, Article number 5002215, 2021. DOI: https:\/\/doi.org\/10.1109\/TIM.2020.3022438.","DOI":"10.1109\/TIM.2020.3022438"},{"key":"1613_CR40","doi-asserted-by":"publisher","first-page":"11","DOI":"10.1016\/j.inffus.2018.09.004","volume":"48","author":"J Ma","year":"2019","unstructured":"J. Ma, W. Yu, P. Liang, C. Li, J. Jiang. FusionGAN: A generative adversarial network for infrared and visible image fusion. Information Fusion, vol. 48, pp. 11\u201326, 2019. DOI: https:\/\/doi.org\/10.1016\/j.inffus.2018.09.004.","journal-title":"Information Fusion"},{"key":"1613_CR41","doi-asserted-by":"publisher","first-page":"183","DOI":"10.1016\/j.neucom.2022.02.025","volume":"483","author":"A Song","year":"2022","unstructured":"A. Song, H. Duan, H. Pei, L. Ding. Triple-discriminator generative adversarial network for infrared and visible image fusion. Neurocomputing, vol. 483, pp. 183\u2013194, 2022. DOI: https:\/\/doi.org\/10.1016\/j.neucom.2022.02.025.","journal-title":"Neurocomputing"},{"issue":"8","key":"1613_CR42","doi-asserted-by":"publisher","first-page":"1982","DOI":"10.1109\/TMM.2019.2895292","volume":"21","author":"X Guo","year":"2019","unstructured":"X. Guo, R. Nie, J. Cao, D. Zhou, L. Mei, K. He. FuseGAN: Learning to fuse multi-focus image via conditional generative adversarial network. IEEE Transactions on Multimedia, vol. 21, no. 8, pp. 1982\u20131996, 2019. DOI: https:\/\/doi.org\/10.1109\/TMM.2019.2895292.","journal-title":"IEEE Transactions on Multimedia"},{"key":"1613_CR43","doi-asserted-by":"publisher","first-page":"7203","DOI":"10.1109\/TIP.2020.2999855","volume":"29","author":"H Xu","year":"2020","unstructured":"H. Xu, J. Ma, X. P. Zhang. MEF-GAN: Multi-exposure image fusion via generative adversarial networks. IEEE Transactions on Image Processing, vol. 29, pp. 7203\u20137216, 2020. DOI: https:\/\/doi.org\/10.1109\/TIP.2020.2999855.","journal-title":"IEEE Transactions on Image Processing"},{"key":"1613_CR44","doi-asserted-by":"publisher","first-page":"6000","DOI":"10.5555\/3295222.3295349","volume-title":"Proceedings of the 31st International Conference on Neural Information Processing Systems","author":"A Vaswani","year":"2017","unstructured":"A. Vaswani, N. Shazeer, N. Parmar, J. Uszkoreit, L. Jones, A. N. Gomez, L. Kaiser, I. Polosukhin. Attention is all you need. In Proceedings of the 31st International Conference on Neural Information Processing Systems, Long Beach, USA, pp. 6000\u20136010, 2017. DOI: https:\/\/doi.org\/10.5555\/3295222.3295349."},{"key":"1613_CR45","doi-asserted-by":"publisher","first-page":"2126","DOI":"10.1609\/aaai.v36i2.20109","volume-title":"Proceedings of the 36th AAAI Conference on Artificial Intelligence","author":"L Qu","year":"2022","unstructured":"L. Qu, S. Liu, M. Wang, Z. Song. TransMEF: A transformer-based multi-exposure image fusion framework using self-supervised multi-task learning. In Proceedings of the 36th AAAI Conference on Artificial Intelligence, pp. 2126\u20132134, 2022. DOI: https:\/\/doi.org\/10.1609\/aaai.v36i2.20109."},{"issue":"7","key":"1613_CR46","doi-asserted-by":"publisher","first-page":"1200","DOI":"10.1109\/JAS.2022.105686","volume":"9","author":"J Ma","year":"2022","unstructured":"J. Ma, L. Tang, F. Fan, J. Huang, X. Mei, Y. Ma. Swin-Fusion: Cross-domain long-range learning for general image fusion via swin transformer. IEEE\/CAA Journal of Automatica Sinica, vol. 9, no. 7, pp. 1200\u20131217, 2022. DOI: https:\/\/doi.org\/10.1109\/JAS.2022.105686.","journal-title":"IEEE\/CAA Journal of Automatica Sinica"},{"key":"1613_CR47","doi-asserted-by":"publisher","unstructured":"J. Li, J. Zhu, C. Li, X. Chen, B. Yang. CGTF: Convolution-guided transformer for infrared and visible image fusion. IEEE Transactions on Instrumentation and Measurement, vol. 71, Article number 5012314, 2022. DOI: https:\/\/doi.org\/10.1109\/TIM.2022.3175055.","DOI":"10.1109\/TIM.2022.3175055"},{"key":"1613_CR48","doi-asserted-by":"publisher","first-page":"5134","DOI":"10.1109\/TIP.2022.3193288","volume":"31","author":"W Tang","year":"2022","unstructured":"W. Tang, F. He, Y. Liu, Y. Duan. MATR: Multimodal medical image fusion via multiscale adaptive transformer. IEEE Transactions on Image Processing, vol. 31, pp. 5134\u20135149, 2022. DOI: https:\/\/doi.org\/10.1109\/TIP.2022.3193288.","journal-title":"IEEE Transactions on Image Processing"},{"issue":"5","key":"1613_CR49","doi-asserted-by":"publisher","first-page":"456","DOI":"10.1007\/s11633-022-1375-7","volume":"19","author":"K Y Liu","year":"2022","unstructured":"K. Y. Liu, X. Y. Li, Y. R. Lai, H. Su, J. C. Wang, C. X. Guo, H. Xie, J. S. Guan, Y. Zhou. Denoised internal models: A brain-inspired autoencoder against adversarial attacks. Machine Intelligence Research, vol. 19, no. 5, pp. 456\u2013471, 2022. DOI: https:\/\/doi.org\/10.1007\/s11633-022-1375-7.","journal-title":"Machine Intelligence Research"},{"issue":"9","key":"1613_CR50","doi-asserted-by":"publisher","first-page":"8945","DOI":"10.1109\/TCSVT.2025.3553135","volume":"35","author":"C Xie","year":"2025","unstructured":"C. Xie, X. Zhang, L. Li, Y. Fu, B. Gong, T. Li, K. Zhang. MAT: Multi-range attention transformer for efficient image super-resolution. IEEE Transactions on Circuits and Systems for Video Technology, vol. 35, no. 9, pp. 8945\u20138957, 2025. DOI: https:\/\/doi.org\/10.1109\/TCSVT.2025.3553135.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"key":"1613_CR51","doi-asserted-by":"publisher","first-page":"770","DOI":"10.1109\/CVPR.2016.90","volume-title":"Proceedings of IEEE Conference on Computer Vision and Pattern Recognition","author":"K He","year":"2016","unstructured":"K. He, X. Zhang, S. Ren, J. Sun. Deep residual learning for image recognition. In Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, Las Vegas, USA, pp. 770\u2013778, 2016. DOI: https:\/\/doi.org\/10.1109\/CVPR.2016.90."},{"key":"1613_CR52","doi-asserted-by":"publisher","first-page":"5108","DOI":"10.1109\/IROS.2017.8206396","volume-title":"Proceedings of IEEE\/RSJ International Conference on Intelligent Robots and Systems","author":"Q Ha","year":"2017","unstructured":"Q. Ha, K. Watanabe, T. Karasawa, Y. Ushiku, T. Harada. MFNet: Towards real-time semantic segmentation for autonomous vehicles with multi-spectral scenes. In Proceedings of IEEE\/RSJ International Conference on Intelligent Robots and Systems, Vancouver, Canada, pp. 5108\u20135115, 2017. DOI: https:\/\/doi.org\/10.1109\/IROS.2017.8206396."},{"key":"1613_CR53","doi-asserted-by":"publisher","first-page":"779","DOI":"10.1109\/CVPR.2016.91","volume-title":"Proceedings of IEEE Conference on Computer Vision and Pattern Recognition","author":"J Redmon","year":"2016","unstructured":"J. Redmon, S. Divvala, R. Girshick, A. Farhadi. You only look once: Unified, real-time object detection. In Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, Las Vegas, USA, pp. 779\u2013788, 2016. DOI: https:\/\/doi.org\/10.1109\/CVPR.2016.91."},{"key":"1613_CR54","doi-asserted-by":"publisher","first-page":"5792","DOI":"10.1109\/CVPR52688.2022.00571","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"J Liu","year":"2022","unstructured":"J. Liu, X. Fan, Z. Huang, G. Wu, R. Liu, W. Zhong, Z. Luo. Target-aware dual adversarial learning and a multi-scenario multi-modality benchmark to fuse infrared and visible for object detection. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, New Orleans, USA, pp. 5792\u20135801, 2022. DOI: https:\/\/doi.org\/10.1109\/CVPR52688.2022.00571."},{"key":"1613_CR55","doi-asserted-by":"publisher","DOI":"10.1109\/ICME57554.2024.10688097","volume-title":"Proceedings of IEEE International Conference on Multimedia and Expo","author":"X Guo","year":"2024","unstructured":"X. Guo, Z. Ma, Q. Wang, P. Wei. Towards real-world continuous super-resolution: Benchmark and method. In Proceedings of IEEE International Conference on Multimedia and Expo, Niagara Falls, Canada, 2024. DOI: https:\/\/doi.org\/10.1109\/ICME57554.2024.10688097."},{"key":"1613_CR56","doi-asserted-by":"publisher","unstructured":"H. Li, Y. Xiao, C. Cheng, Z. Shen, X. Song. DePF: A novel fusion approach based on decomposition pooling for infrared and visible images. IEEE Transactions on Instrumentation and Measurement, vol. 72, Article number 5031014, 2023. DOI: https:\/\/doi.org\/10.1109\/TIM.2023.3326252.","DOI":"10.1109\/TIM.2023.3326252"},{"key":"1613_CR57","doi-asserted-by":"publisher","first-page":"4776","DOI":"10.1109\/TMM.2023.3326296","volume":"26","author":"L Tang","year":"2024","unstructured":"L. Tang, Z. Chen, J. Huang, J. Ma. CAMF: An interpretable infrared and visible image fusion network based on class activation mapping. IEEE Transactions on Multimedia, vol. 26, pp. 4776\u20134791, 2024. DOI: https:\/\/doi.org\/10.1109\/TMM.2023.3326296.","journal-title":"IEEE Transactions on Multimedia"},{"key":"1613_CR58","volume-title":"Proceedings of the 39th International Conference on Machine Learning","author":"J Li","year":"2022","unstructured":"J. Li, D. Li, C. Xiong, S. Hoi. BLIP: Bootstrapping language-image pre-training for unified vision-language understanding and generation. In Proceedings of the 39th International Conference on Machine Learning, Baltimore, USA, Article number 162, 2022."}],"container-title":["Machine Intelligence Research"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11633-025-1613-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11633-025-1613-x","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11633-025-1613-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,7,23]],"date-time":"2026-07-23T08:02:46Z","timestamp":1784793766000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11633-025-1613-x"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,6,17]]},"references-count":58,"journal-issue":{"issue":"4","published-print":{"date-parts":[[2026,8]]}},"alternative-id":["1613"],"URL":"https:\/\/doi.org\/10.1007\/s11633-025-1613-x","relation":{},"ISSN":["2731-538X","2731-5398"],"issn-type":[{"value":"2731-538X","type":"print"},{"value":"2731-5398","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,6,17]]},"assertion":[{"value":"5 June 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"3 November 2025","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"17 June 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"The authors declared that they have no conflicts of interest to this work.","order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations of conflict of interest"}}]}}