{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,25]],"date-time":"2026-07-25T15:59:34Z","timestamp":1784995174833,"version":"3.55.0"},"reference-count":66,"publisher":"Springer Science and Business Media LLC","issue":"5","license":[{"start":{"date-parts":[[2023,11,30]],"date-time":"2023-11-30T00:00:00Z","timestamp":1701302400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2023,11,30]],"date-time":"2023-11-30T00:00:00Z","timestamp":1701302400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62161015"],"award-info":[{"award-number":["62161015"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62176081"],"award-info":[{"award-number":["62176081"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int J Comput Vis"],"published-print":{"date-parts":[[2024,5]]},"DOI":"10.1007\/s11263-023-01948-x","type":"journal-article","created":{"date-parts":[[2023,11,30]],"date-time":"2023-11-30T13:02:28Z","timestamp":1701349348000},"page":"1625-1644","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":124,"title":["A Deep Learning Framework for Infrared and Visible Image Fusion Without Strict Registration"],"prefix":"10.1007","volume":"132","author":[{"given":"Huafeng","family":"Li","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Junyu","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yafei","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-2211-3535","authenticated-orcid":false,"given":"Yu","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2023,11,30]]},"reference":[{"issue":"1","key":"1948_CR1","doi-asserted-by":"publisher","first-page":"12","DOI":"10.1016\/j.biosystemseng.2009.02.009","volume":"103","author":"D Bulanon","year":"2009","unstructured":"Bulanon, D., Burks, T., & Alchanatis, V. (2009). Image fusion of visible and thermal images for fruit detection. Biosystems Engineering, 103(1), 12\u201322.","journal-title":"Biosystems Engineering"},{"key":"1948_CR2","doi-asserted-by":"crossref","unstructured":"Carion, N., Massa, F., Synnaeve, G., et al. (2020). End-to-end object detection with transformers. In European Conference on Computer Vision (ECCV) (pp. 213\u2013229). Cham: Springer.","DOI":"10.1007\/978-3-030-58452-8_13"},{"key":"1948_CR3","unstructured":"Chen, C.F., Panda, R., & Fan, Q. (2022). Regionvit: Regional-to-local attention for vision transformers. In Proceedings of the international conference on learning representations (ICLR)."},{"issue":"2","key":"1948_CR4","doi-asserted-by":"publisher","first-page":"193","DOI":"10.1016\/j.inffus.2005.10.001","volume":"8","author":"H Chen","year":"2007","unstructured":"Chen, H., & Varshney, P. K. (2007). A human perception inspired quality metric for image fusion based on regional information. Information Fusion, 8(2), 193\u2013207.","journal-title":"Information Fusion"},{"key":"1948_CR5","doi-asserted-by":"crossref","unstructured":"Chen, H., Wang, Y., Guo, T., et\u00a0al. (2021). Pre-trained image processing transformer. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR) (pp 12,299\u201312,310).","DOI":"10.1109\/CVPR46437.2021.01212"},{"issue":"10","key":"1948_CR6","doi-asserted-by":"publisher","first-page":"1421","DOI":"10.1016\/j.imavis.2007.12.002","volume":"27","author":"Y Chen","year":"2009","unstructured":"Chen, Y., & Blum, R. S. (2009). A new automated quality assessment algorithm for image fusion. Image and Vision Computing, 27(10), 1421\u20131432.","journal-title":"Image and Vision Computing"},{"key":"1948_CR7","unstructured":"Dosovitskiy, A., Beyer, L., & Kolesnikov, A. (2021). An image is worth 16$$\\times $$16 words: Transformers for image recognition at scale. In Proceedings of the international conference on learning representations (ICLR)."},{"issue":"6","key":"1948_CR8","doi-asserted-by":"publisher","first-page":"820","DOI":"10.3390\/s16060820","volume":"16","author":"A Gonz\u00e1lez","year":"2016","unstructured":"Gonz\u00e1lez, A., Fang, Z., Socarras, Y., et al. (2016). Pedestrian detection at day\/night time with visible and fir cameras: A comparison. Sensors, 16(6), 820.","journal-title":"Sensors"},{"key":"1948_CR9","doi-asserted-by":"crossref","unstructured":"Gu, J., Lu, H., Zuo, W. et\u00a0al. (2019). Blind super-resolution with iterative kernel correction. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR) (pp 1604\u20131613)","DOI":"10.1109\/CVPR.2019.00170"},{"issue":"6","key":"1948_CR10","doi-asserted-by":"publisher","first-page":"1771","DOI":"10.1016\/j.patcog.2006.11.010","volume":"40","author":"J Han","year":"2007","unstructured":"Han, J., & Bhanu, B. (2007). Fusion of color and infrared video for moving human detection. Pattern Recognition, 40(6), 1771\u20131784.","journal-title":"Pattern Recognition"},{"key":"1948_CR11","first-page":"15908","volume":"34","author":"K Han","year":"2021","unstructured":"Han, K., Xiao, A., Wu, E., et al. (2021). Transformer in transformer. Advances in Neural Information Processing Systems, 34, 15908\u201315919.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"1948_CR12","doi-asserted-by":"crossref","unstructured":"Hwang, S., Park, J., Kim, N., et\u00a0al. (2015). Multispectral pedestrian detection: Benchmark dataset and baseline. In Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR) (pp 1037\u20131045).","DOI":"10.1109\/CVPR.2015.7298706"},{"key":"1948_CR13","unstructured":"Kingma, D.P., & Ba, J. (2015). Adam: A method for stochastic optimization. In International conference on learning representations (ICLR)."},{"key":"1948_CR14","doi-asserted-by":"crossref","unstructured":"Kristan, M., Leonardis, A., Matas, J., et\u00a0al. (2020). The eighth visual object tracking vot2020 challenge results. In European conference on computer vision (ECCV) (pp 547\u2013601). Springer.","DOI":"10.1007\/978-3-030-68238-5_39"},{"key":"1948_CR15","doi-asserted-by":"crossref","unstructured":"Kumar, P., Mittal, A., & Kumar, P. (2006). Fusion of thermal infrared and visible spectrum video for robust surveillance. In Computer vision, graphics and image processing (pp. 528\u2013539). Springer.","DOI":"10.1007\/11949619_47"},{"issue":"5","key":"1948_CR16","doi-asserted-by":"publisher","first-page":"2614","DOI":"10.1109\/TIP.2018.2887342","volume":"28","author":"H Li","year":"2018","unstructured":"Li, H., & Wu, X. (2018). Densefuse: A fusion approach to infrared and visible images. IEEE Transactions on Image Processing, 28(5), 2614\u20132623.","journal-title":"IEEE Transactions on Image Processing"},{"key":"1948_CR17","doi-asserted-by":"publisher","first-page":"174","DOI":"10.1016\/j.infrared.2016.02.005","volume":"76","author":"H Li","year":"2016","unstructured":"Li, H., Qiu, H., Yu, Z., et al. (2016). Infrared and visible image fusion scheme based on NSCT and low-level visual features. Infrared Physics & Technology, 76, 174\u2013184.","journal-title":"Infrared Physics & Technology"},{"issue":"4","key":"1948_CR18","doi-asserted-by":"publisher","first-page":"1082","DOI":"10.1109\/TIM.2019.2912239","volume":"69","author":"H Li","year":"2020","unstructured":"Li, H., Wang, Y., Yang, Z., et al. (2020). Discriminative dictionary learning-based multiple component decomposition for detail-preserving noisy image fusion. IEEE Transactions on Instrumentation and Measurement, 69(4), 1082\u20131102.","journal-title":"IEEE Transactions on Instrumentation and Measurement"},{"issue":"12","key":"1948_CR19","doi-asserted-by":"publisher","first-page":"9645","DOI":"10.1109\/TIM.2020.3005230","volume":"69","author":"H Li","year":"2020","unstructured":"Li, H., Wu, X. J., & Durrani, T. (2020). Nestfuse: An infrared and visible image fusion architecture based on nest connection and spatial\/channel attention models. IEEE Transactions on Instrumentation and Measurement, 69(12), 9645\u20139656.","journal-title":"IEEE Transactions on Instrumentation and Measurement"},{"key":"1948_CR20","doi-asserted-by":"publisher","first-page":"4070","DOI":"10.1109\/TIP.2021.3069339","volume":"30","author":"H Li","year":"2021","unstructured":"Li, H., Cen, Y., Liu, Y., et al. (2021). Different input resolutions and arbitrary output resolution: A meta learning-based deep framework for infrared and visible image fusion. IEEE Transactions on Image Processing, 30, 4070\u20134083.","journal-title":"IEEE Transactions on Image Processing"},{"key":"1948_CR21","doi-asserted-by":"publisher","first-page":"72","DOI":"10.1016\/j.inffus.2021.02.023","volume":"73","author":"H Li","year":"2021","unstructured":"Li, H., Wu, J., & Kittler, J. (2021). Rfn-nest: An end-to-end residual fusion network for infrared and visible images. Information Fusion, 73, 72\u201386.","journal-title":"Information Fusion"},{"issue":"9","key":"1948_CR22","doi-asserted-by":"publisher","first-page":"11040","DOI":"10.1109\/TPAMI.2023.3268209","volume":"45","author":"H Li","year":"2023","unstructured":"Li, H., Xu, T., Wu, X., et al. (2023). Lrrnet: A novel representation learning guided fusion network for infrared and visible images. IEEE Transactions on Pattern Analysis and Machine Intelligence, 45(9), 11040\u201311052.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"issue":"5","key":"1948_CR23","doi-asserted-by":"publisher","first-page":"347","DOI":"10.1049\/iet-ipr.2014.0311","volume":"9","author":"Y Liu","year":"2015","unstructured":"Liu, Y., Chen, X., Ward, R. K., et al. (2015). Simultaneous image fusion and denoising with adaptive sparse representation. IET Image Processing, 9(5), 347\u2013357.","journal-title":"IET Image Processing"},{"key":"1948_CR24","doi-asserted-by":"publisher","first-page":"174","DOI":"10.1016\/j.inffus.2014.09.004","volume":"24","author":"Y Liu","year":"2015","unstructured":"Liu, Y., Liu, S., & Wang, Z. (2015). A general framework for image fusion based on multi-scale transform and sparse representation. Information Fusion, 24, 174\u2013164.","journal-title":"Information Fusion"},{"issue":"12","key":"1948_CR25","doi-asserted-by":"publisher","first-page":"1882","DOI":"10.1109\/LSP.2016.2618776","volume":"23","author":"Y Liu","year":"2016","unstructured":"Liu, Y., Chen, X., Ward, R. K., et al. (2016). Image fusion with convolutional sparse representation. IEEE Signal Processing Letters, 23(12), 1882\u20131886.","journal-title":"IEEE Signal Processing Letters"},{"key":"1948_CR26","doi-asserted-by":"publisher","first-page":"191","DOI":"10.1016\/j.inffus.2016.12.001","volume":"36","author":"Y Liu","year":"2017","unstructured":"Liu, Y., Chen, X., Peng, H., et al. (2017). Multi-focus image fusion with a deep convolutional neural network. Information Fusion, 36, 191\u2013207.","journal-title":"Information Fusion"},{"key":"1948_CR27","doi-asserted-by":"crossref","unstructured":"Liu, Z., Lin, Y., Cao, Y., et\u00a0al. (2021). Swin transformer: Hierarchical vision transformer using shifted windows. In Proceedings of the IEEE\/CVF international conference on computer vision (ICCV) (pp. 10012\u201310022).","DOI":"10.1109\/ICCV48922.2021.00986"},{"key":"1948_CR28","doi-asserted-by":"publisher","first-page":"153","DOI":"10.1016\/j.inffus.2018.02.004","volume":"45","author":"J Ma","year":"2019","unstructured":"Ma, J., Ma, Y., & Li, C. (2019). Infrared and visible image fusion methods and applications: A survey. Information Fusion, 45, 153\u2013178.","journal-title":"Information Fusion"},{"key":"1948_CR29","doi-asserted-by":"publisher","first-page":"11","DOI":"10.1016\/j.inffus.2018.09.004","volume":"48","author":"J Ma","year":"2019","unstructured":"Ma, J., Yu, W., Liang, P., et al. (2019). Fusiongan: A generative adversarial network for infrared and visible image fusion. Information Fusion, 48, 11\u201326.","journal-title":"Information Fusion"},{"key":"1948_CR30","doi-asserted-by":"publisher","first-page":"4980","DOI":"10.1109\/TIP.2020.2977573","volume":"29","author":"J Ma","year":"2020","unstructured":"Ma, J., Xu, H., Jiang, J., et al. (2020). Ddcgan: A dual-discriminator conditional generative adversarial network for multi-resolution image fusion. IEEE Transactions on Image Processing, 29, 4980\u20134995.","journal-title":"IEEE Transactions on Image Processing"},{"issue":"7","key":"1948_CR31","doi-asserted-by":"publisher","first-page":"1200","DOI":"10.1109\/JAS.2022.105686","volume":"9","author":"J Ma","year":"2022","unstructured":"Ma, J., Tang, L., Fan, F., et al. (2022). Swinfusion: Cross-domain long-range learning for general image fusion via swin transformer. IEEE\/CAA Journal of Automatica Sinica, 9(7), 1200\u20131217.","journal-title":"IEEE\/CAA Journal of Automatica Sinica"},{"issue":"11","key":"1948_CR32","first-page":"2579","volume":"9","author":"L Van der Maaten","year":"2008","unstructured":"Van der Maaten, L., & Hinton, G. (2008). Visualizing data using t-sne. Journal of machine learning research, 9(11), 2579\u20132605.","journal-title":"Journal of machine learning research"},{"key":"1948_CR33","doi-asserted-by":"crossref","unstructured":"Meinhardt, T., Kirillov, A., Leal-Taix\u00e9, L., et\u00a0al. (2021). Trackformer: Multi-object tracking with transformers. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR) (pp. 8844\u20138854).","DOI":"10.1109\/CVPR52688.2022.00864"},{"key":"1948_CR34","unstructured":"Paszke, A., Gross, S., Massa, F., et\u00a0al. (2019). Pytorch: An imperative style, high-performance deep learning library. (pp. 8026\u20138037)."},{"key":"1948_CR35","doi-asserted-by":"crossref","unstructured":"Ram\u00a0Prabhakar, K., Sai\u00a0Srikar, V., & Venkatesh\u00a0Babu, R. (2017). Deepfuse: A deep unsupervised approach for exposure fusion with extreme exposure image pairs. In Proceedings of the IEEE international conference on computer vision (ICCV) (pp. 4714\u20134722).","DOI":"10.1109\/ICCV.2017.505"},{"issue":"1","key":"1948_CR36","doi-asserted-by":"publisher","DOI":"10.1117\/1.2945910","volume":"2","author":"JW Roberts","year":"2008","unstructured":"Roberts, J. W., Van Aardt, J. A., & Ahmed, F. B. (2008). Assessment of image fusion procedures using entropy, image quality, and multispectral classification. Journal of Applied Remote Sensing, 2(1), 023522.","journal-title":"Journal of Applied Remote Sensing"},{"key":"1948_CR37","unstructured":"Simonyan, K., & Zisserman, A. (2015). Very deep convolutional networks for large-scale image recognition. In International conference on learning representations (ICLR)."},{"issue":"3","key":"1948_CR38","doi-asserted-by":"publisher","first-page":"880","DOI":"10.1016\/j.patcog.2007.06.022","volume":"41","author":"R Singh","year":"2008","unstructured":"Singh, R., Vatsa, M., & Noore, A. (2008). Integrated multilevel image fusion and match score fusion of visible and infrared face images for robust face recognition. Pattern Recognition, 41(3), 880\u2013893.","journal-title":"Pattern Recognition"},{"key":"1948_CR39","doi-asserted-by":"crossref","unstructured":"Srinivas, A., Lin, T.Y., Parmar, N., et\u00a0al. (2021). Bottleneck transformers for visual recognition. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR) (pp. 16,519\u201316,529)","DOI":"10.1109\/CVPR46437.2021.01625"},{"key":"1948_CR40","doi-asserted-by":"crossref","unstructured":"Tang, H., Li, Z., Peng, Z., et\u00a0al. (2020). Blockmix: Meta regularization and self-calibrated inference for metric-based meta-learning. In ACM Multimedia (pp. 610\u2013618).","DOI":"10.1145\/3394171.3413884"},{"issue":"108","key":"1948_CR41","first-page":"792","volume":"130","author":"H Tang","year":"2022","unstructured":"Tang, H., Yuan, C., Li, Z., et al. (2022). Learning attention-guided pyramidal features for few-shot fine-grained recognition. Pattern Recognition, 130(108), 792.","journal-title":"Pattern Recognition"},{"issue":"12","key":"1948_CR42","doi-asserted-by":"publisher","first-page":"2121","DOI":"10.1109\/JAS.2022.106082","volume":"9","author":"L Tang","year":"2022","unstructured":"Tang, L., Deng, Y., Ma, Y., et al. (2022). Superfusion: A versatile image registration and fusion network with semantic awareness. IEEE\/CAA Journal of Automatica Sinica, 9(12), 2121\u20132137.","journal-title":"IEEE\/CAA Journal of Automatica Sinica"},{"key":"1948_CR43","doi-asserted-by":"publisher","first-page":"28","DOI":"10.1016\/j.inffus.2021.12.004","volume":"82","author":"L Tang","year":"2022","unstructured":"Tang, L., Yuan, J., & Ma, J. (2022). Image fusion in the loop of high-level vision tasks: A semantic-aware real-time infrared and visible image fusion network. Information Fusion, 82, 28\u201342.","journal-title":"Information Fusion"},{"issue":"7","key":"1948_CR44","doi-asserted-by":"publisher","first-page":"3159","DOI":"10.1109\/TCSVT.2023.3234340","volume":"33","author":"W Tang","year":"2023","unstructured":"Tang, W., He, F., Liu, Y., et al. (2023). Datfuse: Infrared and visible image fusion via dual attention transformer. IEEE Transactions on Circuits and Systems for Video Technology, 33(7), 3159\u20133172.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"key":"1948_CR45","unstructured":"Vaswani, A., Shazeer, N., Parmar, N., et\u00a0al. (2017). Attention is all you need. In Advances in neural information processing systems."},{"key":"1948_CR46","doi-asserted-by":"crossref","unstructured":"Vs, V., Valanarasu, J.M.J., Oza, P., et\u00a0al. (2022). Image fusion transformer. In 2022 IEEE International conference on image processing (ICIP) (pp. 3566\u20133570).","DOI":"10.1109\/ICIP46576.2022.9897280"},{"key":"1948_CR47","doi-asserted-by":"crossref","unstructured":"Wang, D., Liu, J., Fan, X., et\u00a0al. (2022). Unsupervised misaligned infrared and visible image fusion via cross-modality image generation and registration. In Proceedings of the thirty-first international joint conference on artificial intelligence (IJCAI) (pp. 3508\u20133515).","DOI":"10.24963\/ijcai.2022\/487"},{"key":"1948_CR48","doi-asserted-by":"crossref","unstructured":"Wang, X., Yu, K., Dong, C., et\u00a0al. (2018). Recovering realistic texture in image super-resolution by deep spatial feature transform. In Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR) (pp. 606\u2013615).","DOI":"10.1109\/CVPR.2018.00070"},{"issue":"4","key":"1948_CR49","doi-asserted-by":"publisher","first-page":"600","DOI":"10.1109\/TIP.2003.819861","volume":"13","author":"Z Wang","year":"2004","unstructured":"Wang, Z., Bovik, A. C., Sheikh, H. R., et al. (2004). Image quality assessment: From error visibility to structural similarity. IEEE Transactions on Image Processing, 13(4), 600\u2013612.","journal-title":"IEEE Transactions on Image Processing"},{"key":"1948_CR50","doi-asserted-by":"crossref","unstructured":"Wu, H., Xiao, B., & Codella, N. (2021). Cvt: Introducing convolutions to vision transformers. In Proceedings of the IEEE\/CVF international conference on computer vision (ICCV) (pp. 22\u201331).","DOI":"10.1109\/ICCV48922.2021.00009"},{"key":"1948_CR51","first-page":"1","volume":"71","author":"W Xiao","year":"2022","unstructured":"Xiao, W., Zhang, Y., Wang, H., et al. (2022). Heterogeneous knowledge distillation for simultaneous infrared-visible image fusion and super-resolution. IEEE Transactions on Instrumentation and Measurement, 71, 1\u201315.","journal-title":"IEEE Transactions on Instrumentation and Measurement"},{"issue":"1","key":"1948_CR52","doi-asserted-by":"publisher","first-page":"502","DOI":"10.1109\/TPAMI.2020.3012548","volume":"44","author":"H Xu","year":"2022","unstructured":"Xu, H., Ma, J., Jiang, J., et al. (2022). U2fusion: A unified unsupervised image fusion network. IEEE Transactions on Pattern Analysis and Machine Intelligence, 44(1), 502\u2013518.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"1948_CR53","doi-asserted-by":"crossref","unstructured":"Xu, H., Ma, J., Yuan, J., et\u00a0al. (2022b). Rfnet: Unsupervised network for mutually reinforcing multi-modal image registration and fusion. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR) (pp. 19,679\u201319,688).","DOI":"10.1109\/CVPR52688.2022.01906"},{"issue":"10","key":"1948_CR54","doi-asserted-by":"crossref","first-page":"12148","DOI":"10.1109\/TPAMI.2023.3283682","volume":"45","author":"H Xu","year":"2023","unstructured":"Xu, H., Yuan, J., & Ma, J. (2023). Murf: Mutually reinforcing multi-modal image registration and fusion. IEEE Transactions on Pattern Analysis and Machine Intelligence, 45(10), 12148\u201312166.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"issue":"4","key":"1948_CR55","doi-asserted-by":"publisher","first-page":"308","DOI":"10.1049\/el:20000267","volume":"36","author":"CS Xydeas","year":"2000","unstructured":"Xydeas, C. S., Petrovic, V., et al. (2000). Objective image fusion performance measure. Electronics Letters, 36(4), 308\u2013309.","journal-title":"Electronics Letters"},{"issue":"11","key":"1948_CR56","first-page":"1727","volume":"46","author":"Y Yao","year":"2021","unstructured":"Yao, Y., Zhang, Y., Wan, Y., et al. (2021). Heterologous images matching considering anisotropic weighted moment and absolute phase orientation. Geomatics and Information Science of Wuhan University, 46(11), 1727\u20131736.","journal-title":"Geomatics and Information Science of Wuhan University"},{"key":"1948_CR57","doi-asserted-by":"crossref","unstructured":"Yi, P., Wang, Z., & Jiang, K. (2021). Omniscient video super-resolution. In Proceedings of the IEEE\/CVF international conference on computer vision (CVPR) (pp. 4429\u20134438).","DOI":"10.1109\/ICCV48922.2021.00439"},{"key":"1948_CR58","doi-asserted-by":"publisher","first-page":"3051","DOI":"10.1007\/s11263-021-01515-2","volume":"129","author":"C Yu","year":"2021","unstructured":"Yu, C., Gao, C., Wang, J., et al. (2021). Bisenet v2: Bilateral network with guided aggregation for real-time semantic segmentation. International Journal of Computer Vision, 129, 3051\u20133068.","journal-title":"International Journal of Computer Vision"},{"key":"1948_CR59","doi-asserted-by":"crossref","unstructured":"Yuan, K., Guo, S., Liu, Z., et\u00a0al. (2021). Incorporating convolution designs into visual transformers. In Proceedings of the IEEE\/CVF international conference on computer vision (ICCV) (pp. 579\u2013588)","DOI":"10.1109\/ICCV48922.2021.00062"},{"key":"1948_CR60","doi-asserted-by":"crossref","unstructured":"Zhang, G., Zhang, P., Qi, J., et\u00a0al. (2021). Hat: Hierarchical aggregation transformers for person re-identification. In Proceedings of the 29th ACM international conference on multimedia (ACMMM). (pp. 516\u2013525).","DOI":"10.1145\/3474085.3475202"},{"key":"1948_CR61","doi-asserted-by":"crossref","unstructured":"Zhang, H., & Ma, J. (2021). Sdnet: A versatile squeeze-and-decomposition network for real-time image fusion. International Journal of Computer Vision, 129(10), 2761\u20132785.","DOI":"10.1007\/s11263-021-01501-8"},{"key":"1948_CR62","unstructured":"Zhang, H., Xu, H., Xiao, Y., et\u00a0al. (2020a). Rethinking the image fusion: A fast unified image fusion network based on proportional maintenance of gradient and intensity. In Proceedings of the AAAI conference on artificial intelligence (pp. 12,797\u201312,804)."},{"key":"1948_CR63","doi-asserted-by":"publisher","first-page":"99","DOI":"10.1016\/j.inffus.2019.07.011","volume":"54","author":"Y Zhang","year":"2020","unstructured":"Zhang, Y., Liu, Y., Sun, P., et al. (2020). IFCNN: A general image fusion framework based on convolutional neural network. Information Fusion, 54, 99\u2013118.","journal-title":"Information Fusion"},{"issue":"3","key":"1948_CR64","doi-asserted-by":"publisher","first-page":"1186","DOI":"10.1109\/TCSVT.2021.3075745","volume":"32","author":"Z Zhao","year":"2022","unstructured":"Zhao, Z., Xu, S., Zhang, J., et al. (2022). Efficient and model-based infrared and visible image fusion via algorithm unrolling. IEEE Transactions on Circuits and Systems for Video Technology, 32(3), 1186\u20131196.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"key":"1948_CR65","doi-asserted-by":"crossref","unstructured":"Zheng, S., Lu, J., & Zhao, H. (2021). Rethinking semantic segmentation from a sequence-to-sequence perspective with transformers. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR) (pp. 6881\u20136890).","DOI":"10.1109\/CVPR46437.2021.00681"},{"key":"1948_CR66","unstructured":"Zhu, X., Su, W., & Lu, L. et\u00a0al. (2021). Deformable detr: Deformable transformers for end-to-end object detection. In International conference on learning representations (ICLR)."}],"container-title":["International Journal of Computer Vision"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-023-01948-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11263-023-01948-x\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-023-01948-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,11,4]],"date-time":"2024-11-04T13:12:14Z","timestamp":1730725934000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11263-023-01948-x"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2023,11,30]]},"references-count":66,"journal-issue":{"issue":"5","published-print":{"date-parts":[[2024,5]]}},"alternative-id":["1948"],"URL":"https:\/\/doi.org\/10.1007\/s11263-023-01948-x","relation":{},"ISSN":["0920-5691","1573-1405"],"issn-type":[{"value":"0920-5691","type":"print"},{"value":"1573-1405","type":"electronic"}],"subject":[],"published":{"date-parts":[[2023,11,30]]},"assertion":[{"value":"2 August 2022","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"31 October 2023","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"30 November 2023","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}]}}