{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,10,9]],"date-time":"2025-10-09T06:29:19Z","timestamp":1759991359770,"version":"3.44.0"},"reference-count":38,"publisher":"Springer Science and Business Media LLC","issue":"10","license":[{"start":{"date-parts":[[2025,6,26]],"date-time":"2025-06-26T00:00:00Z","timestamp":1750896000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,6,26]],"date-time":"2025-06-26T00:00:00Z","timestamp":1750896000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100005047","name":"Natural Science Foundation of Liaoning Province","doi-asserted-by":"publisher","award":["No. 2022BS105","No. 2022BS105"],"award-info":[{"award-number":["No. 2022BS105","No. 2022BS105"]}],"id":[{"id":"10.13039\/501100005047","id-type":"DOI","asserted-by":"publisher"}]},{"name":"the Fundamental Research Funds for Technical Study of Ministry of Public Security of China","award":["No. 2023JSYJC23"],"award-info":[{"award-number":["No. 2023JSYJC23"]}]},{"name":"the Public Security Theory and Soft Science Foundation of Ministry of Public Security of China","award":["No. 2023LL21","No. 2023LL21"],"award-info":[{"award-number":["No. 2023LL21","No. 2023LL21"]}]},{"name":"the Fundamental Research Funds for the Central Universities of Criminal Investigation Police University of China with Grant","award":["No.C2023008"],"award-info":[{"award-number":["No.C2023008"]}]},{"name":"the Basic Research Projects of Liaoning Provincial Department of Education with Grant","award":["No. JYTMS20231411"],"award-info":[{"award-number":["No. JYTMS20231411"]}]},{"name":"the Key Research Foundation of Criminal Investigation Police University of China","award":["No. D2025009"],"award-info":[{"award-number":["No. D2025009"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["SIViP"],"published-print":{"date-parts":[[2025,10]]},"DOI":"10.1007\/s11760-025-04390-3","type":"journal-article","created":{"date-parts":[[2025,6,26]],"date-time":"2025-06-26T10:05:37Z","timestamp":1750932337000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":2,"title":["A novel multi-feature fusion network for semantic segmentation"],"prefix":"10.1007","volume":"19","author":[{"given":"Baoyu","family":"Wang","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Hegui","family":"Zhu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Pingping","family":"Cao","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Libo","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Aihong","family":"Shen","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,6,26]]},"reference":[{"issue":"12","key":"4390_CR1","doi-asserted-by":"publisher","first-page":"16601","DOI":"10.1007\/s11042-022-12577-w","volume":"81","author":"M Afif","year":"2022","unstructured":"Afif, M., Ayachi, R., Said, Y., Pissaloux, E., Atri, M.: An efficient object detection system for indoor assistance navigation using deep learning techniques. Multimedia Tools and Applications 81(12), 16601\u201316618 (2022)","journal-title":"Multimedia Tools and Applications"},{"issue":"7","key":"4390_CR2","doi-asserted-by":"publisher","first-page":"10051","DOI":"10.1007\/s11042-022-12042-8","volume":"81","author":"R Das","year":"2022","unstructured":"Das, R., Singh, T.D.: Assamese news image caption generation using attention mechanism. Multimedia Tools and Applications 81(7), 10051\u201310069 (2022)","journal-title":"Multimedia Tools and Applications"},{"key":"4390_CR3","first-page":"1","volume":"20","author":"J Jin","year":"2023","unstructured":"Jin, J., Zhou, W., Yang, R., Ye, L., Yu, L.: Edge detection guide network for semantic segmentation of remote-sensing images. IEEE Geosci. Remote Sens. Lett. 20, 1\u20135 (2023)","journal-title":"IEEE Geosci. Remote Sens. Lett."},{"key":"4390_CR4","doi-asserted-by":"publisher","first-page":"33","DOI":"10.1016\/j.neucom.2023.02.025","volume":"532","author":"TH Tsai","year":"2023","unstructured":"Tsai, T.H., Tseng, Y.W.: Bisenet v3: Bilateral segmentation network with coordinate attention for real-time semantic segmentation. Neurocomputing 532, 33\u201342 (2023)","journal-title":"Neurocomputing"},{"key":"4390_CR5","doi-asserted-by":"crossref","unstructured":"Ma, X., Zhang, X., Pun, M.O., Liu, M.: A multilevel multimodal fusion transformer for remote sensing semantic segmentation. IEEE Transactions on Geoscience and Remote Sensing (2024)","DOI":"10.1109\/TGRS.2024.3373033"},{"issue":"2","key":"4390_CR6","doi-asserted-by":"publisher","first-page":"297","DOI":"10.1049\/cit2.12118","volume":"8","author":"G Zhao","year":"2023","unstructured":"Zhao, G., Zhang, Y., Ge, M., Yu, M.: Bilateral u-net semantic segmentation with spatial attention mechanism. CAAI Transactions on Intelligence Technology 8(2), 297\u2013307 (2023)","journal-title":"CAAI Transactions on Intelligence Technology"},{"issue":"1","key":"4390_CR7","doi-asserted-by":"publisher","first-page":"016,518","DOI":"10.1117\/1.JRS.17.016518","volume":"17","author":"H Jin","year":"2023","unstructured":"Jin, H., Bao, Z., Chang, X., Zhang, T., Chen, C.: Semantic segmentation of remote sensing images based on dilated convolution and spatial-channel attention mechanism. J. Appl. Remote Sens. 17(1), 016,518-016,518 (2023)","journal-title":"J. Appl. Remote Sens."},{"key":"4390_CR8","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2024.127625","volume":"587","author":"KI Rashid","year":"2024","unstructured":"Rashid, K.I., Yang, C., Huang, C.: Fast-dsagcn: Enhancing semantic segmentation with multifaceted attention mechanisms. Neurocomputing 587, 127625 (2024)","journal-title":"Neurocomputing"},{"issue":"7","key":"4390_CR9","doi-asserted-by":"publisher","DOI":"10.1088\/1361-6501\/ad1ddb","volume":"35","author":"Y Tan","year":"2024","unstructured":"Tan, Y., Li, X., Lai, J., Ai, J.: Real-time tunnel lining leakage image semantic segmentation via multiple attention mechanisms. Meas. Sci. Technol. 35(7), 075204 (2024)","journal-title":"Meas. Sci. Technol."},{"issue":"2","key":"4390_CR10","doi-asserted-by":"publisher","first-page":"361","DOI":"10.3390\/rs15020361","volume":"15","author":"M Yuan","year":"2023","unstructured":"Yuan, M., Ren, D., Feng, Q., Wang, Z., Dong, Y., Lu, F., Wu, X.: Mcafnet: a multiscale channel attention fusion network for semantic segmentation of remote sensing images. Remote Sensing 15(2), 361 (2023)","journal-title":"Remote Sensing"},{"key":"4390_CR11","first-page":"1","volume":"62","author":"J Liu","year":"2024","unstructured":"Liu, J., Hua, W., Zhang, W., Liu, F., Xiao, L.: Stair fusion network with context-refined attention for remote sensing image semantic segmentation. IEEE Trans. Geosci. Remote Sens. 62, 1\u201317 (2024)","journal-title":"IEEE Trans. Geosci. Remote Sens."},{"issue":"12","key":"4390_CR12","doi-asserted-by":"publisher","first-page":"3587","DOI":"10.1049\/ipr2.13196","volume":"18","author":"B Wang","year":"2024","unstructured":"Wang, B., Shen, A., Dong, X., Cao, P.: Cf-net: Cross fusion network for semantic segmentation. IET Image Proc. 18(12), 3587\u20133599 (2024)","journal-title":"IET Image Proc."},{"key":"4390_CR13","first-page":"1","volume":"21","author":"H Wu","year":"2023","unstructured":"Wu, H., Huang, P., Zhang, M., Tang, W.: Ctfnet: Cnn-transformer fusion network for remote-sensing image semantic segmentation. IEEE Geosci. Remote Sens. Lett. 21, 1\u20135 (2023)","journal-title":"IEEE Geosci. Remote Sens. Lett."},{"key":"4390_CR14","first-page":"1","volume":"61","author":"H Wu","year":"2023","unstructured":"Wu, H., Huang, P., Zhang, M., Tang, W., Yu, X.: Cmtfnet: Cnn and multiscale transformer fusion network for remote-sensing image semantic segmentation. IEEE Trans. Geosci. Remote Sens. 61, 1\u201312 (2023)","journal-title":"IEEE Trans. Geosci. Remote Sens."},{"key":"4390_CR15","doi-asserted-by":"crossref","unstructured":"Wu, H., Zhang, M., Huang, P., Tang, W.: Cmlformer: Cnn and multi-scale local-context transformer network for remote sensing images semantic segmentation. IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing (2024)","DOI":"10.1109\/JSTARS.2024.3375313"},{"issue":"6","key":"4390_CR16","doi-asserted-by":"publisher","first-page":"125","DOI":"10.3390\/jimaging10060125","volume":"10","author":"N Detsikas","year":"2024","unstructured":"Detsikas, N., Mitianoudis, N., Pratikakis, I.: Mrestnet: A multi-resolution transformer framework with cnn extensions for semantic segmentation. Journal of Imaging 10(6), 125 (2024)","journal-title":"Journal of Imaging"},{"key":"4390_CR17","doi-asserted-by":"crossref","unstructured":"Shao, Y., Sun, L., Jiao, L., Liu, X., Liu, F., Li, L., Yang, S.: Cot: Contourlet transformer for hierarchical semantic segmentation. IEEE Transactions on Neural Networks and Learning Systems (2024)","DOI":"10.1109\/TNNLS.2024.3367901"},{"key":"4390_CR18","first-page":"1","volume":"61","author":"Z Xu","year":"2023","unstructured":"Xu, Z., Geng, J., Jiang, W.: Mmt: Mixed-mask transformer for remote sensing image semantic segmentation. IEEE Trans. Geosci. Remote Sens. 61, 1\u201315 (2023)","journal-title":"IEEE Trans. Geosci. Remote Sens."},{"key":"4390_CR19","doi-asserted-by":"publisher","first-page":"490","DOI":"10.1016\/j.neucom.2021.06.067","volume":"458","author":"H Zhu","year":"2021","unstructured":"Zhu, H., Zeng, H., Liu, J., Zhang, X.: Logish: A new nonlinear nonmonotonic activation function for convolutional neural network. Neurocomputing 458, 490\u2013499 (2021)","journal-title":"Neurocomputing"},{"key":"4390_CR20","doi-asserted-by":"crossref","unstructured":"Xie, B., Cao, J., Xie, J., Khan, F.S., Pang, Y.: Sed: A simple encoder-decoder for open-vocabulary semantic segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp. 3426\u20133436 (2024)","DOI":"10.1109\/CVPR52733.2024.00329"},{"key":"4390_CR21","doi-asserted-by":"crossref","unstructured":"Wang, Y., Zhou, Q., Liu, J., Xiong, J., Gao, G., Wu, X., Latecki, L.J.: Lednet: A lightweight encoder-decoder network for real-time semantic segmentation. In: 2019 IEEE international conference on image processing (ICIP), pp. 1860\u20131864. IEEE (2019)","DOI":"10.1109\/ICIP.2019.8803154"},{"key":"4390_CR22","doi-asserted-by":"crossref","unstructured":"Zhang, R., Leng, L., Che, K., Zhang, H., Cheng, J., Guo, Q., Liao, J., Cheng, R.: Accurate and efficient event-based semantic segmentation using adaptive spiking encoder\u2013decoder network. IEEE Transactions on Neural Networks and Learning Systems (2024)","DOI":"10.1109\/TNNLS.2024.3437415"},{"key":"4390_CR23","doi-asserted-by":"crossref","unstructured":"Hoang, T.M., Zhou, J., Fan, Y.: Image compression with encoder-decoder matched semantic segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition workshops, pp. 160\u2013161 (2020)","DOI":"10.1109\/CVPRW50498.2020.00088"},{"key":"4390_CR24","doi-asserted-by":"crossref","unstructured":"Li, X., Xu, F., Li, L., Xu, N., Liu, F., Yuan, C., Chen, Z., Lyu, X.: Aaformer: Attention-attended transformer for semantic segmentation of remote sensing images. IEEE Geoscience and Remote Sensing Letters (2024)","DOI":"10.1109\/LGRS.2024.3397851"},{"key":"4390_CR25","first-page":"1","volume":"61","author":"X Li","year":"2023","unstructured":"Li, X., Xu, F., Liu, F., Lyu, X., Tong, Y., Xu, Z., Zhou, J.: A synergistical attention model for semantic segmentation of remote sensing images. IEEE Trans. Geosci. Remote Sens. 61, 1\u201316 (2023)","journal-title":"IEEE Trans. Geosci. Remote Sens."},{"key":"4390_CR26","doi-asserted-by":"crossref","unstructured":"Shu, J., Zhang, J., et\u00a0al.: Cfsa-net: Efficient large-scale point cloud semantic segmentation based on cross-fusion self-attention. Computers, Materials & Continua 77(3), (2023)","DOI":"10.32604\/cmc.2023.045818"},{"key":"4390_CR27","doi-asserted-by":"publisher","DOI":"10.1016\/j.engappai.2022.105554","volume":"117","author":"Y Zhou","year":"2023","unstructured":"Zhou, Y., Ji, A., Zhang, L., Xue, X.: Sampling-attention deep learning network with transfer learning for large-scale urban point cloud semantic segmentation. Eng. Appl. Artif. Intell. 117, 105554 (2023)","journal-title":"Eng. Appl. Artif. Intell."},{"key":"4390_CR28","first-page":"1","volume":"61","author":"T Zeng","year":"2023","unstructured":"Zeng, T., Luo, F., Guo, T., Gong, X., Xue, J., Li, H.: Recurrent residual dual attention network for airborne laser scanning point cloud semantic segmentation. IEEE Trans. Geosci. Remote Sens. 61, 1\u201314 (2023)","journal-title":"IEEE Trans. Geosci. Remote Sens."},{"key":"4390_CR29","first-page":"12077","volume":"34","author":"E Xie","year":"2021","unstructured":"Xie, E., Wang, W., Yu, Z., Anandkumar, A., Alvarez, J.M., Luo, P.: Segformer: Simple and efficient design for semantic segmentation with transformers. Adv. Neural. Inf. Process. Syst. 34, 12077\u201312090 (2021)","journal-title":"Adv. Neural. Inf. Process. Syst."},{"key":"4390_CR30","unstructured":"Yan, H., Zhang, C., Wu, M.: Lawin transformer: Improving semantic segmentation transformer with multi-scale representations via large window attention. arXiv preprint arXiv:2201.01615 (2022)"},{"key":"4390_CR31","first-page":"1","volume":"61","author":"T Xiao","year":"2023","unstructured":"Xiao, T., Liu, Y., Huang, Y., Li, M., Yang, G.: Enhancing multiscale representations with transformer for remote sensing image semantic segmentation. IEEE Trans. Geosci. Remote Sens. 61, 1\u201316 (2023)","journal-title":"IEEE Trans. Geosci. Remote Sens."},{"key":"4390_CR32","doi-asserted-by":"publisher","first-page":"3023","DOI":"10.1109\/JSTARS.2024.3349657","volume":"17","author":"M Yao","year":"2024","unstructured":"Yao, M., Zhang, Y., Liu, G., Pang, D.: Ssnet: A novel transformer and cnn hybrid network for remote sensing semantic segmentation. IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing 17, 3023\u20133037 (2024)","journal-title":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing"},{"issue":"9","key":"4390_CR33","doi-asserted-by":"publisher","first-page":"6186","DOI":"10.1109\/TCSVT.2022.3162599","volume":"32","author":"J Nie","year":"2022","unstructured":"Nie, J., Wu, H., He, Z., Gao, M., Dong, Z.: Spreading fine-grained prior knowledge for accurate tracking. IEEE Trans. Circuits Syst. Video Technol. 32(9), 6186\u20136199 (2022)","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"4390_CR34","doi-asserted-by":"publisher","first-page":"6194","DOI":"10.1109\/TMM.2022.3206668","volume":"25","author":"J Nie","year":"2022","unstructured":"Nie, J., He, Z., Yang, Y., Gao, M., Dong, Z.: Learning localization-aware target confidence for siamese visual tracking. IEEE Trans. Multimedia 25, 6194\u20136206 (2022)","journal-title":"IEEE Trans. Multimedia"},{"key":"4390_CR35","doi-asserted-by":"crossref","unstructured":"Nie, J., He, Z., Yang, Y., Gao, M., Zhang, J.: Glt-t: Global-local transformer voting for 3d single object tracking in point clouds. In: proceedings of the AAAI conference on artificial intelligence, 37, 1957\u20131965 (2023)","DOI":"10.1609\/aaai.v37i2.25287"},{"issue":"2","key":"4390_CR36","doi-asserted-by":"publisher","first-page":"88","DOI":"10.1016\/j.patrec.2008.04.005","volume":"30","author":"GJ Brostow","year":"2009","unstructured":"Brostow, G.J., Fauqueur, J., Cipolla, R.: Semantic object classes in video: A high-definition ground truth database. Pattern Recogn. Lett. 30(2), 88\u201397 (2009)","journal-title":"Pattern Recogn. Lett."},{"key":"4390_CR37","doi-asserted-by":"crossref","unstructured":"Cordts, M., Omran, M., Ramos, S., Rehfeld, T., Enzweiler, M., Benenson, R., Franke, U., Roth, S., Schiele, B.: The cityscapes dataset for semantic urban scene understanding. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp. 3213\u20133223 (2016)","DOI":"10.1109\/CVPR.2016.350"},{"issue":"1","key":"4390_CR38","doi-asserted-by":"publisher","first-page":"98","DOI":"10.1007\/s11263-014-0733-5","volume":"111","author":"M Everingham","year":"2015","unstructured":"Everingham, M., Eslami, S., Van Gool, L., Williams, C.K., Winn, J., Zisserman, A.: The pascal visual object classes challenge: A retrospective. Int. J. Comput. Vision 111(1), 98\u2013136 (2015)","journal-title":"Int. J. Comput. Vision"}],"container-title":["Signal, Image and Video Processing"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11760-025-04390-3.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11760-025-04390-3\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11760-025-04390-3.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,9,6]],"date-time":"2025-09-06T22:44:42Z","timestamp":1757198682000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11760-025-04390-3"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,6,26]]},"references-count":38,"journal-issue":{"issue":"10","published-print":{"date-parts":[[2025,10]]}},"alternative-id":["4390"],"URL":"https:\/\/doi.org\/10.1007\/s11760-025-04390-3","relation":{},"ISSN":["1863-1703","1863-1711"],"issn-type":[{"type":"print","value":"1863-1703"},{"type":"electronic","value":"1863-1711"}],"subject":[],"published":{"date-parts":[[2025,6,26]]},"assertion":[{"value":"25 February 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"4 June 2025","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"12 June 2025","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"26 June 2025","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflicts of Interest"}},{"value":"The data used in this paper are from publicly available datasets and do not violate any ethical guidelines.","order":3,"name":"Ethics","group":{"name":"EthicsHeading","label":"Ethical and informed consent for data used"}}],"article-number":"785"}}