{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,2,11]],"date-time":"2026-02-11T08:55:12Z","timestamp":1770800112651,"version":"3.50.0"},"reference-count":85,"publisher":"Springer Science and Business Media LLC","issue":"1","license":[{"start":{"date-parts":[[2025,12,15]],"date-time":"2025-12-15T00:00:00Z","timestamp":1765756800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,12,15]],"date-time":"2025-12-15T00:00:00Z","timestamp":1765756800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"Project Supported by Natural Science Foundation of Henan","award":["232300421023"],"award-info":[{"award-number":["232300421023"]}]},{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"crossref","award":["62176113"],"award-info":[{"award-number":["62176113"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"crossref"}]},{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"crossref","award":["62176113"],"award-info":[{"award-number":["62176113"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"crossref"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimedia Systems"],"published-print":{"date-parts":[[2026,2]]},"DOI":"10.1007\/s00530-025-02110-y","type":"journal-article","created":{"date-parts":[[2025,12,15]],"date-time":"2025-12-15T04:29:46Z","timestamp":1765772986000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Attention-guided trilateral network for real-time semantic segmentation"],"prefix":"10.1007","volume":"32","author":[{"given":"Siming","family":"Jia","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yongsheng","family":"Dong","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Lintao","family":"Zheng","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Chongchong","family":"Mao","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Lin","family":"Wang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Guoyong","family":"Wang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,12,15]]},"reference":[{"issue":"6","key":"2110_CR1","doi-asserted-by":"publisher","first-page":"17233","DOI":"10.1007\/s11042-023-16262-4","volume":"83","author":"Z Jiang","year":"2024","unstructured":"Jiang, Z., Zaheer, W., Wali, A., Gilani, S.: Visual sentiment analysis using data-augmented deep transfer learning techniques. Multimedia Tool Appl. 83(6), 17233\u201317249 (2024)","journal-title":"Multimedia Tool Appl."},{"issue":"6","key":"2110_CR2","doi-asserted-by":"publisher","first-page":"3461","DOI":"10.1109\/TSMC.2022.3225381","volume":"53","author":"Z Zhuang","year":"2023","unstructured":"Zhuang, Z., Tao, H., Chen, Y., Stojanovic, V., Paszke, W.: An optimal iterative learning control approach for linear systems with nonuniform trial lengths under input constraints. IEEE Trans. Syst. Man. Cybern. Syst. 53(6), 3461\u20133473 (2023)","journal-title":"IEEE Trans. Syst. Man. Cybern. Syst."},{"issue":"9","key":"2110_CR3","doi-asserted-by":"publisher","first-page":"4496","DOI":"10.1109\/TCSVT.2023.3278131","volume":"33","author":"F Zhang","year":"2023","unstructured":"Zhang, F., Chen, G., Wang, H., Li, J., Zhang, C.: Multi-scale video super-resolution transformer with polynomial approximation. IEEE Trans. Circuits Syst. Video Technol. 33(9), 4496\u20134506 (2023)","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"2110_CR4","doi-asserted-by":"crossref","unstructured":"Mettes, P., Ghadimi\u00a0Atigh, M., Keller-Ressel, M., Gu, J., Yeung, S.: Hyperbolic deep learning in computer vision: a survey. Int. J. Comput. Vis. 1\u201325 (2024)","DOI":"10.1007\/s11263-024-02043-5"},{"issue":"3","key":"2110_CR5","doi-asserted-by":"publisher","first-page":"593","DOI":"10.1007\/s41095-023-0369-x","volume":"10","author":"F Zhang","year":"2024","unstructured":"Zhang, F., Chen, G., Wang, H., Zhang, C.: CF-DAN: facial-expression recognition based on cross-fusion dual-attention network. Comput. Visual Media 10(3), 593\u2013608 (2024)","journal-title":"Comput. Visual Media"},{"issue":"5","key":"2110_CR6","doi-asserted-by":"publisher","first-page":"272","DOI":"10.1007\/s00530-024-01459-w","volume":"30","author":"B Ge","year":"2024","unstructured":"Ge, B., Zhu, X., Tang, Z., Xia, C., Lu, Y., Chen, Z.: Triple fusion and feature pyramid decoder for rgb-d semantic segmentation. Multimedia Syst. 30(5), 272 (2024)","journal-title":"Multimedia Syst."},{"key":"2110_CR7","doi-asserted-by":"crossref","unstructured":"Ji, W., Li, J., Bian, C., Zhou, Z., Zhao, J., Yuille, A.L., Cheng, L.: Multispectral video semantic segmentation: A benchmark dataset and baseline. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 1094\u20131104 (2023)","DOI":"10.1109\/CVPR52729.2023.00112"},{"issue":"1","key":"2110_CR8","first-page":"1","volume":"31","author":"X Yu","year":"2025","unstructured":"Yu, X., Liu, H., Zhang, D., Wu, J., Zheng, J.: Hierarchical region-level decoupling knowledge distillation for semantic segmentation. Multimedia Syst. 31(1), 1\u201318 (2025)","journal-title":"Multimedia Syst."},{"issue":"1","key":"2110_CR9","doi-asserted-by":"publisher","first-page":"76","DOI":"10.1007\/s00530-024-01654-9","volume":"31","author":"T Zhou","year":"2025","unstructured":"Zhou, T., He, H., Wang, Y., Liao, Y.: Improved gated recurrent units together with fusion for semantic segmentation of remote sensing images based on parallel hybrid network. Multimedia Syst. 31(1), 76 (2025)","journal-title":"Multimedia Syst."},{"key":"2110_CR10","doi-asserted-by":"crossref","unstructured":"Long, J., Shelhamer, E., Darrell, T.: Fully convolutional networks for semantic segmentation. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 3431\u20133440 (2015)","DOI":"10.1109\/CVPR.2015.7298965"},{"issue":"6","key":"2110_CR11","doi-asserted-by":"publisher","first-page":"364","DOI":"10.1007\/s00530-024-01497-4","volume":"30","author":"P Liu","year":"2024","unstructured":"Liu, P., Tian, S., Gao, Y., Xie, Y., Hao, S.: Efficientfusion: simple and efficient learning with pixel-level fusion for semantic segmentation. Multimedia Syst. 30(6), 364 (2024)","journal-title":"Multimedia Syst."},{"key":"2110_CR12","doi-asserted-by":"crossref","unstructured":"Zhao, H., Shi, J., Qi, X., Wang, X., Jia, J.: Pyramid scene parsing network. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 2881\u20132890 (2017)","DOI":"10.1109\/CVPR.2017.660"},{"key":"2110_CR13","doi-asserted-by":"crossref","unstructured":"He, H., Cai, J., Pan, Z., Liu, J., Zhang, J., Tao, D., Zhuang, B.: Dynamic focus-aware positional queries for semantic segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 11299\u201311308 (2023)","DOI":"10.1109\/CVPR52729.2023.01087"},{"key":"2110_CR14","doi-asserted-by":"crossref","unstructured":"Chen, J., Lu, J., Zhu, X., Zhang, L.: Generative semantic segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 7111\u20137120 (2023)","DOI":"10.1109\/CVPR52729.2023.00687"},{"key":"2110_CR15","unstructured":"Paszke, A., Chaurasia, A., Kim, S., Culurciello, E.: ENet: a deep neural network architecture for real-time semantic segmentation. arXiv preprint arXiv:1606.02147 (2016)"},{"key":"2110_CR16","unstructured":"Howard, A.G., Zhu, M., Chen, B., Kalenichenko, D., Wang, W., Weyand, T., Andreetto, M., Adam, H.: MobileNets: efficient convolutional neural networks for mobile vision applications. arXiv preprint arXiv:1704.04861 (2017)"},{"key":"2110_CR17","doi-asserted-by":"crossref","unstructured":"Sandler, M., Howard, A., Zhu, M., Zhmoginov, A., Chen, L.-C.: MobileNetV2: inverted residuals and linear bottlenecks. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 4510\u20134520 (2018)","DOI":"10.1109\/CVPR.2018.00474"},{"key":"2110_CR18","doi-asserted-by":"crossref","unstructured":"Howard, A., Sandler, M., Chu, G., Chen, L.-C., Chen, B., Tan, M., Wang, W., Zhu, Y., Pang, R., Vasudevan, V., et al.: Searching for MobileNetV3. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 1314\u20131324 (2019)","DOI":"10.1109\/ICCV.2019.00140"},{"key":"2110_CR19","doi-asserted-by":"crossref","unstructured":"Yu, C., Wang, J., Peng, C., Gao, C., Yu, G., Sang, N.: BiSeNet: bilateral segmentation network for real-time semantic segmentation. In: Proceedings of the European Conference on Computer Vision, pp. 325\u2013341 (2018)","DOI":"10.1007\/978-3-030-01261-8_20"},{"key":"2110_CR20","doi-asserted-by":"crossref","unstructured":"Fan, M., Lai, S., Huang, J., Wei, X., Chai, Z., Luo, J., Wei, X.: Rethinking bisenet for real-time semantic segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 9716\u20139725 (2021)","DOI":"10.1109\/CVPR46437.2021.00959"},{"key":"2110_CR21","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2024.127653","volume":"586","author":"H Li","year":"2024","unstructured":"Li, H., Zhang, Y., Xiong, Z., Sun, X.: Deep multi-threshold spiking-UNet for image processing. Neurocomputing 586, 127653 (2024)","journal-title":"Neurocomputing"},{"key":"2110_CR22","doi-asserted-by":"crossref","unstructured":"Min, J., Lee, J., Ponce, J., Cho, M.: Hyperpixel flow: Semantic correspondence with multi-layer neural features. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 3395\u20133404 (2019)","DOI":"10.1109\/ICCV.2019.00349"},{"key":"2110_CR23","doi-asserted-by":"crossref","unstructured":"Zhou, W., Fan, X., Yan, W., Shan, S., Jiang, Q., Hwang, J.-N.: Graph attention guidance network with knowledge distillation for semantic segmentation of remote sensing images. IEEE Trans. Geosci. Remote Sens. (2023)","DOI":"10.1109\/TGRS.2023.3311480"},{"key":"2110_CR24","first-page":"379","volume":"5","author":"MA Hameed","year":"2024","unstructured":"Hameed, M.A., Hassaballah, M., Abdelazim, R., Sahu, A.K.: A novel medical steganography technique based on adversarial neural cryptography and digital signature using least significant bit replacement. Int. J. Cogn. Comput. Eng. 5, 379\u2013397 (2024)","journal-title":"Int. J. Cogn. Comput. Eng."},{"issue":"3","key":"2110_CR25","doi-asserted-by":"publisher","first-page":"129","DOI":"10.1007\/s00530-024-01332-w","volume":"30","author":"MA Hameed","year":"2024","unstructured":"Hameed, M.A., Hassaballah, M., Qiao, T.: IS-DGM: an improved steganography method based on a deep generative model and hyper logistic map encryption via social media networks. Multimedia Syst. 30(3), 129 (2024)","journal-title":"Multimedia Syst."},{"issue":"1","key":"2110_CR26","first-page":"2247675","volume":"2022","author":"M Abdel Hameed","year":"2022","unstructured":"Abdel Hameed, M., Hassaballah, M., Hosney, M.E., Alqahtani, A.: [Retracted] an ai-enabled internet of things based autism care system for improving cognitive ability of children with autism spectrum disorders. Comput. Intell. Neurosci. 2022(1), 2247675 (2022)","journal-title":"Comput. Intell. Neurosci."},{"issue":"5","key":"2110_CR27","doi-asserted-by":"publisher","first-page":"4639","DOI":"10.1007\/s12652-022-04366-y","volume":"14","author":"MA Hameed","year":"2023","unstructured":"Hameed, M.A., Abdel-Aleem, O.A., Hassaballah, M.: A secure data hiding approach based on least-significant-bit and nature-inspired optimization techniques. J. Ambient. Intell. Humaniz. Comput. 14(5), 4639\u20134657 (2023)","journal-title":"J. Ambient. Intell. Humaniz. Comput."},{"key":"2110_CR28","doi-asserted-by":"crossref","unstructured":"Hassaballah, M., Hameed, M.A., Alkinani, M.H.: Introduction to digital image steganography. In: Proceedings of the Digital Media Steganography, pp. 1\u201315. Elsevier, (2020)","DOI":"10.1016\/B978-0-12-819438-6.00009-8"},{"key":"2110_CR29","doi-asserted-by":"crossref","unstructured":"Hassaballah, M., Hameed, M.A., Aly, S., AbdelRady, A.: A color image steganography method based on adpvd and hog techniques. In: Proceedings of the Digital Media Steganography, pp. 17\u201340. Elsevier, (2020)","DOI":"10.1016\/B978-0-12-819438-6.00010-4"},{"issue":"11","key":"2110_CR30","doi-asserted-by":"publisher","first-page":"7743","DOI":"10.1109\/TII.2021.3053595","volume":"17","author":"M Hassaballah","year":"2021","unstructured":"Hassaballah, M., Hameed, M.A., Awad, A.I., Muhammad, K.: A novel image steganography method for industrial internet of things security. IEEE Trans. Industr. Inf. 17(11), 7743\u20137751 (2021)","journal-title":"IEEE Trans. Industr. Inf."},{"key":"2110_CR31","doi-asserted-by":"publisher","first-page":"185189","DOI":"10.1109\/ACCESS.2019.2960254","volume":"7","author":"MA Hameed","year":"2019","unstructured":"Hameed, M.A., Hassaballah, M., Aly, S., Awad, A.I.: An adaptive image steganography method based on histogram of oriented gradient and PVD-LSB techniques. IEEE Access 7, 185189\u2013185204 (2019)","journal-title":"IEEE Access"},{"issue":"12","key":"2110_CR32","doi-asserted-by":"publisher","first-page":"14705","DOI":"10.1007\/s11042-017-5056-4","volume":"77","author":"M Abdel Hameed","year":"2018","unstructured":"Abdel Hameed, M., Aly, S., Hassaballah, M.: An efficient data hiding method based on adaptive directional pixel value differencing (ADPVD). Multimedia Tool Appl. 77(12), 14705\u201314723 (2018)","journal-title":"Multimedia Tool Appl."},{"key":"2110_CR33","doi-asserted-by":"crossref","unstructured":"Bekhet, S., Hassaballah, M., Kenk, M.A., Hameed, M.A.: An artificial intelligence based technique for covid-19 diagnosis from chest x-ray. In: Proceedings of the Novel Intelligent and Leading Emerging Sciences, IEEE. pp. 191\u2013195 (2020)","DOI":"10.1109\/NILES50944.2020.9257930"},{"key":"2110_CR34","doi-asserted-by":"crossref","unstructured":"Kenk, M.A., Hassaballah, M., Hameed, M.A., Bekhet, S.: Visibility enhancer: adaptable for distorted traffic scenes by dusty weather. In: Proceedings of the Novel Intelligent and Leading Emerging Sciences, IEEE. pp. 213\u2013218 (2020)","DOI":"10.1109\/NILES50944.2020.9257952"},{"key":"2110_CR35","unstructured":"Hassaballah, M., Aly, S., Abdel\u00a0Rady, A.S., et al.: A high payload steganography method based on pixel value differencing (2018)"},{"key":"2110_CR36","doi-asserted-by":"crossref","unstructured":"Hafiz, A.M., Hassaballah, M., Alqahtani, A., Alsubai, S., Hameed, M.A.: Reinforcement learning with an ensemble of binary action deep q-networks. Comput. syst. sci. eng. 46(3) (2023)","DOI":"10.32604\/csse.2023.031720"},{"issue":"2","key":"2110_CR37","first-page":"1","volume":"4","author":"M Abdel Hameed","year":"2024","unstructured":"Abdel Hameed, M.: Scrambled encryption approach for color images based on 9-d chaotic systems with 3-D substitution bit levels. Aswan Univ. J. Sci. Technol. 4(2), 1\u201318 (2024)","journal-title":"Aswan Univ. J. Sci. Technol."},{"key":"2110_CR38","doi-asserted-by":"crossref","unstructured":"Hameed, M.A., Hassaballah, M., Bekhet, S., Kenk, M.A., et al.: A high quality secure medical image steganography method. In: Proceedings of the International Conference on Computing and Information Technology, IEEE. pp. 465\u2013470 (2023)","DOI":"10.1109\/ICCIT58132.2023.10273950"},{"key":"2110_CR39","doi-asserted-by":"crossref","unstructured":"Hassaballah, M., Hameed, M.A.: Novel metaheuristic algorithms and their applications to efficient detection of diabetic retinopathy. J. Artif. Intell. Soft Comput. Res. 15 (2025)","DOI":"10.2478\/jaiscr-2025-0009"},{"key":"2110_CR40","unstructured":"Chen, L.-C., Papandreou, G., Kokkinos, I., Murphy, K., Yuille, A.L.: Semantic image segmentation with deep convolutional nets and fully connected crfs. arXiv preprint arXiv:1412.7062 (2014)"},{"issue":"4","key":"2110_CR41","doi-asserted-by":"publisher","first-page":"834","DOI":"10.1109\/TPAMI.2017.2699184","volume":"40","author":"L-C Chen","year":"2017","unstructured":"Chen, L.-C., Papandreou, G., Kokkinos, I., Murphy, K., Yuille, A.L.: DeepLab: semantic image segmentation with deep convolutional nets, atrous convolution, and fully connected crfs. IEEE Trans. Pattern Anal. Mach. Intell. 40(4), 834\u2013848 (2017)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"2110_CR42","unstructured":"Chen, L.-C., Papandreou, G., Schroff, F., Adam, H.: Rethinking atrous convolution for semantic image segmentation. arXiv preprint arXiv:1706.05587 (2017)"},{"key":"2110_CR43","doi-asserted-by":"crossref","unstructured":"Chen, L.-C., Zhu, Y., Papandreou, G., Schroff, F., Adam, H.: Encoder-decoder with atrous separable convolution for semantic image segmentation. In: Proceedings of the European Conference on Computer Vision, pp. 801\u2013818 (2018)","DOI":"10.1007\/978-3-030-01234-2_49"},{"key":"2110_CR44","doi-asserted-by":"crossref","unstructured":"Dai, J., Qi, H., Xiong, Y., Li, Y., Zhang, G., Hu, H., Wei, Y.: Deformable convolutional networks. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 764\u2013773 (2017)","DOI":"10.1109\/ICCV.2017.89"},{"key":"2110_CR45","doi-asserted-by":"crossref","unstructured":"Zhu, X., Hu, H., Lin, S., Dai, J.: Deformable convnets v2: More deformable, better results. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 9308\u20139316 (2019)","DOI":"10.1109\/CVPR.2019.00953"},{"key":"2110_CR46","doi-asserted-by":"crossref","unstructured":"Wang, W., Dai, J., Chen, Z., Huang, Z., Li, Z., Zhu, X., Hu, X., Lu, T., Lu, L., Li, H., et al.: Internimage: Exploring large-scale vision foundation models with deformable convolutions. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 14408\u201314419 (2023)","DOI":"10.1109\/CVPR52729.2023.01385"},{"key":"2110_CR47","doi-asserted-by":"crossref","unstructured":"Liu, Z., Lin, Y., Cao, Y., Hu, H., Wei, Y., Zhang, Z., Lin, S., Guo, B.: Swin transformer: Hierarchical vision transformer using shifted windows. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 10012\u201310022 (2021)","DOI":"10.1109\/ICCV48922.2021.00986"},{"key":"2110_CR48","first-page":"12077","volume":"34","author":"E Xie","year":"2021","unstructured":"Xie, E., Wang, W., Yu, Z., Anandkumar, A., Alvarez, J.M., Luo, P.: SegFormer: simple and efficient design for semantic segmentation with transformers. Adv. Neural. Inf. Process. Syst. 34, 12077\u201312090 (2021)","journal-title":"Adv. Neural. Inf. Process. Syst."},{"key":"2110_CR49","doi-asserted-by":"crossref","unstructured":"Zhang, W., Huang, Z., Luo, G., Chen, T., Wang, X., Liu, W., Yu, G., Shen, C.: TopFormer: token pyramid transformer for mobile semantic segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 12083\u201312093 (2022)","DOI":"10.1109\/CVPR52688.2022.01177"},{"key":"2110_CR50","doi-asserted-by":"crossref","unstructured":"Dong, B., Wang, P., Wang, F.: Head-free lightweight semantic segmentation with linear transformer. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 37, pp. 516\u2013524 (2023)","DOI":"10.1609\/aaai.v37i1.25126"},{"key":"2110_CR51","doi-asserted-by":"crossref","unstructured":"Shang, C., Li, H., Meng, F., Wu, Q., Qiu, H., Wang, L.: Incrementer: Transformer for class-incremental semantic segmentation with knowledge distillation focusing on old class. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 7214\u20137224 (2023)","DOI":"10.1109\/CVPR52729.2023.00697"},{"key":"2110_CR52","doi-asserted-by":"publisher","DOI":"10.1016\/j.engappai.2023.106634","volume":"124","author":"P Song","year":"2023","unstructured":"Song, P., Yang, Z., Li, J., Fan, H.: DPCTN: dual path context-aware transformer network for medical image segmentation. Eng. Appl. Artif. Intell. 124, 106634 (2023)","journal-title":"Eng. Appl. Artif. Intell."},{"key":"2110_CR53","doi-asserted-by":"crossref","unstructured":"Zhao, H., Qi, X., Shen, X., Shi, J., Jia, J.: ICNet for real-time semantic segmentation on high-resolution images. In: Proceedings of the European Conference on Computer Vision, pp. 405\u2013420 (2018)","DOI":"10.1007\/978-3-030-01219-9_25"},{"key":"2110_CR54","doi-asserted-by":"publisher","first-page":"3051","DOI":"10.1007\/s11263-021-01515-2","volume":"129","author":"C Yu","year":"2021","unstructured":"Yu, C., Gao, C., Wang, J., Yu, G., Shen, C., Sang, N.: BiSeNet V2: bilateral network with guided aggregation for real-time semantic segmentation. Int. J. Comput. Vision 129, 3051\u20133068 (2021)","journal-title":"Int. J. Comput. Vision"},{"key":"2110_CR55","doi-asserted-by":"publisher","first-page":"33","DOI":"10.1016\/j.neucom.2023.02.025","volume":"532","author":"T-H Tsai","year":"2023","unstructured":"Tsai, T.-H., Tseng, Y.-W.: BiSeNet V3: bilateral segmentation network with coordinate attention for real-time semantic segmentation. Neurocomputing 532, 33\u201342 (2023)","journal-title":"Neurocomputing"},{"issue":"3","key":"2110_CR56","first-page":"3448","volume":"24","author":"Y Hong","year":"2022","unstructured":"Hong, Y., Pan, H., Sun, W., Jia, Y.: Deep dual-resolution networks for real-time and accurate semantic segmentation of road scenes. IEEE Trans. Intell. Transp. Syst. 24(3), 3448\u20133460 (2022)","journal-title":"IEEE Trans. Intell. Transp. Syst."},{"key":"2110_CR57","first-page":"7423","volume":"35","author":"J Wang","year":"2022","unstructured":"Wang, J., Gou, C., Wu, Q., Feng, H., Han, J., Ding, E., Wang, J.: RTFormer: efficient design for real-time semantic segmentation with transformer. Adv. Neural. Inf. Process. Syst. 35, 7423\u20137436 (2022)","journal-title":"Adv. Neural. Inf. Process. Syst."},{"key":"2110_CR58","doi-asserted-by":"crossref","unstructured":"Xu, J., Xiong, Z., Bhattacharyya, S.P.: PIDNet: a real-time semantic segmentation network inspired by pid controllers. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 19529\u201319539 (2023)","DOI":"10.1109\/CVPR52729.2023.01871"},{"key":"2110_CR59","doi-asserted-by":"crossref","unstructured":"Xu, Z., Wu, D., Yu, C., Chu, X., Sang, N., Gao, C.: SCTNet: single-branch cnn with transformer semantic information for real-time segmentation. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 38, pp. 6378\u20136386 (2024)","DOI":"10.1609\/aaai.v38i6.28457"},{"issue":"5","key":"2110_CR60","doi-asserted-by":"publisher","first-page":"513","DOI":"10.3390\/diagnostics15050513","volume":"15","author":"S Umirzakova","year":"2025","unstructured":"Umirzakova, S., Muksimova, S., Baltayev, J., Cho, Y.I.: Force map-enhanced segmentation of a lightweight model for the early detection of cervical cancer. Diagnostics 15(5), 513 (2025)","journal-title":"Diagnostics"},{"key":"2110_CR61","doi-asserted-by":"crossref","unstructured":"Huang, S., Lu, Z., Cheng, R., He, C.: FaPN: feature-aligned pyramid network for dense image prediction. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 864\u2013873 (2021)","DOI":"10.1109\/ICCV48922.2021.00090"},{"key":"2110_CR62","doi-asserted-by":"crossref","unstructured":"Li, X., You, A., Zhu, Z., Zhao, H., Yang, M., Yang, K., Tan, S., Tong, Y.: Semantic flow for fast and accurate scene parsing. In: Proceedings of the European Conference on Computer Vision, pp. 775\u2013793 (2020)","DOI":"10.1007\/978-3-030-58452-8_45"},{"issue":"6","key":"2110_CR63","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1007\/s40747-023-01063-x","volume":"9","author":"Y Dong","year":"2023","unstructured":"Dong, Y., Yang, H., Pei, Y., Shen, L., Zheng, L., Li, P.: Compact interactive dual-branch network for real-time semantic segmentation. Complex Intell. Syst. 9(6), 1\u201314 (2023)","journal-title":"Complex Intell. Syst."},{"key":"2110_CR64","unstructured":"Wang, Y., Chen, S., Bian, H., Li, W., Lu, Q.: Spatial-assistant encoder-decoder network for real time semantic segmentation. arXiv preprint arXiv:2309.10519 (2023)"},{"key":"2110_CR65","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., Sun, J.: Deep residual learning for image recognition. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 770\u2013778 (2016)","DOI":"10.1109\/CVPR.2016.90"},{"key":"2110_CR66","unstructured":"Peng, J., Liu, Y., Tang, S., Hao, Y., Chu, L., Chen, G., Wu, Z., Chen, Z., Yu, Z., Du, Y., et al.: PP-LiteSeg: a superior real-time semantic segmentation model. arXiv preprint arXiv:2204.02681 (2022)"},{"key":"2110_CR67","doi-asserted-by":"crossref","unstructured":"Misra, D., Nalamada, T., Arasanipalai, A.U., Hou, Q.: Rotate to attend: Convolutional triplet attention module. In: Proceedings of the IEEE\/CVF Winter Conference on Applications of Computer Vision, pp. 3139\u20133148 (2021)","DOI":"10.1109\/WACV48630.2021.00318"},{"key":"2110_CR68","doi-asserted-by":"crossref","unstructured":"Cordts, M., Omran, M., Ramos, S., Rehfeld, T., Enzweiler, M., Benenson, R., Franke, U., Roth, S., Schiele, B.: The cityscapes dataset for semantic urban scene understanding. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 3213\u20133223 (2016)","DOI":"10.1109\/CVPR.2016.350"},{"issue":"2","key":"2110_CR69","doi-asserted-by":"publisher","first-page":"88","DOI":"10.1016\/j.patrec.2008.04.005","volume":"30","author":"GJ Brostow","year":"2009","unstructured":"Brostow, G.J., Fauqueur, J., Cipolla, R.: Semantic object classes in video: a high-definition ground truth database. Pattern Recogn. Lett. 30(2), 88\u201397 (2009)","journal-title":"Pattern Recogn. Lett."},{"key":"2110_CR70","doi-asserted-by":"publisher","first-page":"302","DOI":"10.1007\/s11263-018-1140-0","volume":"127","author":"B Zhou","year":"2019","unstructured":"Zhou, B., Zhao, H., Puig, X., Xiao, T., Fidler, S., Barriuso, A., Torralba, A.: Semantic understanding of scenes through the ade20k dataset. Int. J. Comput. Vision 127, 302\u2013321 (2019)","journal-title":"Int. J. Comput. Vision"},{"key":"2110_CR71","unstructured":"Vanholder, H.: Efficient inference with tensorrt. In: Proceedings of the GPU Technology Conference, vol. 1 (2016)"},{"key":"2110_CR72","unstructured":"Mao, A., Mohri, M., Zhong, Y.: Cross-entropy loss functions: Theoretical analysis and applications. In: International Conference on Machine Learning, PMLR. pp. 23803\u201323828 (2023)"},{"key":"2110_CR73","doi-asserted-by":"crossref","unstructured":"Shrivastava, A., Gupta, A., Girshick, R.: Training region-based object detectors with online hard example mining. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 761\u2013769 (2016)","DOI":"10.1109\/CVPR.2016.89"},{"key":"2110_CR74","unstructured":"Wan, Q., Huang, Z., Lu, J., Yu, G., Zhang, L.: SeaFormer: squeeze-enhanced axial transformer for mobile semantic segmentation. arXiv preprint arXiv:2301.13156 (2023)"},{"key":"2110_CR75","doi-asserted-by":"publisher","DOI":"10.1016\/j.engappai.2024.107988","volume":"133","author":"X Song","year":"2024","unstructured":"Song, X., Fang, X., Meng, X., Fang, X., Lv, M., Zhuo, Y.: Real-time semantic segmentation network with an enhanced backbone based on atrous spatial pyramid pooling module. Eng. Appl. Artif. Intell. 133, 107988 (2024)","journal-title":"Eng. Appl. Artif. Intell."},{"key":"2110_CR76","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2025.130099","volume":"637","author":"S Jia","year":"2025","unstructured":"Jia, S., Dong, Y., Mao, C., Zheng, L., Li, Y., Liu, K.: Multi-path feature enhancement network for real-time semantic segmentation. Neurocomputing 637, 130099 (2025)","journal-title":"Neurocomputing"},{"key":"2110_CR77","doi-asserted-by":"publisher","DOI":"10.1016\/j.patcog.2025.112019","volume":"170","author":"Y Dong","year":"2026","unstructured":"Dong, Y., Mao, C., Zheng, L., Wu, Q., Zhang, M., Li, X.: AFPN: Alignment feature pyramid network for real-time semantic segmentation. Pattern Recogn. 170, 112019 (2026)","journal-title":"Pattern Recogn."},{"issue":"1","key":"2110_CR78","doi-asserted-by":"publisher","first-page":"872","DOI":"10.1038\/s41598-024-84685-6","volume":"15","author":"B Ye","year":"2025","unstructured":"Ye, B., Xue, R., Wu, Q.: A hybrid attention multi-scale fusion network for real-time semantic segmentation. Sci. Rep. 15(1), 872 (2025)","journal-title":"Sci. Rep."},{"key":"2110_CR79","doi-asserted-by":"crossref","unstructured":"Ye, Z., Yan, H., Sun, Y., Li, B., Liu, L., Wu, W.: MSPNet: real-time semantic segmentation with large kernel and atrous convolutions. The Visual Computer, 1\u201316 (2025)","DOI":"10.1007\/s00371-025-03853-5"},{"key":"2110_CR80","doi-asserted-by":"crossref","unstructured":"Luo, H., Liu, C., Shark, L.-K.: Saba: Scale-adaptive attention and boundary aware network for real-time semantic segmentation. Expert Syst. Appl. 127680 (2025)","DOI":"10.1016\/j.eswa.2025.127680"},{"key":"2110_CR81","doi-asserted-by":"crossref","unstructured":"Xing, J., Jia, S., Zheng, L., Dong, Y.: Attention-nested dual-branch network in real-time semantic segmentation. Int. J. Mach. Learn. Cybernet. 1\u201316 (2025)","DOI":"10.1007\/s13042-025-02774-y"},{"key":"2110_CR82","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2024.128991","volume":"617","author":"Y Dong","year":"2025","unstructured":"Dong, Y., Mao, C., Zheng, L., Wu, Q.: DMANet: dual-branch multiscale attention network for real-time semantic segmentation. Neurocomputing 617, 128991 (2025)","journal-title":"Neurocomputing"},{"key":"2110_CR83","doi-asserted-by":"crossref","unstructured":"Li, H., Xiong, P., Fan, H., Sun, J.: DFANet: deep feature aggregation for real-time semantic segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 9522\u20139531 (2019)","DOI":"10.1109\/CVPR.2019.00975"},{"key":"2110_CR84","doi-asserted-by":"crossref","unstructured":"Cavagnero, N., Rosi, G., Cuttano, C., Pistilli, F., Ciccone, M., Averta, G., Cermelli, F.: Pem: Prototype-based efficient maskformer for image segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 15804\u201315813 (2024)","DOI":"10.1109\/CVPR52733.2024.01496"},{"key":"2110_CR85","unstructured":"Yan, H., Wu, M., Zhang, C.: Multi-scale representations by varying window attention for semantic segmentation. In: Proceedings of the International Conference on Learning Representations (2024)"}],"container-title":["Multimedia Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-025-02110-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00530-025-02110-y","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-025-02110-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,2,11]],"date-time":"2026-02-11T04:19:11Z","timestamp":1770783551000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00530-025-02110-y"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,12,15]]},"references-count":85,"journal-issue":{"issue":"1","published-print":{"date-parts":[[2026,2]]}},"alternative-id":["2110"],"URL":"https:\/\/doi.org\/10.1007\/s00530-025-02110-y","relation":{},"ISSN":["0942-4962","1432-1882"],"issn-type":[{"value":"0942-4962","type":"print"},{"value":"1432-1882","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,12,15]]},"assertion":[{"value":"5 March 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"19 November 2025","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"15 December 2025","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"49"}}