{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,12]],"date-time":"2026-06-12T23:30:27Z","timestamp":1781307027450,"version":"3.54.1"},"reference-count":43,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2024,6,20]],"date-time":"2024-06-20T00:00:00Z","timestamp":1718841600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,6,20]],"date-time":"2024-06-20T00:00:00Z","timestamp":1718841600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["61906049"],"award-info":[{"award-number":["61906049"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"name":"Guangzhou Higher Education Teaching Reform Projec","award":["2022JXGG016"],"award-info":[{"award-number":["2022JXGG016"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Vis Comput"],"published-print":{"date-parts":[[2025,2]]},"DOI":"10.1007\/s00371-024-03448-6","type":"journal-article","created":{"date-parts":[[2024,6,20]],"date-time":"2024-06-20T08:02:11Z","timestamp":1718870531000},"page":"1543-1554","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":13,"title":["ZMNet: feature fusion and semantic boundary supervision for real-time semantic segmentation"],"prefix":"10.1007","volume":"41","author":[{"given":"Ya","family":"Li","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Ziming","family":"Li","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Huiwang","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Qing","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2024,6,20]]},"reference":[{"key":"3448_CR1","unstructured":"Howard, A.G., Zhu, M., Chen, B., Kalenichenko, D., Wang, W., Weyand, T., Andreetto, M., Adam, H.: Mobilenets: efficient convolutional neural networks for mobile vision applications. arXiv preprint arXiv:1704.04861 (2017)"},{"key":"3448_CR2","doi-asserted-by":"crossref","unstructured":"Yu, C., Wang, J., Peng, C., Gao, C., Yu, G., Sang, N.: Bisenet: bilateral segmentation network for real-time semantic segmentation. In: Proceedings of the European Conference on Computer Vision (ECCV), pp. 325\u2013341 (2018)","DOI":"10.1007\/978-3-030-01261-8_20"},{"issue":"11","key":"3448_CR3","doi-asserted-by":"publisher","first-page":"3051","DOI":"10.1007\/s11263-021-01515-2","volume":"129","author":"C Yu","year":"2021","unstructured":"Yu, C., Gao, C., Wang, J., Yu, G., Shen, C., Sang, N.: Bisenet v2: bilateral network with guided aggregation for real-time semantic segmentation. Int. J. Comput. Vis. 129(11), 3051\u20133068 (2021)","journal-title":"Int. J. Comput. Vis."},{"key":"3448_CR4","doi-asserted-by":"crossref","unstructured":"Fan, M., Lai, S., Huang, J., Wei, X., Chai, Z., Luo, J., Wei, X.: Rethinking bisenet for real-time semantic segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 9716\u20139725 (2021)","DOI":"10.1109\/CVPR46437.2021.00959"},{"key":"3448_CR5","doi-asserted-by":"crossref","unstructured":"Dai, Y., Gieseke, F., Oehmcke, S., Wu, Y., Barnard, K.: Attentional feature fusion. In: Proceedings of the IEEE\/CVF Winter Conference on Applications of Computer Vision, pp. 3560\u20133569 (2021)","DOI":"10.1109\/WACV48630.2021.00360"},{"key":"3448_CR6","doi-asserted-by":"crossref","unstructured":"Yang, M., Yu, K., Zhang, C., Li, Z., Yang, K.: Denseaspp for semantic segmentation in street scenes. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 3684\u20133692 (2018)","DOI":"10.1109\/CVPR.2018.00388"},{"key":"3448_CR7","doi-asserted-by":"crossref","unstructured":"Yu, F., Wang, D., Shelhamer, E., Darrell, T.: Deep layer aggregation. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 2403\u20132412 (2018)","DOI":"10.1109\/CVPR.2018.00255"},{"issue":"2","key":"3448_CR8","doi-asserted-by":"publisher","first-page":"22","DOI":"10.1007\/s00138-023-01373-7","volume":"34","author":"H Yin","year":"2023","unstructured":"Yin, H., Xie, W., Zhang, J., Zhang, Y., Zhu, W., Gao, J., Shao, Y., Li, Y.: Dual context network for real-time semantic segmentation. Mach. Vis. Appl. 34(2), 22 (2023)","journal-title":"Mach. Vis. Appl."},{"key":"3448_CR9","doi-asserted-by":"crossref","unstructured":"Zhen, M., Wang, J., Zhou, L., Li, S., Shen, T., Shang, J., Fang, T., Quan, L.: Joint semantic segmentation and boundary detection using iterative pyramid contexts. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 13666\u201313675 (2020)","DOI":"10.1109\/CVPR42600.2020.01368"},{"key":"3448_CR10","doi-asserted-by":"publisher","first-page":"27","DOI":"10.1016\/j.neucom.2022.11.084","volume":"521","author":"J Liu","year":"2023","unstructured":"Liu, J., Zhang, F., Zhou, Z., Wang, J.: Bfmnet: bilateral feature fusion network with multi-scale context aggregation for real-time semantic segmentation. Neurocomputing 521, 27\u201340 (2023)","journal-title":"Neurocomputing"},{"key":"3448_CR11","doi-asserted-by":"crossref","unstructured":"Zhang, X., Zhou, X., Lin, M., Sun, J.: Shufflenet: an extremely efficient convolutional neural network for mobile devices. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 6848\u20136856 (2018)","DOI":"10.1109\/CVPR.2018.00716"},{"key":"3448_CR12","unstructured":"Poudel, R.P., Liwicki, S., Cipolla, R.: Fast-scnn: fast semantic segmentation network. arXiv preprint arXiv:1902.04502 (2019)"},{"key":"3448_CR13","doi-asserted-by":"publisher","first-page":"205","DOI":"10.1016\/j.neucom.2022.11.075","volume":"520","author":"D Luo","year":"2023","unstructured":"Luo, D., Kang, H., Long, J., Zhang, J., Liu, X., Quan, T.: Gdn: guided down-sampling network for real-time semantic segmentation. Neurocomputing 520, 205\u2013215 (2023)","journal-title":"Neurocomputing"},{"key":"3448_CR14","unstructured":"Hong, Y., Pan, H., Sun, W., Jia, Y.: Deep dual-resolution networks for real-time and accurate semantic segmentation of road scenes. arXiv preprint arXiv:2101.06085 (2021)"},{"key":"3448_CR15","doi-asserted-by":"crossref","unstructured":"Zhang, W., Huang, Z., Luo, G., Chen, T., Wang, X., Liu, W., Yu, G., Shen, C.: Topformer: Token pyramid transformer for mobile semantic segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 12083\u201312093 (2022)","DOI":"10.1109\/CVPR52688.2022.01177"},{"issue":"7","key":"3448_CR16","doi-asserted-by":"publisher","first-page":"4444","DOI":"10.1109\/TCSVT.2021.3121680","volume":"32","author":"X Weng","year":"2021","unstructured":"Weng, X., Yan, Y., Chen, S., Xue, J.-H., Wang, H.: Stage-aware feature alignment network for real-time semantic segmentation of street scenes. IEEE Trans. Circuits Syst. Video Technol. 32(7), 4444\u20134459 (2021)","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"issue":"10","key":"3448_CR17","doi-asserted-by":"publisher","first-page":"17224","DOI":"10.1109\/TITS.2022.3150350","volume":"23","author":"X Weng","year":"2022","unstructured":"Weng, X., Yan, Y., Dong, G., Shu, C., Wang, B., Wang, H., Zhang, J.: Deep multi-branch aggregation network for real-time semantic segmentation in street scenes. IEEE Trans. Intell. Transp. Syst. 23(10), 17224\u201317240 (2022)","journal-title":"IEEE Trans. Intell. Transp. Syst."},{"key":"3448_CR18","first-page":"1438","volume":"36","author":"Y Li","year":"2022","unstructured":"Li, Y., Chang, Y., Yu, C., Yan, L.: Close the loop: a unified bottom-up and top-down paradigm for joint image deraining and segmentation. Proc. AAAI Conf. Artif. Intell. 36, 1438\u20131446 (2022)","journal-title":"Proc. AAAI Conf. Artif. Intell."},{"key":"3448_CR19","doi-asserted-by":"crossref","unstructured":"Zhao, S., Huang, W., Yang, M., Liu, W.: real rainy scene analysis: A dual-module benchmark for image deraining and segmentation. In: 2023 IEEE International Conference on Multimedia and Expo Workshops (ICMEW), 69\u201374 (2023). IEEE","DOI":"10.1109\/ICMEW59549.2023.00018"},{"key":"3448_CR20","doi-asserted-by":"crossref","unstructured":"Sun, S., Ren, W., Li, J., Zhang, K., Liang, M., Cao, X.: Event-aware video deraining via multi-patch progressive learning. IEEE Trans. Image Process. (2023)","DOI":"10.1109\/TIP.2023.3272283"},{"key":"3448_CR21","doi-asserted-by":"crossref","unstructured":"Wang, F., Jiang, M., Qian, C., Yang, S., Li, C., Zhang, H., Wang, X., Tang, X.: Residual attention network for image classification. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 3156\u20133164 (2017)","DOI":"10.1109\/CVPR.2017.683"},{"key":"3448_CR22","doi-asserted-by":"crossref","unstructured":"Hu, J., Shen, L., Sun, G.: Squeeze-and-excitation networks. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 7132\u20137141 (2018)","DOI":"10.1109\/CVPR.2018.00745"},{"key":"3448_CR23","unstructured":"Park, J., Woo, S., Lee, J.-Y., Kweon, I.S.: Bam: Bottleneck attention module. arXiv preprint arXiv:1807.06514 (2018)"},{"key":"3448_CR24","doi-asserted-by":"crossref","unstructured":"Woo, S., Park, J., Lee, J.-Y., Kweon, I.S.: Cbam: convolutional block attention module. In: Proceedings of the European Conference on Computer Vision (ECCV), pp. 3\u201319 (2018)","DOI":"10.1007\/978-3-030-01234-2_1"},{"key":"3448_CR25","doi-asserted-by":"crossref","unstructured":"Zhang, H., Wu, C., Zhang, Z., Zhu, Y., Lin, H., Zhang, Z., Sun, Y., He, T., Mueller, J., Manmatha, R., : Resnest: split-attention networks. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 2736\u20132746 (2022)","DOI":"10.1109\/CVPRW56347.2022.00309"},{"key":"3448_CR26","doi-asserted-by":"crossref","unstructured":"Huang, Y., Shi, P., He, H., He, H., Zhao, B.: Senet: spatial information enhancement for semantic segmentation neural networks. Vis. Comput. 1\u201314 (2023)","DOI":"10.1007\/s00371-023-03043-1"},{"issue":"7","key":"3448_CR27","doi-asserted-by":"publisher","first-page":"2473","DOI":"10.1007\/s00371-021-02124-3","volume":"38","author":"M Jiang","year":"2022","unstructured":"Jiang, M., Zhai, F., Kong, J.: Sparse attention module for optimizing semantic segmentation performance combined with a multi-task feature extraction network. Vis. Comput. 38(7), 2473\u20132488 (2022)","journal-title":"Vis. Comput."},{"key":"3448_CR28","doi-asserted-by":"crossref","unstructured":"Li, X., You, A., Zhu, Z., Zhao, H., Yang, M., Yang, K., Tan, S., Tong, Y.: Semantic flow for fast and accurate scene parsing. In: European Conference on Computer Vision, pp. 775\u2013793 (2020). Springer","DOI":"10.1007\/978-3-030-58452-8_45"},{"key":"3448_CR29","doi-asserted-by":"crossref","unstructured":"Li, H., Xiong, P., Fan, H., Sun, J.: Dfanet: deep feature aggregation for real-time semantic segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 9522\u20139531 (2019)","DOI":"10.1109\/CVPR.2019.00975"},{"issue":"1","key":"3448_CR30","first-page":"550","volume":"44","author":"Z Huang","year":"2021","unstructured":"Huang, Z., Wei, Y., Wang, X., Liu, W., Huang, T.S., Shi, H.: Alignseg: feature-aligned segmentation networks. IEEE Trans. Pattern Anal. Mach. Intell. 44(1), 550\u2013557 (2021)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"3448_CR31","doi-asserted-by":"crossref","unstructured":"Shrivastava, A., Gupta, A., Girshick, R.: Training region-based object detectors with online hard example mining. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 761\u2013769 (2016)","DOI":"10.1109\/CVPR.2016.89"},{"key":"3448_CR32","doi-asserted-by":"crossref","unstructured":"Lin, T.-Y., Goyal, P., Girshick, R., He, K., Doll\u00e1r, P.: Focal loss for dense object detection. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 2980\u20132988 (2017)","DOI":"10.1109\/ICCV.2017.324"},{"key":"3448_CR33","doi-asserted-by":"crossref","unstructured":"Cordts, M., Omran, M., Ramos, S., Rehfeld, T., Enzweiler, M., Benenson, R., Franke, U., Roth, S., Schiele, B.: The cityscapes dataset for semantic urban scene understanding. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 3213\u20133223 (2016)","DOI":"10.1109\/CVPR.2016.350"},{"key":"3448_CR34","doi-asserted-by":"crossref","unstructured":"Brostow, G.J., Shotton, J., Fauqueur, J., Cipolla, R.: Segmentation and recognition using structure from motion point clouds. In: European Conference on Computer Vision, pp. 44\u201357 (2008). Springer","DOI":"10.1007\/978-3-540-88682-2_5"},{"key":"3448_CR35","doi-asserted-by":"crossref","unstructured":"Orsic, M., Kreso, I., Bevandic, P., Segvic, S.: In defense of pre-trained imagenet architectures for real-time semantic segmentation of road-driving images. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 12607\u201312616 (2019)","DOI":"10.1109\/CVPR.2019.01289"},{"issue":"6","key":"3448_CR36","doi-asserted-by":"publisher","first-page":"84","DOI":"10.1145\/3065386","volume":"60","author":"A Krizhevsky","year":"2017","unstructured":"Krizhevsky, A., Sutskever, I., Hinton, G.E.: Imagenet classification with deep convolutional neural networks. Commun. ACM 60(6), 84\u201390 (2017)","journal-title":"Commun. ACM"},{"issue":"4","key":"3448_CR37","doi-asserted-by":"publisher","first-page":"834","DOI":"10.1109\/TPAMI.2017.2699184","volume":"40","author":"L-C Chen","year":"2017","unstructured":"Chen, L.-C., Papandreou, G., Kokkinos, I., Murphy, K., Yuille, A.L.: Deeplab: semantic image segmentation with deep convolutional nets, atrous convolution, and fully connected crfs. IEEE Trans. Pattern Anal. Mach. Intell. 40(4), 834\u2013848 (2017)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"3448_CR38","unstructured":"Chen, L.-C., Papandreou, G., Schroff, F., Adam, H.: rethinking atrous convolution for semantic image segmentation. arXiv preprint arXiv:1706.05587 (2017)"},{"key":"3448_CR39","unstructured":"Li, P., Dong, X., Yu, X., Yang, Y.: When humans meet machines: towards efficient segmentation networks. In: The 31st British Machine Vision Virtual Conference (2020)"},{"key":"3448_CR40","doi-asserted-by":"publisher","first-page":"94","DOI":"10.1016\/j.neucom.2022.08.049","volume":"509","author":"P Sheng","year":"2022","unstructured":"Sheng, P., Shi, Y., Liu, X., Jin, H.: Lsnet: real-time attention semantic segmentation network with linear complexity. Neurocomputing 509, 94\u2013101 (2022)","journal-title":"Neurocomputing"},{"key":"3448_CR41","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., Sun, J.: Deep residual learning for image recognition. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 770\u2013778 (2016)","DOI":"10.1109\/CVPR.2016.90"},{"key":"3448_CR42","doi-asserted-by":"crossref","unstructured":"Deng, J., Dong, W., Socher, R., Li, L.-J., Li, K., Fei-Fei, L.: Imagenet: a large-scale hierarchical image database. In: 2009 IEEE Conference on Computer Vision and Pattern Recognition, pp. 248\u2013255 (2009). IEEE","DOI":"10.1109\/CVPR.2009.5206848"},{"key":"3448_CR43","unstructured":"Paszke, A., Chaurasia, A., Kim, S., Culurciello, E.: Enet: a deep neural network architecture for real-time semantic segmentation. arXiv preprint arXiv:1606.02147 (2016)"}],"container-title":["The Visual Computer"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00371-024-03448-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00371-024-03448-6\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00371-024-03448-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,2,12]],"date-time":"2025-02-12T14:56:02Z","timestamp":1739372162000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00371-024-03448-6"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,6,20]]},"references-count":43,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2025,2]]}},"alternative-id":["3448"],"URL":"https:\/\/doi.org\/10.1007\/s00371-024-03448-6","relation":{},"ISSN":["0178-2789","1432-2315"],"issn-type":[{"value":"0178-2789","type":"print"},{"value":"1432-2315","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,6,20]]},"assertion":[{"value":"1 May 2024","order":1,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"20 June 2024","order":2,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}]}}