{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,3]],"date-time":"2026-06-03T04:13:38Z","timestamp":1780460018328,"version":"3.54.1"},"reference-count":67,"publisher":"Springer Science and Business Media LLC","issue":"1","license":[{"start":{"date-parts":[[2024,4,2]],"date-time":"2024-04-02T00:00:00Z","timestamp":1712016000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,4,2]],"date-time":"2024-04-02T00:00:00Z","timestamp":1712016000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100012166","name":"National Key Research and Development Program of China","doi-asserted-by":"publisher","award":["2022ZD0160400"],"award-info":[{"award-number":["2022ZD0160400"]}],"id":[{"id":"10.13039\/501100012166","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100001809","name":"Natural Science Foundation of China","doi-asserted-by":"crossref","award":["62271346"],"award-info":[{"award-number":["62271346"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"crossref"}]},{"DOI":"10.13039\/501100006606","name":"Natural Science Foundation of Tianjin City","doi-asserted-by":"publisher","award":["21JCQNJC00420"],"award-info":[{"award-number":["21JCQNJC00420"]}],"id":[{"id":"10.13039\/501100006606","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Vis Comput"],"published-print":{"date-parts":[[2025,1]]},"DOI":"10.1007\/s00371-024-03336-z","type":"journal-article","created":{"date-parts":[[2024,4,2]],"date-time":"2024-04-02T20:25:04Z","timestamp":1712089504000},"page":"423-435","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":2,"title":["SSNet: a joint learning network for semantic segmentation and disparity estimation"],"prefix":"10.1007","volume":"41","author":[{"given":"Dayu","family":"Jia","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yanwei","family":"Pang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jiale","family":"Cao","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Pan","family":"Jing","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2024,4,2]]},"reference":[{"key":"3336_CR1","doi-asserted-by":"publisher","first-page":"483","DOI":"10.1109\/TIP.2022.3229614","volume":"32","author":"W Yu","year":"2023","unstructured":"Yu, W., Zhu, M., Wang, N., Wang, X., Gao, X.: An efficient transformer based on global and local self-attention for face photo-sketch synthesis. IEEE Trans. Image Process. 32, 483\u2013495 (2023)","journal-title":"IEEE Trans. Image Process."},{"key":"3336_CR2","doi-asserted-by":"crossref","unstructured":"Li, H., Wang, N., Yang, X., Wang, X., Gao, X.: Towards semi-supervised deep facial expression recognition with an adaptive confidence margin. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition (2022)","DOI":"10.1109\/CVPR52688.2022.00413"},{"key":"3336_CR3","doi-asserted-by":"publisher","first-page":"597","DOI":"10.1007\/s00371-021-02360-7","volume":"39","author":"Z Lin","year":"2022","unstructured":"Lin, Z., Sun, W., Tang, B., Li, J., Yao, X., Li, Y.: Semantic segmentation network with multi-path structure, attention reweighting and multi-scale encoding. Vis. Comput. 39, 597 (2022)","journal-title":"Vis. Comput."},{"key":"3336_CR4","doi-asserted-by":"publisher","first-page":"2473","DOI":"10.1007\/s00371-021-02124-3","volume":"38","author":"M Jiang","year":"2021","unstructured":"Jiang, M., Zhai, F., Kong, J.: Sparse attention module for optimizing semantic segmentation performance combined with a multi-task feature extraction network. Vis. Comput. 38, 2473 (2021)","journal-title":"Vis. Comput."},{"key":"3336_CR5","doi-asserted-by":"publisher","first-page":"1427","DOI":"10.1007\/s00371-018-01622-1","volume":"35","author":"L Tian","year":"2019","unstructured":"Tian, L., Liu, J., Ling, H., Guo, W.: Disparity estimation in stereo video sequence with adaptive spatiotemporally consistent constraints. Vis. Comput. 35, 1427 (2019)","journal-title":"Vis. Comput."},{"key":"3336_CR6","doi-asserted-by":"publisher","first-page":"3881","DOI":"10.1007\/s00371-021-02228-w","volume":"38","author":"X Li","year":"2021","unstructured":"Li, X., Fan, Y., Lv, G., Ma, H.: Area-based correlation and non-local attention network for stereo matching. Vis. Comput. 38, 3881 (2021)","journal-title":"Vis. Comput."},{"key":"3336_CR7","doi-asserted-by":"crossref","unstructured":"Li, Y., Huang, J.-B., Ahuja, N., Yang, M.-H.: Deep joint image filtering. In: Proceedings of European Conference on Computer Vision, pp. 154\u2013169. Springer (2016)","DOI":"10.1007\/978-3-319-46493-0_10"},{"issue":"11","key":"3336_CR8","first-page":"8355","volume":"44","author":"J Dong","year":"2022","unstructured":"Dong, J., Pan, J., Ren, J.S., Lin, L., Tang, J., Yang, M.-H.: Learning spatially variant linear representation models for joint filtering. IEEE Trans. Pattern Anal. Mach. Intell. 44(11), 8355\u20138370 (2022)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"3336_CR9","doi-asserted-by":"crossref","unstructured":"Zhang, F., Prisacariu, V., Yang, R., Torr, P.H.: Ga-net: guided aggregation net for end-to-end stereo matching. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 185\u2013194 (2019)","DOI":"10.1109\/CVPR.2019.00027"},{"key":"3336_CR10","doi-asserted-by":"crossref","unstructured":"Xu, H., Zhang, J.: Aanet: adaptive aggregation network for efficient stereo matching. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 1956\u20131965 (2020)","DOI":"10.1109\/CVPR42600.2020.00203"},{"key":"3336_CR11","doi-asserted-by":"crossref","unstructured":"Yang, G., Zhao, H., Shi, J., Deng, Z., Jia, J.: Segstereo: exploiting semantic information for disparity estimation. In: Proceedings of the European Conference on Computer Vision, pp. 636\u2013651 (2018)","DOI":"10.1007\/978-3-030-01234-2_39"},{"key":"3336_CR12","doi-asserted-by":"crossref","unstructured":"Zhan, W., Ou, X., Yang, Y., Chen, L.: Dsnet: joint learning for scene segmentation and disparity estimation. In: Proceedings of the IEEE International Conference on Robotics and Automation, pp. 2946\u20132952 (2019)","DOI":"10.1109\/ICRA.2019.8793573"},{"key":"3336_CR13","doi-asserted-by":"crossref","unstructured":"Liu, Z., Lin, Y., Cao, Y., Hu, H., Wei, Y., Zhang, Z., Lin, S., Guo, B.: Swin transformer: hierarchical vision transformer using shifted windows. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 9992\u201310002 (2021)","DOI":"10.1109\/ICCV48922.2021.00986"},{"key":"3336_CR14","unstructured":"Krizhevsky, A., Sutskever, I., Hinton, G.E.: Imagenet classification with deep convolutional neural networks. In: Advances in Neural Information Processing Systems, pp. 1106\u20131114 (2012)"},{"key":"3336_CR15","unstructured":"Simonyan, K., Zisserman, A.: Very deep convolutional networks for large-scale image recognition. arXiv:1409.1556 (2014)"},{"key":"3336_CR16","doi-asserted-by":"crossref","unstructured":"Long, J., Shelhamer, E., Darrell, T.: Fully convolutional networks for semantic segmentation. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, pp. 3431\u20133440 (2015)","DOI":"10.1109\/CVPR.2015.7298965"},{"key":"3336_CR17","unstructured":"Badrinarayanan, V., Kendall, A., Cipolla, R.: Segnet: a deep convolutional encoder-decoder architecture for image segmentation. arXiv:1511.00561 (2015)"},{"key":"3336_CR18","doi-asserted-by":"crossref","unstructured":"Noh, H., Hong, S., Han, B.: Learning deconvolution network for semantic segmentation. In: Proceedings of IEEE International Conference on Computer Vision, pp. 1520\u20131528 (2015)","DOI":"10.1109\/ICCV.2015.178"},{"key":"3336_CR19","doi-asserted-by":"crossref","unstructured":"Lin, G., Milan, A., Shen, C., Reid, I.D.: Refinenet: multi-path refinement networks for high-resolution semantic segmentation. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, pp. 5168\u20135177 (2017)","DOI":"10.1109\/CVPR.2017.549"},{"key":"3336_CR20","doi-asserted-by":"crossref","unstructured":"Ronneberger, O., Fischer, P., Brox, T.: U-net: convolutional networks for biomedical image segmentation. In: Proceedings of International Conference on Medical Image Computing and Computer Assisted Intervention, pp. 234\u2013241 (2015)","DOI":"10.1007\/978-3-319-24574-4_28"},{"issue":"6","key":"3336_CR21","doi-asserted-by":"publisher","first-page":"1856","DOI":"10.1109\/TMI.2019.2959609","volume":"39","author":"Z Zhou","year":"2020","unstructured":"Zhou, Z., Siddiquee, M.M.R., Tajbakhsh, N., Liang, J.: Unet++: redesigning skip connections to exploit multiscale features in image segmentation. IEEE Trans. Med. Imaging 39(6), 1856\u20131867 (2020)","journal-title":"IEEE Trans. Med. Imaging"},{"key":"3336_CR22","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., Sun, J.: Spatial pyramid pooling in deep convolutional networks for visual recognition. In: Proceedings of European Conference on Computer Vision, pp. 346\u2013361. Springer (2014)","DOI":"10.1007\/978-3-319-10578-9_23"},{"key":"3336_CR23","doi-asserted-by":"crossref","unstructured":"Zhao, H., Shi, J., Qi, X., Wang, X., Jia, J.: Pyramid scene parsing network. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, pp. 6230\u20136239 (2017)","DOI":"10.1109\/CVPR.2017.660"},{"key":"3336_CR24","unstructured":"Chen, L.-C., Papandreou, G., Kokkinos, I., Murphy, K., Yuille, A.L.: Semantic image segmentation with deep convolutional nets and fully connected crfs. arXiv:1412.7062 (2014)"},{"issue":"4","key":"3336_CR25","doi-asserted-by":"publisher","first-page":"834","DOI":"10.1109\/TPAMI.2017.2699184","volume":"40","author":"L-C Chen","year":"2018","unstructured":"Chen, L.-C., Papandreou, G., Kokkinos, I., Murphy, K., Yuille, A.L.: Deeplab: semantic image segmentation with deep convolutional nets, atrous convolution, and fully connected crfs. IEEE Trans. Pattern Anal. Mach. Intell. 40(4), 834\u2013848 (2018)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"3336_CR26","unstructured":"Chen, L.-C., Papandreou, G., Schroff, F., Adam, H.: Rethinking atrous convolution for semantic image segmentation. arXiv:1706.05587 (2017)"},{"key":"3336_CR27","doi-asserted-by":"crossref","unstructured":"Chen, L.-C., Zhu, Y., Papandreou, G., Schroff, F., Adam, H.: Encoder-decoder with atrous separable convolution for semantic image segmentation. arXiv:1802.02611 (2018)","DOI":"10.1007\/978-3-030-01234-2_49"},{"key":"3336_CR28","doi-asserted-by":"crossref","unstructured":"Yang, M., Yu, K., Zhang, C., Li, Z., Yang, K.: Denseaspp for semantic segmentation in street scenes. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, pp. 3684\u20133692 (2018)","DOI":"10.1109\/CVPR.2018.00388"},{"key":"3336_CR29","doi-asserted-by":"crossref","unstructured":"Huang, G., Liu, Z., Maaten, L.V.D., Weinberger, K.Q.: Densely connected convolutional networks. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, pp. 2261\u20132269 (2017)","DOI":"10.1109\/CVPR.2017.243"},{"key":"3336_CR30","doi-asserted-by":"crossref","unstructured":"Pang, Y., Li, Y., Shen, J., Shao, L.: Towards bridging semantic gap to improve semantic segmentation. In: Proceedings of IEEE International Conference on Computer Vision, pp. 4229\u20134238 (2019)","DOI":"10.1109\/ICCV.2019.00433"},{"issue":"8","key":"3336_CR31","doi-asserted-by":"publisher","first-page":"2011","DOI":"10.1109\/TPAMI.2019.2913372","volume":"42","author":"J Hu","year":"2020","unstructured":"Hu, J., Shen, L., Albanie, S., Sun, G., Wu, E.: Squeeze-and-excitation networks. IEEE Trans. Pattern Anal. Mach. Intell. 42(8), 2011\u20132023 (2020)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"3336_CR32","doi-asserted-by":"crossref","unstructured":"Zhang, H., Dana, K., Shi, J., Zhang, Z., Wang, X., Tyagi, A., Agrawal, A.: Context encoding for semantic segmentation. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, pp. 7151\u20137160 (2018)","DOI":"10.1109\/CVPR.2018.00747"},{"key":"3336_CR33","doi-asserted-by":"crossref","unstructured":"Fu, J., Liu, J., Tian, H., Fang, Z., Lu, H.: Dual attention network for scene segmentation. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, pp. 3146\u20133154 (2019)","DOI":"10.1109\/CVPR.2019.00326"},{"key":"3336_CR34","unstructured":"Vaswani, A., Shazeer, N., Parmar, N., Uszkoreit, J., Jones, L., Gomez, A.N., Kaiser, L., Polosukhin, I.: Attention is all you need. In: Advances in Neural Information Processing Systems, pp. 6000\u20136010 (2017)"},{"key":"3336_CR35","doi-asserted-by":"crossref","unstructured":"Zheng, S., Lu, J., Zhao, H., Zhu, X., Luo, Z., Wang, Y., Fu, Y., Feng, J., Xiang, T., Torr, P.H., Zhang, L.: Rethinking semantic segmentation from a sequence-to-sequence perspective with transformers. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition (2021)","DOI":"10.1109\/CVPR46437.2021.00681"},{"key":"3336_CR36","unstructured":"Dosovitskiy, A., Beyer, L., Kolesnikov, A., Weissenborn, D., Zhai, X., Unterthiner, T., Dehghani, M., Minderer, M., Heigold, G., Gelly, S., Uszkoreit, J., Houlsby, N.: An image is worth 16x16 words: transformers for image recognition at scale. In: Proceedings of International Conference on Learning Representations (2021)"},{"key":"3336_CR37","unstructured":"Xie, E., Wang, W., Yu, Z., Anandkumar, A., Alvarez, J.M., Luo, P.: Segformer: simple and efficient design for semantic segmentation with transformers. arXiv:2105.15203 (2021)"},{"key":"3336_CR38","doi-asserted-by":"crossref","unstructured":"Strudel, R., Garcia, R., Laptev, I., Schmid, C.: Segmenter: transformer for semantic segmentation. In: Proceedings of IEEE International Conference on Computer Vision (2021)","DOI":"10.1109\/ICCV48922.2021.00717"},{"key":"3336_CR39","doi-asserted-by":"crossref","unstructured":"Valanarasu, J.M.J., Oza, P., Hacihaliloglu, I., Patel, V.M.: Medical transformer: gated axial-attention for medical image segmentation. In: Proceedings of International Conference on Medical Image Computing and Computer Assisted Intervention, pp. 36\u201346 (2021)","DOI":"10.1007\/978-3-030-87193-2_4"},{"key":"3336_CR40","first-page":"1","volume":"2016","author":"HR Affendi","year":"2016","unstructured":"Affendi, H.R., Haidi, I.: Literature survey on stereo vision disparity map algorithms. J. Sens. 2016, 1\u201323 (2016)","journal-title":"J. Sens."},{"key":"3336_CR41","doi-asserted-by":"crossref","unstructured":"Mayer, N., Ilg, E., Hausser, P., Fischer, P., Cremers, D., Dosovitskiy, A., Brox, T.: A large dataset to train convolutional networks for disparity, optical flow, and scene flow estimation. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 4040\u20134048 (2016)","DOI":"10.1109\/CVPR.2016.438"},{"key":"3336_CR42","doi-asserted-by":"crossref","unstructured":"Guo, X., Yang, K., Yang, W., Wang, X., Li, H.: Group-wise correlation stereo network. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 3273\u20133282 (2019)","DOI":"10.1109\/CVPR.2019.00339"},{"key":"3336_CR43","doi-asserted-by":"crossref","unstructured":"Chang, J.-R., Chen, Y.-S.: Pyramid stereo matching network. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 5410\u20135418 (2018)","DOI":"10.1109\/CVPR.2018.00567"},{"key":"3336_CR44","doi-asserted-by":"crossref","unstructured":"Dai, J., Qi, H., Xiong, Y., Li, Y., Zhang, G., Hu, H., Wei, Y.: Deformable convolutional networks. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 764\u2013773 (2017)","DOI":"10.1109\/ICCV.2017.89"},{"key":"3336_CR45","doi-asserted-by":"crossref","unstructured":"Li, Z., Liu, X., Drenkow, N., Ding, A., Creighton, F.X., Taylor, R.H., Unberath, M.: Revisiting stereo depth estimation from a sequence-to-sequence perspective with transformers. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 6197\u20136206 (2021)","DOI":"10.1109\/ICCV48922.2021.00614"},{"key":"3336_CR46","doi-asserted-by":"crossref","unstructured":"Li, P., Chen, X., Shen, S.: Stereo r-cnn based 3d object detection for autonomous driving. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 7644\u20137652 (2019)","DOI":"10.1109\/CVPR.2019.00783"},{"key":"3336_CR47","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., Sun, J.: Deep residual learning for image recognition. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, pp. 770\u2013778 (2016)","DOI":"10.1109\/CVPR.2016.90"},{"issue":"6","key":"3336_CR48","doi-asserted-by":"publisher","first-page":"1137","DOI":"10.1109\/TPAMI.2016.2577031","volume":"39","author":"S Ren","year":"2017","unstructured":"Ren, S., He, K., Girshick, R., Sun, J.: Faster r-cnn: towards real-time object detection with region proposal networks. IEEE Trans. Pattern Anal. Mach. Intell. 39(6), 1137\u20131149 (2017)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"3336_CR49","doi-asserted-by":"crossref","unstructured":"Qin, Z., Wang, J., Lu, Y.: Triangulation learning network: from monocular to stereo 3d object detection. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 7607\u20137615 (2019)","DOI":"10.1109\/CVPR.2019.00780"},{"key":"3336_CR50","doi-asserted-by":"crossref","unstructured":"Wang, Y., Chao, W.-L., Garg, D., Hariharan, B., Campbell, M., Weinberger, K.Q.: Pseudo-lidar from visual depth estimation: bridging the gap in 3d object detection for autonomous driving. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 8445\u20138453 (2019)","DOI":"10.1109\/CVPR.2019.00864"},{"key":"3336_CR51","unstructured":"You, Y., Wang, Y., Chao, W.-L., Garg, D., Pleiss, G., Hariharan, B., Campbell, M., Weinberger, K.Q.: Pseudo-lidar++: Accurate depth for 3d object detection in autonomous driving. CoRR arXiv:1906.06310 (2019)"},{"key":"3336_CR52","doi-asserted-by":"crossref","unstructured":"Dvornik, N., Shmelkov, K., Mairal, J., Schmid, C.: Blitznet: A real-time deep network for scene understanding. In: Proceedings of IEEE International Conference on Computer Vision, pp. 4174\u20134182 (2017)","DOI":"10.1109\/ICCV.2017.447"},{"key":"3336_CR53","doi-asserted-by":"crossref","unstructured":"Cao, J., Pang, Y., Li, X.: Triply supervised decoder networks for joint detection and segmentation. In: Proceedings of IEEE International Conference on Computer Vision, pp. 7384\u20137393 (2019)","DOI":"10.1109\/CVPR.2019.00757"},{"key":"3336_CR54","unstructured":"Zeng, Y., Zhuge, Y., Lu, H., Zhang, L.: Joint learning of saliency detection and weakly supervised semantic segmentation. In: Proceedings of International Conference on Computer Vision, pp. 7222\u20137232 (2019)"},{"key":"3336_CR55","unstructured":"Ba, J.L., Kiros, J.R., Hinton, G.E.: Layer normalization. arXiv:1607.06450 (2016)"},{"key":"3336_CR56","doi-asserted-by":"crossref","unstructured":"Ding, X., Guo, Y., Ding, G., Han, J.: Acnet: strengthening the kernel skeletons for powerful cnn via asymmetric convolution blocks. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 1911\u20131920 (2019)","DOI":"10.1109\/ICCV.2019.00200"},{"key":"3336_CR57","doi-asserted-by":"crossref","unstructured":"Liu, X., Zheng, Y., Killeen, B., Ishii, M., Hager, G.D., Taylor, R.H., Unberath, M.: Extremely dense point correspondences using a learned feature descriptor. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 4846\u20134855 (2020)","DOI":"10.1109\/CVPR42600.2020.00490"},{"key":"3336_CR58","doi-asserted-by":"crossref","unstructured":"Cordts, M., Omran, M., Ramos, S., Rehfeld, T., Enzweiler, M., Benenson, R., Franke, U., Roth, S., Schiele, B.: The cityscapes dataset for semantic urban scene understanding. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, pp. 3213\u20133223 (2016)","DOI":"10.1109\/CVPR.2016.350"},{"issue":"2","key":"3336_CR59","doi-asserted-by":"publisher","first-page":"328","DOI":"10.1109\/TPAMI.2007.1166","volume":"30","author":"H Hirschmuller","year":"2008","unstructured":"Hirschmuller, H.: Stereo processing by semiglobal matching and mutual information. IEEE Trans. Pattern Anal. Mach. Intell. 30(2), 328\u2013341 (2008)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"3336_CR60","doi-asserted-by":"crossref","unstructured":"Menze, M., Geiger, A.: Object scene flow for autonomous vehicles. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 3061\u20133070 (2015)","DOI":"10.1109\/CVPR.2015.7298925"},{"key":"3336_CR61","unstructured":"Paszke, A., Gross, S., Massa, F., Lerer, A., Bradbury, J., Chanan, G., Killeen, T., Lin, Z., Gimelshein, N., Antiga, L., etc: Pytorch: an imperative style, high-performance deep learning library. In: Advances in Neural Information Processing Systems (2019)"},{"key":"3336_CR62","unstructured":"Sun, K., Zhao, Y., Jiang, B., Cheng, T., Xiao, B., Liu, D., Mu, Y., Wang, X., Liu, W., Wang, J.: High-resolution representations for labeling pixels and regions. In: arXiv Preprint arXiv:1904.04514 (2019)"},{"key":"3336_CR63","doi-asserted-by":"crossref","unstructured":"Yuan, Y., Chen, X., Wang, J.: Object-contextual representations for semantic segmentation. In: Proceedings of European Conference on Computer Vision (2020)","DOI":"10.1007\/978-3-030-58539-6_11"},{"key":"3336_CR64","doi-asserted-by":"crossref","unstructured":"Neuhold, G., Ollmann, T., Rota\u00a0Bul\u00f2, S., Kontschieder, P.: The mapillary vistas dataset for semantic understanding of street scenes. In: Proceedings of IEEE International Conference on Computer Vision (2017)","DOI":"10.1109\/ICCV.2017.534"},{"key":"3336_CR65","doi-asserted-by":"crossref","unstructured":"Klingner, M., Term\u00f6hlen, J.-A., Mikolajczyk, J., Fingscheidt, T.: Self-supervised monocular depth estimation: solving the dynamic object problem by semantic guidance. In: Proceedings of the European Conference on Computer Vision, pp. 582\u2013600 (2020)","DOI":"10.1007\/978-3-030-58565-5_35"},{"key":"3336_CR66","doi-asserted-by":"crossref","unstructured":"Lambert, J., Liu, Z., Sener, O., Hays, J., Koltun, V.: Mseg: A composite dataset for multi-domain semantic segmentation. arXiv:2112.13762 (2021)","DOI":"10.1109\/CVPR42600.2020.00295"},{"key":"3336_CR67","unstructured":"Bevandi\u0107, P., Or\u0161i\u0107, M., Grubi\u0161i\u0107, I., \u0160ari\u0107, J., \u0160egvi\u0107, S.: Multi-domain semantic segmentation with overlapping labels. arXiv:2108.11224 (2021)"}],"container-title":["The Visual Computer"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00371-024-03336-z.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00371-024-03336-z\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00371-024-03336-z.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,1,24]],"date-time":"2025-01-24T12:58:46Z","timestamp":1737723526000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00371-024-03336-z"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,4,2]]},"references-count":67,"journal-issue":{"issue":"1","published-print":{"date-parts":[[2025,1]]}},"alternative-id":["3336"],"URL":"https:\/\/doi.org\/10.1007\/s00371-024-03336-z","relation":{},"ISSN":["0178-2789","1432-2315"],"issn-type":[{"value":"0178-2789","type":"print"},{"value":"1432-2315","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,4,2]]},"assertion":[{"value":"22 February 2024","order":1,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"2 April 2024","order":2,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}]}}