{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,20]],"date-time":"2026-05-20T18:04:49Z","timestamp":1779300289190,"version":"3.51.4"},"reference-count":33,"publisher":"Springer Science and Business Media LLC","issue":"2","license":[{"start":{"date-parts":[[2020,5,30]],"date-time":"2020-05-30T00:00:00Z","timestamp":1590796800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springer.com\/tdm"},{"start":{"date-parts":[[2020,5,30]],"date-time":"2020-05-30T00:00:00Z","timestamp":1590796800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springer.com\/tdm"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["J Intell Robot Syst"],"published-print":{"date-parts":[[2020,11]]},"DOI":"10.1007\/s10846-020-01205-0","type":"journal-article","created":{"date-parts":[[2020,5,30]],"date-time":"2020-05-30T05:02:42Z","timestamp":1590814962000},"page":"455-463","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":15,"title":["Semi-Supervised Monocular Depth Estimation Based on Semantic Supervision"],"prefix":"10.1007","volume":"100","author":[{"ORCID":"https:\/\/orcid.org\/0000-0003-0808-2963","authenticated-orcid":false,"given":"Min","family":"Yue","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Guangyuan","family":"Fu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Ming","family":"Wu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Hongqiao","family":"Wang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2020,5,30]]},"reference":[{"issue":"5","key":"1205_CR1","doi-asserted-by":"publisher","first-page":"1147","DOI":"10.1109\/TRO.2015.2463671","volume":"31","author":"R Mur-Artal","year":"2015","unstructured":"Mur-Artal, R., Montiel, J.M.M., Tard\u00f3s, J.D.: ORB-SLAM: a versatile and accurate monocular SLAM system. IEEE Trans. Robot. 31(5), 1147\u20131163 (2015)","journal-title":"IEEE Trans. Robot."},{"issue":"5","key":"1205_CR2","doi-asserted-by":"publisher","first-page":"1255","DOI":"10.1109\/TRO.2017.2705103","volume":"33","author":"R Mur-Artal","year":"2017","unstructured":"Mur-Artal, R., Tard\u00f3s, J.D.: ORB-SLAM2: an open-source SLAM system for monocular, stereo, and RGB-D cameras. IEEE Trans. Robot. 33(5), 1255\u20131262 (2017)","journal-title":"IEEE Trans. Robot."},{"key":"1205_CR3","doi-asserted-by":"crossref","unstructured":"G. Klein and D. Murray, \"Parallel Tracking and Mapping for Small AR Workspaces,\" Presented at the Proceedings of the 2007 6th IEEE and ACM International Symposium on Mixed and Augmented Reality, 2007","DOI":"10.1109\/ISMAR.2007.4538852"},{"issue":"2","key":"1205_CR4","doi-asserted-by":"publisher","first-page":"249","DOI":"10.1109\/TRO.2016.2623335","volume":"33","author":"C Forster","year":"2017","unstructured":"Forster, C., Zhang, Z., Gassner, M., Werlberger, M., Scaramuzza, D.: SVO: Semidirect visual Odometry for monocular and multicamera systems. IEEE Trans. Robot. 33(2), 249\u2013265 (2017)","journal-title":"IEEE Trans. Robot."},{"key":"1205_CR5","doi-asserted-by":"crossref","unstructured":"J. Engel, T. Sch\u00f6ps, and D. Cremers, \"LSD-SLAM: Large-Scale Direct Monocular SLAM,\" in Computer Vision \u2013 ECCV 2014, Cham, 2014, pp. 834\u2013849: Springer International Publishing","DOI":"10.1007\/978-3-319-10605-2_54"},{"key":"1205_CR6","doi-asserted-by":"crossref","unstructured":"J. Jiao, Y. Cao, Y. Song, and R. Lau, \"Look deeper into depth: Monocular depth estimation with semantic booster and attention-driven loss,\" in Proceedings of the European Conference on Computer Vision (ECCV), 2018, pp. 53\u201369","DOI":"10.1007\/978-3-030-01267-0_4"},{"key":"1205_CR7","doi-asserted-by":"crossref","unstructured":"W. Chen, S. Qian, and J. Deng, \"Learning single-image depth from videos using quality assessment networks,\" in Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 2019, pp. 5604\u20135613","DOI":"10.1109\/CVPR.2019.00575"},{"key":"1205_CR8","doi-asserted-by":"crossref","unstructured":"Z. Zhang, Z. Cui, C. Xu, Z. Jie, X. Li, and J. Yang, \"Joint task-recursive learning for semantic segmentation and depth estimation,\" in Proceedings of the European Conference on Computer Vision (ECCV), 2018, pp. 235\u2013251","DOI":"10.1007\/978-3-030-01249-6_15"},{"key":"1205_CR9","doi-asserted-by":"crossref","unstructured":"A. Kendall, M. Grimes, and R. Cipolla, \"PoseNet: A Convolutional Network for Real-Time 6-DOF Camera Relocalization,\" in 2015 IEEE International Conference on Computer Vision (ICCV), 2015, pp. 2938\u20132946","DOI":"10.1109\/ICCV.2015.336"},{"key":"1205_CR10","unstructured":"D. Eigen, C. Puhrsch, and R. Fergus, \"Depth Map Prediction from a Single Image using a Multi-Scale Deep Network,\" in Advances in Neural Information Processing Systems 27, 2014, pp. 2366--2374: Curran Associates, Inc."},{"key":"1205_CR11","doi-asserted-by":"crossref","unstructured":"D. Eigen and R. Fergus, \"Predicting Depth, Surface Normals and Semantic Labels with a Common Multi-Scale Convolutional Architecture,\" The IEEE International Conference on Computer Vision (ICCV), 2015","DOI":"10.1109\/ICCV.2015.304"},{"key":"1205_CR12","doi-asserted-by":"crossref","unstructured":"S. Wang, R. Clark, H. Wen, and N. Trigoni, \"DeepVO: Towards End-to-End Visual Odometry with Deep Recurrent Convolutional Neural Networks,\" in 2017 IEEE International Conference on Robotics and Automation (ICRA), 2017, Pp. 2043-2050","DOI":"10.1109\/ICRA.2017.7989236"},{"key":"1205_CR13","doi-asserted-by":"crossref","unstructured":"B. Ummenhofer et al., \"DeMoN: Depth and Motion Network for Learning Monocular Stereo,\" in The IEEE Conference on Computer Vision and Pattern Recognition (CVPR), 2017","DOI":"10.1109\/CVPR.2017.596"},{"key":"1205_CR14","doi-asserted-by":"crossref","unstructured":"C. Godard, O. Mac Aodha, and G. J. Brostow, \"Unsupervised monocular depth estimation with left-right consistency,\" in Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 2017, pp. 270\u2013279","DOI":"10.1109\/CVPR.2017.699"},{"key":"1205_CR15","doi-asserted-by":"crossref","unstructured":"T. Zhou, M. Brown, N. Snavely, and D. G. Lowe, \"Unsupervised Learning of Depth and Ego-Motion from Video,\" in 2017 IEEE Conference on Computer Vision and Pattern Recognition (CVPR), 2017, pp. 6612\u20136619","DOI":"10.1109\/CVPR.2017.700"},{"key":"1205_CR16","doi-asserted-by":"crossref","unstructured":"Z. Yin and J. Shi, \"GeoNet: Unsupervised Learning of Dense Depth, Optical Flow and Camera Pose,\" in 2018 IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 2018, Pp. 1983-1992","DOI":"10.1109\/CVPR.2018.00212"},{"key":"1205_CR17","doi-asserted-by":"crossref","unstructured":"A. Wong and S. Soatto, \"Bilateral cyclic constraint and adaptive regularization for unsupervised monocular depth prediction,\" in Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 2019, pp. 5644\u20135653","DOI":"10.1109\/CVPR.2019.00579"},{"key":"1205_CR18","doi-asserted-by":"crossref","unstructured":"Y. Kuznietsov, J. Stuckler, and B. Leibe, \"Semi-Supervised Deep Learning for Monocular Depth Map Prediction,\" in 2017 IEEE Conference on Computer Vision and Pattern Recognition (CVPR), 2017, pp. 2215\u20132223","DOI":"10.1109\/CVPR.2017.238"},{"key":"1205_CR19","doi-asserted-by":"crossref","unstructured":"N. Yang, R. Wang, J. St\u00fcckler, and D. Cremers, \"Deep Virtual Stereo Odometry: Leveraging Deep Depth Prediction for Monocular Direct Sparse Odometry: 15th European Conference, Munich, Germany, September 8\u201314, 2018, Proceedings, Part VIII,\" 2018, pp. 835\u2013852","DOI":"10.1007\/978-3-030-01237-3_50"},{"key":"1205_CR20","volume-title":"Real-Time Joint Semantic Segmentation and Depth Estimation Using Asymmetric Annotations","author":"V Nekrasov","year":"2018","unstructured":"V. Nekrasov, T. Dharmasiri, A. Spek, T. Drummond, and I. Reid, \"Real-Time Joint Semantic Segmentation and Depth Estimation Using Asymmetric Annotations,\" 2018"},{"key":"1205_CR21","doi-asserted-by":"crossref","unstructured":"A. Atapour-Abarghouei and T. P. Breckon, \"Veritatem dies aperit-temporally consistent depth prediction enabled by a multi-task geometric and semantic scene understanding approach,\" in Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 2019, pp. 3373\u20133384","DOI":"10.1109\/CVPR.2019.00349"},{"key":"1205_CR22","doi-asserted-by":"crossref","unstructured":"P. Z. Ramirez, M. Poggi, F. Tosi, S. Mattoccia, and L. Di Stefano, \"Geometry meets semantics for semi-supervised monocular depth estimation,\" in Asian Conference on Computer Vision, 2018, pp. 298\u2013313: Springer","DOI":"10.1007\/978-3-030-20893-6_19"},{"key":"1205_CR23","doi-asserted-by":"crossref","unstructured":"A. Atapour-Abarghouei and T. P. Breckon, \"Monocular segment-wise depth: Monocular depth estimation based on a semantic segmentation prior,\" in 2019 IEEE International Conference on Image Processing (ICIP), 2019, pp. 4295\u20134299: IEEE","DOI":"10.1109\/ICIP.2019.8803551"},{"key":"1205_CR24","doi-asserted-by":"crossref","unstructured":"G. Ros, L. Sellart, J. Materzynska, D. Vazquez, and A. M. Lopez, \"The synthia dataset: A large collection of synthetic images for semantic segmentation of urban scenes,\" in Proceedings of the IEEE conference on computer vision and pattern recognition, 2016, pp. 3234\u20133243","DOI":"10.1109\/CVPR.2016.352"},{"key":"1205_CR25","doi-asserted-by":"crossref","unstructured":"G. Lin, A. Milan, C. Shen, and I. Reid, RefineNet: Multi-path Refinement Networks for High-Resolution Semantic Segmentation. 2017, pp. 5168\u20135177","DOI":"10.1109\/CVPR.2017.549"},{"key":"1205_CR26","doi-asserted-by":"publisher","first-page":"05\/01","DOI":"10.1117\/12.524839","volume":"5291","author":"C Fehn","year":"2004","unstructured":"Fehn, C.: Depth-image-based rendering (DIBR), compression and transmission for a new approach on 3D-TV. Proc. SPIE. 5291, 05\/01 (2004)","journal-title":"Proc. SPIE"},{"key":"1205_CR27","doi-asserted-by":"crossref","unstructured":"T. Zhou, S. Tulsiani, W. Sun, J. Malik, and A. Efros, View Synthesis by Appearance Flow. 2016, pp. 286\u2013301","DOI":"10.1007\/978-3-319-46493-0_18"},{"key":"1205_CR28","doi-asserted-by":"crossref","unstructured":"Z. Wang, A. Bovik, H. R. Sheikh, and E. Simoncelli, \"Image quality assessment: From error visibility to structural similarity,\" IEEE Trans. Image Process., vol. 13, pp. 600\u2013612, 01\/01 2014","DOI":"10.1109\/TIP.2003.819861"},{"key":"1205_CR29","unstructured":"M. Abadi et al., TensorFlow : Large-Scale Machine Learning on Heterogeneous Distributed Systems. 2015"},{"key":"1205_CR30","doi-asserted-by":"crossref","unstructured":"A. Geiger, Are we ready for autonomous driving? The KITTI vision benchmark suite. 2012, pp. 3354\u20133361","DOI":"10.1109\/CVPR.2012.6248074"},{"key":"1205_CR31","unstructured":".D. Kingma and J. Ba, \"Adam: a Method for Stochastic Optimization,\" International Conference on Learning Representations, 12\/22 2014"},{"key":"1205_CR32","first-page":"02\/25","volume":"38","author":"F Liu","year":"2015","unstructured":"Liu, F., Shen, C., Lin, G., Reid, I.: Learning depth from single monocular images using deep convolutional neural fields. IEEE Trans. Pattern Anal. Mach. Intell. 38, 02\/25 (2015)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"1205_CR33","doi-asserted-by":"crossref","unstructured":"R. Garg, V. K. B G, G. Carneiro, and I. Reid, Unsupervised CNN for Single View Depth Estimation: Geometry to the Rescue. 2016","DOI":"10.1007\/978-3-319-46484-8_45"}],"container-title":["Journal of Intelligent &amp; Robotic Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10846-020-01205-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10846-020-01205-0\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10846-020-01205-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2021,5,29]],"date-time":"2021-05-29T23:24:38Z","timestamp":1622330678000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10846-020-01205-0"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2020,5,30]]},"references-count":33,"journal-issue":{"issue":"2","published-print":{"date-parts":[[2020,11]]}},"alternative-id":["1205"],"URL":"https:\/\/doi.org\/10.1007\/s10846-020-01205-0","relation":{},"ISSN":["0921-0296","1573-0409"],"issn-type":[{"value":"0921-0296","type":"print"},{"value":"1573-0409","type":"electronic"}],"subject":[],"published":{"date-parts":[[2020,5,30]]},"assertion":[{"value":"12 December 2019","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"16 April 2020","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"30 May 2020","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}}]}}