{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,9]],"date-time":"2026-05-09T17:03:45Z","timestamp":1778346225384,"version":"3.51.4"},"reference-count":44,"publisher":"Springer Science and Business Media LLC","issue":"37","license":[{"start":{"date-parts":[[2024,4,27]],"date-time":"2024-04-27T00:00:00Z","timestamp":1714176000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,4,27]],"date-time":"2024-04-27T00:00:00Z","timestamp":1714176000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100015401","name":"Key Research and Development Projects of Shaanxi Province","doi-asserted-by":"publisher","award":["2020NY-144"],"award-info":[{"award-number":["2020NY-144"]}],"id":[{"id":"10.13039\/501100015401","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimed Tools Appl"],"DOI":"10.1007\/s11042-024-19235-3","type":"journal-article","created":{"date-parts":[[2024,4,27]],"date-time":"2024-04-27T07:03:00Z","timestamp":1714201380000},"page":"84619-84637","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":3,"title":["Single image depth estimation using improved U-Net and edge-guide loss"],"prefix":"10.1007","volume":"83","author":[{"given":"Mengfei","family":"He","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yan","family":"Gao","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yan","family":"Long","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2024,4,27]]},"reference":[{"key":"19235_CR1","doi-asserted-by":"crossref","unstructured":"Huang C-H, Tsung W-N, Yang W-J, Chen C-H (2019) Unsupervised monocular depth estimation for autonomous driving. In: Proceedings of the international display workshops (IDV), pp 128\u2013131","DOI":"10.36463\/idw.2019.0128"},{"key":"19235_CR2","doi-asserted-by":"publisher","first-page":"620","DOI":"10.1016\/j.compeleceng.2016.04.018","volume":"67","author":"C Lai","year":"2018","unstructured":"Lai C, Su K (2018) Development of an intelligent mobile robot localization system using Kinect RGB-D mapping and neural network. Comput Electr Eng 67:620\u2013628","journal-title":"Comput Electr Eng"},{"issue":"9","key":"19235_CR3","doi-asserted-by":"publisher","first-page":"2485","DOI":"10.1167\/jov.21.9.2485","volume":"21","author":"J Lee","year":"2021","unstructured":"Lee J, Joo S (2021) Three-dimensional depth estimation of virtual objects in augmented reality. J Vision 21(9):2485. https:\/\/doi.org\/10.1167\/jov.21.9.2485","journal-title":"J Vision"},{"key":"19235_CR4","doi-asserted-by":"publisher","unstructured":"Smisek J, Jancosek M, Pajdla T (2011) 3D with kinect. In: IEEE international conference on computer vision workshops (ICCV Workshops), pp 1154\u20131160. https:\/\/doi.org\/10.1109\/ICCVW.2011.6130380","DOI":"10.1109\/ICCVW.2011.6130380"},{"issue":"6","key":"19235_CR5","doi-asserted-by":"publisher","first-page":"44","DOI":"10.1093\/jof\/98.6.44","volume":"98","author":"RO Dubayah","year":"2000","unstructured":"Dubayah RO, Drake JB (2000) Lidar remote sensing for forestry. J Forest 98(6):44\u201346","journal-title":"J Forest"},{"key":"19235_CR6","doi-asserted-by":"publisher","unstructured":"Godard C, Aodha O, Firman M, Brostow G (2019) Digging into self-supervised monocular depth estimation. In: Proceedings of the IEEE international conference on computer vision (ICCV), pp 3827\u20133837. https:\/\/doi.org\/10.1109\/ICCV.2019.00393","DOI":"10.1109\/ICCV.2019.00393"},{"key":"19235_CR7","doi-asserted-by":"publisher","first-page":"1862","DOI":"10.4028\/www.scientific.net\/AMM.284-287.1862","volume":"284\u2013287","author":"KY Chen","year":"2013","unstructured":"Chen KY, Chien CC, Tseng CT (2013) Improving the accuracy of depth estimation in binocular vision for robotic applications. Appl Mech Mater 284\u2013287:1862\u20131866. https:\/\/doi.org\/10.4028\/www.scientific.net\/AMM.284-287.1862","journal-title":"Appl Mech Mater"},{"issue":"1","key":"19235_CR8","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1167\/9.1.10","volume":"9","author":"RS Allison","year":"2009","unstructured":"Allison RS, Gillam BJ, Vecellio E (2009) Binocular depth discrimination and estimation beyond interaction space. J Vision 9(1):1\u201314. https:\/\/doi.org\/10.1167\/9.1.10","journal-title":"J Vision"},{"key":"19235_CR9","doi-asserted-by":"publisher","unstructured":"Zhou T, Brown M, Snavely N, Lowe D (2017) Unsupervised learning of depth and ego-motion from video. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR), pp 6612\u20136621. https:\/\/doi.org\/10.1109\/CVPR.2017.700","DOI":"10.1109\/CVPR.2017.700"},{"key":"19235_CR10","doi-asserted-by":"publisher","unstructured":"Wang C, Buenaposada JM, Rui Z, Lucey S (2018) Learning depth from monocular videos using direct methods. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition (CVPR), pp 2022\u20132030. https:\/\/doi.org\/10.1109\/CVPR.2018.00216","DOI":"10.1109\/CVPR.2018.00216"},{"key":"19235_CR11","doi-asserted-by":"publisher","unstructured":"Ranjan A, Jampani V, Balles L, Kim K, Sun D, Wulff J, Black MJ (2019) Competitive collaboration: joint unsupervised learning of depth, camera motion, optical flow and motion segmentation. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition (CVPR), pp 12232\u201312241. https:\/\/doi.org\/10.1109\/CVPR.2019.01252","DOI":"10.1109\/CVPR.2019.01252"},{"key":"19235_CR12","first-page":"2366","volume":"3","author":"D Eigen","year":"2014","unstructured":"Eigen D, Puhrsch C, Fergus R (2014) Depth map prediction from a single image using a multi-scale deep network. Adv Neural Inf Process Syst 3:2366\u20132374","journal-title":"Adv Neural Inf Process Syst"},{"key":"19235_CR13","doi-asserted-by":"publisher","unstructured":"Eigen D, Fergus R (2015) Predicting depth, surface normals and semantic labels with a common multi-scale convolutional architecture. In: Proceedings of the IEEE international conference on computer vision (ICCV), pp 2650\u20132658. https:\/\/doi.org\/10.1109\/ICCV.2015.304","DOI":"10.1109\/ICCV.2015.304"},{"key":"19235_CR14","unstructured":"Alhashim I, Wonka P (2018) High quality monocular depth estimation via transfer learning. arXiv:1812.11941"},{"issue":"4","key":"19235_CR15","doi-asserted-by":"publisher","first-page":"1738","DOI":"10.1109\/TPAMI.2020.3032602","volume":"44","author":"H Laga","year":"2022","unstructured":"Laga H, Jospin L, Boussaid F, Bennamoun M (2022) A survey on deep learning techniques for stereo-based depth estimation. IEEE T Pattern Anal 44(4):1738\u20131764","journal-title":"IEEE T Pattern Anal"},{"issue":"1","key":"19235_CR16","doi-asserted-by":"publisher","first-page":"7","DOI":"10.1007\/BF00128525","volume":"1","author":"RC Bolles","year":"1987","unstructured":"Bolles RC, Baker HH, Marimont DH (1987) Epipolar-plane image analysis: An approach to determining structure from motion. Int J Comput Vis 1(1):7\u201355","journal-title":"Int J Comput Vis"},{"issue":"1\u20132","key":"19235_CR17","doi-asserted-by":"publisher","first-page":"97","DOI":"10.1007\/s11263-005-3844-1","volume":"65","author":"E Prados","year":"2005","unstructured":"Prados E, Faugeras O (2005) A generic and provably convergent shape-from-shading method for orthographic and pinhole cameras. Int J Comput Vision 65(1\u20132):97\u2013125","journal-title":"Int J Comput Vision"},{"issue":"8","key":"19235_CR18","doi-asserted-by":"publisher","first-page":"824","DOI":"10.1109\/34.308479","volume":"16","author":"SK Nayar","year":"1994","unstructured":"Nayar SK, Nakagawa Y (1994) Shape from focus. IEEE T Pattern Anal 16(8):824\u2013831","journal-title":"IEEE T Pattern Anal"},{"issue":"3","key":"19235_CR19","doi-asserted-by":"publisher","first-page":"406","DOI":"10.1109\/TPAMI.2005.43","volume":"27","author":"F Paolo","year":"2005","unstructured":"Paolo F, Stefano S (2005) A geometric approach to shape from defocus. IEEE T Pattern Anal 27(3):406\u2013417","journal-title":"IEEE T Pattern Anal"},{"key":"19235_CR20","doi-asserted-by":"publisher","unstructured":"Huang G, Liu Z, Laurens V, Weinberger KQ (2017) Densely connected convolutional networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR), pp 2261\u20132269. https:\/\/doi.org\/10.1109\/CVPR.2017.243","DOI":"10.1109\/CVPR.2017.243"},{"key":"19235_CR21","doi-asserted-by":"publisher","unstructured":"Hao Z, Li Y, You S, Lu F (2018) Detail preserving depth estimation from a single image using attention guided networks. In: Proceedings of the international conference on 3D vision (3DV), pp 304\u2013313. https:\/\/doi.org\/10.1109\/3DV.2018.00043","DOI":"10.1109\/3DV.2018.00043"},{"key":"19235_CR22","doi-asserted-by":"publisher","unstructured":"Lee J, Kim C (2019) Monocular depth estimation using relative depth maps. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition (CVPR), pp 9721\u20139730. https:\/\/doi.org\/10.1109\/CVPR.2019.00996","DOI":"10.1109\/CVPR.2019.00996"},{"key":"19235_CR23","doi-asserted-by":"publisher","unstructured":"Xue F, Cao J, Zhou Y, Sheng F, Wang Y, Ming A (2021) Boundary-induced and scene-aggregated network for monocular depth prediction. Pattern Recogn 115. https:\/\/doi.org\/10.1016\/j.patcog.2021.107901","DOI":"10.1016\/j.patcog.2021.107901"},{"key":"19235_CR24","doi-asserted-by":"publisher","unstructured":"Laina I, Rupprecht C, Belagiannis V, Tombari F, Navab N (2016) Deeper depth prediction with fully convolutional residual networks. In: Proceedings of the international conference on 3D vision (3DV), pp 239\u2013248. https:\/\/doi.org\/10.1109\/3DV.2016.32","DOI":"10.1109\/3DV.2016.32"},{"key":"19235_CR25","doi-asserted-by":"publisher","unstructured":"Wang L, Zhang J, Wang O, Lin Z, Lu H (2020) SDC-depth: semantic divide-and-conquer network for monocular depth estimation. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition (CVPR), pp 538\u2013547. https:\/\/doi.org\/10.1109\/CVPR42600.2020.00062","DOI":"10.1109\/CVPR42600.2020.00062"},{"key":"19235_CR26","doi-asserted-by":"crossref","unstructured":"Lyu X, Liu L, Wang M, Kong X, Liu L, Liu Y, Chen X, Yuan Y (2021) HR-depth: high resolution self-supervised monocular depth estimation. In: 35th AAAI conference on artificial intelligence (AAAI), pp 2294\u20132301","DOI":"10.1609\/aaai.v35i3.16329"},{"key":"19235_CR27","doi-asserted-by":"publisher","unstructured":"Li B, Shen C, Dai Y, Hengel AVD, He M (2015) Depth and surface normal estimation from monocular images using regression on deep features and hierarchical CRFs. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition (CVPR), pp 1119\u20131127. https:\/\/doi.org\/10.1109\/CVPR.2015.7298715","DOI":"10.1109\/CVPR.2015.7298715"},{"key":"19235_CR28","doi-asserted-by":"publisher","unstructured":"He K, Zhang X, Ren S, Sun J (2016) Deep residual learning for image recognition. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR), pp 770\u2013778. https:\/\/doi.org\/10.1109\/CVPR.2016.90","DOI":"10.1109\/CVPR.2016.90"},{"key":"19235_CR29","doi-asserted-by":"publisher","unstructured":"Hu J, Ozay M, Zhang Y, Okatani T (2019) Revisiting single image depth estimation: toward higher resolution maps with accurate object boundaries. In: Proceedings of the IEEE winter conference on applications of computer vision (WACV), pp 1043\u20131051. https:\/\/doi.org\/10.1109\/WACV.2019.00116","DOI":"10.1109\/WACV.2019.00116"},{"key":"19235_CR30","unstructured":"Chen W, Fu Z, Yang D, Deng J (2016) Single-image depth perception in the wild. In: Proceedings of the annual conference on neural information processing systems (NIPS), pp 730\u2013738"},{"key":"19235_CR31","doi-asserted-by":"publisher","unstructured":"Xian K, Zhang J, Wang O, Mai L, Lin Z, Cao Z (2020) Structure-guided ranking loss for single image depth prediction. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition (CVPR), pp 608\u2013617. https:\/\/doi.org\/10.1109\/CVPR42600.2020.00069","DOI":"10.1109\/CVPR42600.2020.00069"},{"key":"19235_CR32","doi-asserted-by":"publisher","unstructured":"Zeiler M, Taylor GW, Fergus R (2011) Adaptive deconvolutional networks for mid and high level feature learning. In: Proceedings of the IEEE international conference on computer vision (ICCV), pp 2018\u20132025. https:\/\/doi.org\/10.1109\/ICCV.2011.6126474","DOI":"10.1109\/ICCV.2011.6126474"},{"key":"19235_CR33","doi-asserted-by":"publisher","first-page":"818","DOI":"10.1007\/978-3-319-10590-1_53","volume":"8689","author":"M Zeiler","year":"2014","unstructured":"Zeiler M, Fergus R (2014) Visualizing and understanding convolutional networks. Lect Notes Comput Sci 8689:818\u2013833","journal-title":"Lect Notes Comput Sci"},{"issue":"4","key":"19235_CR34","doi-asserted-by":"publisher","first-page":"640","DOI":"10.1109\/TPAMI.2016.2572683","volume":"39","author":"E Shelhamer","year":"2017","unstructured":"Shelhamer E, Long J, Darrell T (2017) Fully convolutional networks for semantic segmentation. IEEE T Pattern Anal 39(4):640\u2013651","journal-title":"IEEE T Pattern Anal"},{"key":"19235_CR35","unstructured":"Yu F, Koltun V (2016) Multi-scale context aggregation by dilated convolutions. In: Proceedings of the international conference on learning representations (ICLR)"},{"key":"19235_CR36","doi-asserted-by":"publisher","unstructured":"Wang P, Chen P, Yuan Y, Liu D, Huang Z, Hou X, Cottrell G (2018) Understanding convolution for semantic segmentation. In: Proceedings of the IEEE winter conference on applications of computer vision (WACV), pp 1451\u20131460. https:\/\/doi.org\/10.1109\/WACV.2018.00163","DOI":"10.1109\/WACV.2018.00163"},{"key":"19235_CR37","doi-asserted-by":"publisher","unstructured":"Zhu S, Brazil G, Liu X (2020) The edge of depth: explicit constraints between segmentation and depth. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition (CVPR), pp 13113\u201313122. https:\/\/doi.org\/10.1109\/CVPR42600.2020.01313","DOI":"10.1109\/CVPR42600.2020.01313"},{"issue":"4","key":"19235_CR38","doi-asserted-by":"publisher","first-page":"600","DOI":"10.1109\/TIP.2003.819861","volume":"13","author":"Z Wang","year":"2004","unstructured":"Wang Z, Bovik A, Sheikh H, Simoncelli E (2004) Image quality assessment: from error visibility to structural similarity. IEEE T Image Process 13(4):600\u2013612","journal-title":"IEEE T Image Process"},{"key":"19235_CR39","doi-asserted-by":"publisher","unstructured":"Godard C, Aodha OM, Brostow GJ (2017) Unsupervised monocular depth estimation with left-right consistency. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR), pp 6602\u20136611. https:\/\/doi.org\/10.1109\/CVPR.2017.699","DOI":"10.1109\/CVPR.2017.699"},{"issue":"6","key":"19235_CR40","doi-asserted-by":"publisher","first-page":"679","DOI":"10.1109\/TPAMI.1986.4767851","volume":"8","author":"J Canny","year":"1986","unstructured":"Canny J (1986) A computational approach to edge detection. IEEE T Pattern Anal 8(6):679\u2013698","journal-title":"IEEE T Pattern Anal"},{"key":"19235_CR41","doi-asserted-by":"publisher","unstructured":"Silberman N, Hoiem D, Kohli P, Fergus R (2012) Indoor segmentation and support inference from RGBD images. In: Proceedings of the european conference on computer vision (ECCV), pp 746\u2013760. https:\/\/doi.org\/10.1007\/978-3-642-33715-4_54","DOI":"10.1007\/978-3-642-33715-4_54"},{"key":"19235_CR42","doi-asserted-by":"publisher","first-page":"689","DOI":"10.1145\/1015706.1015780","volume":"23","author":"A Levin","year":"2004","unstructured":"Levin A, Lischinski D, Weiss Y (2004) Colorization using optimization. Acm T Graphic 23:689\u2013694","journal-title":"Acm T Graphic"},{"key":"19235_CR43","unstructured":"Kingma D, Ba J (2014) Adam: a method for stochastic optimization. In: Proceedings of the international conference on learning representations (ICLR)"},{"key":"19235_CR44","doi-asserted-by":"publisher","unstructured":"Jia D, Wei D, Socher R, Li L, Kai L, Li F (2009) ImageNet: a large-scale hierarchical image database. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR), pp 248\u2013255. https:\/\/doi.org\/10.1109\/CVPR.2009.5206848","DOI":"10.1109\/CVPR.2009.5206848"}],"container-title":["Multimedia Tools and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-024-19235-3.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11042-024-19235-3\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-024-19235-3.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,11,15]],"date-time":"2024-11-15T08:09:08Z","timestamp":1731658148000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11042-024-19235-3"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,4,27]]},"references-count":44,"journal-issue":{"issue":"37","published-online":{"date-parts":[[2024,11]]}},"alternative-id":["19235"],"URL":"https:\/\/doi.org\/10.1007\/s11042-024-19235-3","relation":{},"ISSN":["1573-7721"],"issn-type":[{"value":"1573-7721","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,4,27]]},"assertion":[{"value":"21 October 2022","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"21 February 2024","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"9 April 2024","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"27 April 2024","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors have no relevant financial or non-financial interests to disclose.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}]}}