{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,11]],"date-time":"2026-07-11T11:12:07Z","timestamp":1783768327539,"version":"3.55.0"},"reference-count":64,"publisher":"Springer Science and Business Media LLC","issue":"4","license":[{"start":{"date-parts":[[2026,5,6]],"date-time":"2026-05-06T00:00:00Z","timestamp":1778025600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,5,6]],"date-time":"2026-05-06T00:00:00Z","timestamp":1778025600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"the National Natural Science Foundation of China (NSFC) under Grants","award":["62102003"],"award-info":[{"award-number":["62102003"]}]},{"name":"the Anhui Postdoctoral Science Foundation under Grant","award":["2022B623"],"award-info":[{"award-number":["2022B623"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimedia Systems"],"published-print":{"date-parts":[[2026,6]]},"DOI":"10.1007\/s00530-026-02356-0","type":"journal-article","created":{"date-parts":[[2026,5,6]],"date-time":"2026-05-06T06:00:29Z","timestamp":1778047229000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["SharpDepth: self-supervised monocular depth estimation through edge awareness and wavelet frequency domain fusion"],"prefix":"10.1007","volume":"32","author":[{"given":"Bin","family":"Ge","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Chao","family":"Yang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Chenxing","family":"Xia","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Mengyan","family":"Cheng","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Rui","family":"Sun","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,5,6]]},"reference":[{"issue":"3","key":"2356_CR1","doi-asserted-by":"publisher","first-page":"246","DOI":"10.1007\/s00530-025-01819-0","volume":"31","author":"Y Shi","year":"2025","unstructured":"Shi, Y., Zhou, S., Wang, W., Lu, X.: Depth-free view synthesis from diffusion models for monocular 3d detector in autonomous driving. Multimedia Syst. 31(3), 246 (2025)","journal-title":"Multimedia Syst."},{"issue":"5","key":"2356_CR2","doi-asserted-by":"publisher","first-page":"2114","DOI":"10.1109\/TVCG.2022.3150486","volume":"28","author":"W Mehringer","year":"2022","unstructured":"Mehringer, W., Wirth, M., Roth, D., Michelson, G., Eskofier, B.M.: Stereopsis only: Validation of a monocular depth cues reduced gamified virtual reality with reaction time measurement. IEEE Trans. Visual Comput. Graphics 28(5), 2114\u20132124 (2022)","journal-title":"IEEE Trans. Visual Comput. Graphics"},{"key":"2356_CR3","doi-asserted-by":"crossref","unstructured":"Chen, Y., Inaltekin, H., Gorlatova, M.: Adaptslam: Edge-assisted adaptive slam with resource constraints via uncertainty minimization. In: IEEE INFOCOM 2023-IEEE Conference on Computer Communications, pp. 1\u201310 (2023). IEEE","DOI":"10.1109\/INFOCOM53939.2023.10229009"},{"issue":"4","key":"2356_CR4","doi-asserted-by":"publisher","first-page":"2224","DOI":"10.1109\/TPAMI.2023.3332407","volume":"46","author":"J Bae","year":"2023","unstructured":"Bae, J., Hwang, K., Im, S.: A study on the generality of neural network structures for monocular depth estimation. IEEE Trans. Pattern Anal. Mach. Intell. 46(4), 2224\u20132238 (2023)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"2356_CR5","doi-asserted-by":"crossref","unstructured":"Fang, J., Liu, G.: Semantic and optical flow guided self-supervised monocular depth and ego-motion estimation. In: International Conference on Image and Graphics, pp. 465\u2013477 (2021). Springer","DOI":"10.1007\/978-3-030-87361-5_38"},{"key":"2356_CR6","doi-asserted-by":"crossref","unstructured":"Godard, C., Mac\u00a0Aodha, O., Firman, M., Brostow, G.J.: Digging into self-supervised monocular depth estimation. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 3828\u20133838 (2019)","DOI":"10.1109\/ICCV.2019.00393"},{"issue":"6","key":"2356_CR7","doi-asserted-by":"publisher","first-page":"4989","DOI":"10.1109\/TCSVT.2023.3340948","volume":"34","author":"G Wu","year":"2023","unstructured":"Wu, G., Liu, H., Wang, L., Li, K., Guo, Y., Chen, Z.: Self-supervised multi-frame monocular depth estimation for dynamic scenes. IEEE Trans. Circuits Syst. Video Technol. 34(6), 4989\u20135001 (2023)","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"issue":"2","key":"2356_CR8","doi-asserted-by":"publisher","first-page":"111","DOI":"10.1007\/s00530-025-01700-0","volume":"31","author":"Z Lu","year":"2025","unstructured":"Lu, Z., Chen, Y.: Self-supervised monocular depth estimation via multiple bilateral consistency. Multimedia Syst. 31(2), 111 (2025)","journal-title":"Multimedia Syst."},{"key":"2356_CR9","doi-asserted-by":"crossref","unstructured":"Watson, J., Mac\u00a0Aodha, O., Prisacariu, V., Brostow, G., Firman, M.: The temporal opportunist: Self-supervised multi-frame monocular depth. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 1164\u20131174 (2021)","DOI":"10.1109\/CVPR46437.2021.00122"},{"key":"2356_CR10","doi-asserted-by":"crossref","unstructured":"Lyu, X., Liu, L., Wang, M., Kong, X., Liu, L., Liu, Y., Chen, X., Yuan, Y.: Hr-depth: High resolution self-supervised monocular depth estimation. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 35, pp. 2294\u20132301 (2021)","DOI":"10.1609\/aaai.v35i3.16329"},{"key":"2356_CR11","doi-asserted-by":"crossref","unstructured":"Zhang, N., Nex, F., Vosselman, G., Kerle, N.: Lite-mono: A lightweight cnn and transformer architecture for self-supervised monocular depth estimation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 18537\u201318546 (2023)","DOI":"10.1109\/CVPR52729.2023.01778"},{"key":"2356_CR12","unstructured":"Eigen, D., Puhrsch, C., Fergus, R.: Depth map prediction from a single image using a multi-scale deep network. Advances in neural information processing systems 27 (2014)"},{"issue":"12","key":"2356_CR13","doi-asserted-by":"publisher","first-page":"8883","DOI":"10.1109\/TPAMI.2024.3411571","volume":"46","author":"S Shao","year":"2024","unstructured":"Shao, S., Pei, Z., Chen, W., Chen, P.C., Li, Z.: Nddepth: Normal-distance assisted monocular depth estimation and completion. IEEE Trans. Pattern Anal. Mach. Intell. 46(12), 8883\u20138899 (2024)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"2356_CR14","doi-asserted-by":"publisher","DOI":"10.1016\/j.patcog.2023.109982","volume":"145","author":"S Tang","year":"2024","unstructured":"Tang, S., Lu, T., Liu, X., Zhou, H., Zhang, Y.: Catnet: Convolutional attention and transformer for monocular depth estimation. Pattern Recogn. 145, 109982 (2024)","journal-title":"Pattern Recogn."},{"key":"2356_CR15","doi-asserted-by":"crossref","unstructured":"Zhou, T., Brown, M., Snavely, N., Lowe, D.G.: Unsupervised learning of depth and ego-motion from video. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 1851\u20131858 (2017)","DOI":"10.1109\/CVPR.2017.700"},{"key":"2356_CR16","doi-asserted-by":"crossref","unstructured":"Gordon, A., Li, H., Jonschkowski, R., Angelova, A.: Depth from videos in the wild: Unsupervised monocular depth learning from unknown cameras. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 8977\u20138986 (2019)","DOI":"10.1109\/ICCV.2019.00907"},{"key":"2356_CR17","doi-asserted-by":"crossref","unstructured":"Zhan, H., Garg, R., Weerasekera, C.S., Li, K., Agarwal, H., Reid, I.: Unsupervised learning of monocular depth estimation and visual odometry with deep feature reconstruction. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 340\u2013349 (2018)","DOI":"10.1109\/CVPR.2018.00043"},{"key":"2356_CR18","doi-asserted-by":"crossref","unstructured":"Zhang, R., Qiu, H., Wang, T., Guo, Z., Cui, Z., Qiao, Y., Li, H., Gao, P.: Monodetr: Depth-guided transformer for monocular 3d object detection. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 9155\u20139166 (2023)","DOI":"10.1109\/ICCV51070.2023.00840"},{"key":"2356_CR19","doi-asserted-by":"crossref","unstructured":"Chen, Y., Schmid, C., Sminchisescu, C.: Self-supervised learning with geometric constraints in monocular video: Connecting flow, depth, and camera. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 7063\u20137072 (2019)","DOI":"10.1109\/ICCV.2019.00716"},{"key":"2356_CR20","doi-asserted-by":"crossref","unstructured":"Klodt, M., Vedaldi, A.: Supervising the new with the old: learning sfm from sfm. In: Proceedings of the European Conference on Computer Vision (ECCV), pp. 698\u2013713 (2018)","DOI":"10.1007\/978-3-030-01249-6_43"},{"key":"2356_CR21","doi-asserted-by":"crossref","unstructured":"Han, W., Yin, J., Shen, J.: Self-supervised monocular depth estimation by direction-aware cumulative convolution network. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 8613\u20138623 (2023)","DOI":"10.1109\/ICCV51070.2023.00791"},{"key":"2356_CR22","doi-asserted-by":"crossref","unstructured":"He, M., Hui, L., Bian, Y., Ren, J., Xie, J., Yang, J.: Ra-depth: Resolution adaptive self-supervised monocular depth estimation. In: European Conference on Computer Vision, pp. 565\u2013581 (2022). Springer","DOI":"10.1007\/978-3-031-19812-0_33"},{"issue":"1","key":"2356_CR23","doi-asserted-by":"publisher","first-page":"497","DOI":"10.1109\/TPAMI.2023.3322549","volume":"46","author":"L Sun","year":"2023","unstructured":"Sun, L., Bian, J.-W., Zhan, H., Yin, W., Reid, I., Shen, C.: Sc-depthv3: Robust self-supervised monocular depth estimation for dynamic scenes. IEEE Trans. Pattern Anal. Mach. Intell. 46(1), 497\u2013508 (2023)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"2356_CR24","doi-asserted-by":"crossref","unstructured":"Jin, Z., Qiu, Y., Zhang, K., Li, H., Luo, W.: Mb-taylorformer v2: Improved multi-branch linear transformer expanded by taylor formula for image restoration. IEEE Transactions on Pattern Analysis and Machine Intelligence (2025)","DOI":"10.1109\/TPAMI.2025.3559891"},{"key":"2356_CR25","doi-asserted-by":"publisher","first-page":"7419","DOI":"10.1109\/TIP.2021.3104166","volume":"30","author":"K Zhang","year":"2021","unstructured":"Zhang, K., Li, R., Yu, Y., Luo, W., Li, C.: Deep dense multi-scale network for snow removal using semantic and depth priors. IEEE Trans. Image Process. 30, 7419\u20137431 (2021)","journal-title":"IEEE Trans. Image Process."},{"issue":"1","key":"2356_CR26","doi-asserted-by":"publisher","first-page":"291","DOI":"10.1109\/TIP.2018.2867733","volume":"28","author":"K Zhang","year":"2018","unstructured":"Zhang, K., Luo, W., Zhong, Y., Ma, L., Liu, W., Li, H.: Adversarial spatio-temporal learning for video deblurring. IEEE Trans. Image Process. 28(1), 291\u2013301 (2018)","journal-title":"IEEE Trans. Image Process."},{"key":"2356_CR27","doi-asserted-by":"crossref","unstructured":"Liu, Y., Dong, X., Lin, Y., Ye, M., Zhang, K., Du, B.: Condition-guided diffusion for multi-modal pedestrian trajectory prediction incorporating intention and interaction priors. IEEE Transactions on Pattern Analysis and Machine Intelligence (2025)","DOI":"10.1109\/TPAMI.2025.3645918"},{"key":"2356_CR28","doi-asserted-by":"publisher","DOI":"10.1016\/j.knosys.2023.111156","volume":"282","author":"Y Cui","year":"2023","unstructured":"Cui, Y., Knoll, A.: Exploring the potential of channel interactions for image restoration. Knowl.-Based Syst. 282, 111156 (2023)","journal-title":"Knowl.-Based Syst."},{"key":"2356_CR29","doi-asserted-by":"crossref","unstructured":"Cui, Y., Wang, Q., Li, C., Ren, W., Knoll, A.: Eenet: An effective and efficient network for single image dehazing. pattern recognition 158, 111074 (2025)","DOI":"10.1016\/j.patcog.2024.111074"},{"key":"2356_CR30","doi-asserted-by":"crossref","unstructured":"Cui, Y., Ren, W., Knoll, A.: Omni-kernel network for image restoration. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 38, pp. 1426\u20131434 (2024)","DOI":"10.1609\/aaai.v38i2.27907"},{"key":"2356_CR31","doi-asserted-by":"crossref","unstructured":"Cui, Y., Ren, W., Knoll, A.: Exploring the potential of pooling techniques for universal image restoration. IEEE Transactions on Image Processing (2025)","DOI":"10.1109\/TIP.2025.3572788"},{"key":"2356_CR32","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., Sun, J.: Deep residual learning for image recognition. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 770\u2013778 (2016)","DOI":"10.1109\/CVPR.2016.90"},{"key":"2356_CR33","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2023.127122","volume":"569","author":"Z Zhang","year":"2024","unstructured":"Zhang, Z., Chan, R.K., Wong, K.K.: Glocalfuse-depth: Fusing transformers and cnns for all-day self-supervised monocular depth estimation. Neurocomputing 569, 127122 (2024)","journal-title":"Neurocomputing"},{"key":"2356_CR34","doi-asserted-by":"publisher","DOI":"10.1016\/j.inffus.2024.102363","volume":"108","author":"X Wang","year":"2024","unstructured":"Wang, X., Luo, H., Wang, Z., Zheng, J., Bai, X.: Self-supervised multi-frame depth estimation with visual-inertial pose transformer and monocular guidance. Information Fusion 108, 102363 (2024)","journal-title":"Information Fusion"},{"issue":"3","key":"2356_CR35","doi-asserted-by":"publisher","first-page":"251","DOI":"10.1007\/s00530-025-01842-1","volume":"31","author":"L Ge","year":"2025","unstructured":"Ge, L., Zhang, C., Chen, Z., Lu, K., Feng, C.: Srsa-depth: shape and region similarity awareness for outdoor monocular depth estimation. Multimedia Syst. 31(3), 251 (2025)","journal-title":"Multimedia Syst."},{"key":"2356_CR36","doi-asserted-by":"publisher","DOI":"10.1016\/j.engappai.2025.110026","volume":"144","author":"X Qin","year":"2025","unstructured":"Qin, X., Wang, L., Zhu, Y., Mao, F., Zhang, X., He, C., Dong, Q.: Rectified self-supervised monocular depth estimation loss for nighttime and dynamic scenes. Eng. Appl. Artif. Intell. 144, 110026 (2025)","journal-title":"Eng. Appl. Artif. Intell."},{"key":"2356_CR37","doi-asserted-by":"crossref","unstructured":"Shao, S., Pei, Z., Chen, W., Sun, D., Chen, P.C., Li, Z.: Monodiffusion: Self-supervised monocular depth estimation using diffusion model. IEEE Transactions on Circuits and Systems for Video Technology (2024)","DOI":"10.1109\/TCSVT.2024.3509619"},{"issue":"9","key":"2356_CR38","doi-asserted-by":"publisher","first-page":"110","DOI":"10.1109\/MCOM.2013.6588659","volume":"51","author":"S-Y Wang","year":"2013","unstructured":"Wang, S.-Y., Chou, C.-L., Yang, C.-M.: Estinet openflow network simulator and emulator. IEEE Commun. Mag. 51(9), 110\u2013117 (2013)","journal-title":"IEEE Commun. Mag."},{"key":"2356_CR39","doi-asserted-by":"crossref","unstructured":"Zhou, Z., Dong, Q.: Learning occlusion-aware coarse-to-fine depth map for self-supervised monocular depth estimation. In: Proceedings of the 30th ACM International Conference on Multimedia, pp. 6386\u20136395 (2022)","DOI":"10.1145\/3503161.3548381"},{"key":"2356_CR40","doi-asserted-by":"publisher","DOI":"10.1016\/j.engappai.2025.110155","volume":"145","author":"S Luo","year":"2025","unstructured":"Luo, S., Qian, J., Huang, X., Zhou, Q., Hu, J.: A dual-discriminator network based on sobel gradient operator for digital twin-assisted fault diagnosis. Eng. Appl. Artif. Intell. 145, 110155 (2025)","journal-title":"Eng. Appl. Artif. Intell."},{"key":"2356_CR41","unstructured":"Lee, J.H., Han, M.-K., Ko, D.W., Suh, I.H.: From big to small: Multi-scale local planar guidance for monocular depth estimation. arXiv preprint arXiv:1907.10326 (2019)"},{"key":"2356_CR42","doi-asserted-by":"publisher","first-page":"3125","DOI":"10.1109\/TIP.2022.3164550","volume":"31","author":"Y-H Wu","year":"2022","unstructured":"Wu, Y.-H., Liu, Y., Zhang, L., Cheng, M.-M., Ren, B.: Edn: Salient object detection via extremely-downsampled network. IEEE Trans. Image Process. 31, 3125\u20133136 (2022)","journal-title":"IEEE Trans. Image Process."},{"key":"2356_CR43","doi-asserted-by":"publisher","DOI":"10.1016\/j.engappai.2023.107340","volume":"127","author":"A Barjasteh","year":"2024","unstructured":"Barjasteh, A., Ghafouri, S.H., Hashemi, M.: A hybrid model based on discrete wavelet transform (dwt) and bidirectional recurrent neural networks for wind speed prediction. Eng. Appl. Artif. Intell. 127, 107340 (2024)","journal-title":"Eng. Appl. Artif. Intell."},{"key":"2356_CR44","doi-asserted-by":"crossref","unstructured":"Zhou, Z., Fan, X., Shi, P., Xin, Y.: R-msfm: Recurrent multi-scale feature modulation for monocular depth estimating. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 12777\u201312786 (2021)","DOI":"10.1109\/ICCV48922.2021.01254"},{"issue":"4","key":"2356_CR45","doi-asserted-by":"publisher","first-page":"600","DOI":"10.1109\/TIP.2003.819861","volume":"13","author":"Z Wang","year":"2004","unstructured":"Wang, Z., Bovik, A.C., Sheikh, H.R., Simoncelli, E.P.: Image quality assessment: from error visibility to structural similarity. IEEE Trans. Image Process. 13(4), 600\u2013612 (2004)","journal-title":"IEEE Trans. Image Process."},{"key":"2356_CR46","doi-asserted-by":"crossref","unstructured":"Zhou, H., Greenwood, D., Taylor, S.: Self-supervised monocular depth estimation with internal feature fusion. arXiv preprint arXiv:2110.09482 (2021)","DOI":"10.5244\/C.35.208"},{"issue":"11","key":"2356_CR47","doi-asserted-by":"publisher","first-page":"1231","DOI":"10.1177\/0278364913491297","volume":"32","author":"A Geiger","year":"2013","unstructured":"Geiger, A., Lenz, P., Stiller, C., Urtasun, R.: Vision meets robotics: The kitti dataset. The international journal of robotics research 32(11), 1231\u20131237 (2013)","journal-title":"The international journal of robotics research"},{"key":"2356_CR48","doi-asserted-by":"crossref","unstructured":"Uhrig, J., Schneider, N., Schneider, L., Franke, U., Brox, T., Geiger, A.: Sparsity invariant cnns. In: 2017 International Conference on 3D Vision (3DV), pp. 11\u201320 (2017). IEEE","DOI":"10.1109\/3DV.2017.00012"},{"issue":"5","key":"2356_CR49","doi-asserted-by":"publisher","first-page":"824","DOI":"10.1109\/TPAMI.2008.132","volume":"31","author":"A Saxena","year":"2008","unstructured":"Saxena, A., Sun, M., Ng, A.Y.: Make3d: Learning 3d scene structure from a single still image. IEEE Trans. Pattern Anal. Mach. Intell. 31(5), 824\u2013840 (2008)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"2356_CR50","unstructured":"Paszke, A., Gross, S., Chintala, S., Chanan, G., Yang, E., DeVito, Z., Lin, Z., Desmaison, A., Antiga, L., Lerer, A.: Automatic differentiation in pytorch (2017)"},{"key":"2356_CR51","unstructured":"Loshchilov, I., Hutter, F.: Decoupled weight decay regularization. arXiv preprint arXiv:1711.05101 (2017)"},{"key":"2356_CR52","doi-asserted-by":"crossref","unstructured":"Wang, Y., Liang, Y., Xu, H., Jiao, S., Yu, H.: Sqldepth: Generalizable self-supervised fine-structured monocular depth estimation. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 38, pp. 5713\u20135721 (2024)","DOI":"10.1609\/aaai.v38i6.28383"},{"key":"2356_CR53","unstructured":"Choi, J., Jung, D., Lee, D., Kim, C.: Safenet: Self-supervised monocular depth estimation with semantic-aware feature extraction. arXiv preprint arXiv:2010.02893 (2020)"},{"key":"2356_CR54","doi-asserted-by":"crossref","unstructured":"Zhou, K., Hong, L., Chen, C., Xu, H., Ye, C., Hu, Q., Li, Z.: Devnet: Self-supervised monocular depth learning via density volume construction. In: European Conference on Computer Vision, pp. 125\u2013142 (2022). Springer","DOI":"10.1007\/978-3-031-19842-7_8"},{"key":"2356_CR55","doi-asserted-by":"crossref","unstructured":"Bae, J., Moon, S., Im, S.: Deep digging into the generalization of self-supervised monocular depth estimation. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 37, pp. 187\u2013196 (2023)","DOI":"10.1609\/aaai.v37i1.25090"},{"key":"2356_CR56","doi-asserted-by":"publisher","first-page":"54987","DOI":"10.52202\/075280-2399","volume":"36","author":"Y Sun","year":"2023","unstructured":"Sun, Y., Hariharan, B.: Dynamo-depth: Fixing unsupervised depth estimation for dynamical scenes. Adv. Neural. Inf. Process. Syst. 36, 54987\u201355005 (2023)","journal-title":"Adv. Neural. Inf. Process. Syst."},{"key":"2356_CR57","doi-asserted-by":"publisher","DOI":"10.1016\/j.engappai.2024.109313","volume":"138","author":"Z Cheng","year":"2024","unstructured":"Cheng, Z., Zhang, Y., Yu, Y., Song, Z., Tang, C.: Tinydepth: Lightweight self-supervised monocular depth estimation based on transformer. Eng. Appl. Artif. Intell. 138, 109313 (2024)","journal-title":"Eng. Appl. Artif. Intell."},{"key":"2356_CR58","doi-asserted-by":"publisher","DOI":"10.1016\/j.imavis.2024.105360","volume":"154","author":"L Wu","year":"2025","unstructured":"Wu, L., Wang, L., Wei, G., Yu, Y.: Hpd-depth: High performance decoding network for self-supervised monocular depth estimation. Image Vis. Comput. 154, 105360 (2025)","journal-title":"Image Vis. Comput."},{"key":"2356_CR59","doi-asserted-by":"crossref","unstructured":"Feng, C., Zhang, C., Chen, Z., Hu, W., Lu, K., Ge, L.: Self-supervised monocular depth estimation with dual-path encoders and offset field interpolation. IEEE Transactions on Image Processing (2025)","DOI":"10.1109\/TIP.2025.3533207"},{"key":"2356_CR60","doi-asserted-by":"crossref","unstructured":"Shu, C., Yu, K., Duan, Z., Yang, K.: Feature-metric loss for self-supervised learning of depth and egomotion. In: European Conference on Computer Vision, pp. 572\u2013588 (2020). Springer","DOI":"10.1007\/978-3-030-58529-7_34"},{"key":"2356_CR61","doi-asserted-by":"crossref","unstructured":"Yan, J., Zhao, H., Bu, P., Jin, Y.: Channel-wise attention-based network for self-supervised monocular depth estimation. In: 2021 International Conference on 3D Vision (3DV), pp. 464\u2013473 (2021). IEEE","DOI":"10.1109\/3DV53792.2021.00056"},{"key":"2356_CR62","doi-asserted-by":"crossref","unstructured":"Zhao, C., Zhang, Y., Poggi, M., Tosi, F., Guo, X., Zhu, Z., Huang, G., Tang, Y., Mattoccia, S.: Monovit: Self-supervised monocular depth estimation with a vision transformer. In: 2022 International Conference on 3D Vision (3DV), pp. 668\u2013678 (2022). IEEE","DOI":"10.1109\/3DV57658.2022.00077"},{"key":"2356_CR63","doi-asserted-by":"publisher","DOI":"10.1016\/j.knosys.2024.112552","volume":"304","author":"B Zhao","year":"2024","unstructured":"Zhao, B., He, H., Xu, H., Shi, P., Hao, X., Huang, G.: Lda-mono: A lightweight dual aggregation network for self-supervised monocular depth estimation. Knowl.-Based Syst. 304, 112552 (2024)","journal-title":"Knowl.-Based Syst."},{"key":"2356_CR64","doi-asserted-by":"crossref","unstructured":"Guizilini, V., Ambru?, R., Chen, D., Zakharov, S., Gaidon, A.: Multi-frame self-supervised depth with transformers. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 160\u2013170 (2022)","DOI":"10.1109\/CVPR52688.2022.00026"}],"container-title":["Multimedia Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-026-02356-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00530-026-02356-0","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-026-02356-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,7,11]],"date-time":"2026-07-11T10:19:45Z","timestamp":1783765185000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00530-026-02356-0"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,5,6]]},"references-count":64,"journal-issue":{"issue":"4","published-print":{"date-parts":[[2026,6]]}},"alternative-id":["2356"],"URL":"https:\/\/doi.org\/10.1007\/s00530-026-02356-0","relation":{},"ISSN":["0942-4962","1432-1882"],"issn-type":[{"value":"0942-4962","type":"print"},{"value":"1432-1882","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,5,6]]},"assertion":[{"value":"26 August 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"9 March 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"6 May 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no competing interests.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interests"}}],"article-number":"275"}}