{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,21]],"date-time":"2026-06-21T04:20:25Z","timestamp":1782015625214,"version":"3.54.5"},"reference-count":108,"publisher":"Springer Science and Business Media LLC","issue":"9","license":[{"start":{"date-parts":[[2025,5,23]],"date-time":"2025-05-23T00:00:00Z","timestamp":1747958400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,5,23]],"date-time":"2025-05-23T00:00:00Z","timestamp":1747958400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62072242"],"award-info":[{"award-number":["62072242"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100020771","name":"Young Scientists Fund of the National Natural Science Foundation of China","doi-asserted-by":"crossref","award":["62206134"],"award-info":[{"award-number":["62206134"]}],"id":[{"id":"10.13039\/501100020771","id-type":"DOI","asserted-by":"crossref"}]},{"name":"Tianjin Key Laboratory of Visual Computing and Intelligent Perception (VCIP), and the Fundamental Research Funds for the Central Universities","award":["070-63233084"],"award-info":[{"award-number":["070-63233084"]}]},{"DOI":"10.13039\/501100012226","name":"Fundamental Research Funds for the Central Universities","doi-asserted-by":"publisher","award":["62306238"],"award-info":[{"award-number":["62306238"]}],"id":[{"id":"10.13039\/501100012226","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100012226","name":"Fundamental Research Funds for the Central Universities","doi-asserted-by":"publisher","award":["62376121"],"award-info":[{"award-number":["62376121"]}],"id":[{"id":"10.13039\/501100012226","id-type":"DOI","asserted-by":"publisher"}]},{"name":"ostgraduate Research & Practice Innovation Program of Jiangsu Province","award":["KYCX23_0471"],"award-info":[{"award-number":["KYCX23_0471"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int J Comput Vis"],"published-print":{"date-parts":[[2025,9]]},"DOI":"10.1007\/s11263-025-02470-y","type":"journal-article","created":{"date-parts":[[2025,5,23]],"date-time":"2025-05-23T10:41:46Z","timestamp":1747996906000},"page":"6051-6073","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":5,"title":["RigNet++: Semantic Assisted Repetitive Image Guided Network for Depth Completion"],"prefix":"10.1007","volume":"133","author":[{"ORCID":"https:\/\/orcid.org\/0000-0003-3502-438X","authenticated-orcid":false,"given":"Zhiqiang","family":"Yan","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xiang","family":"Li","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Le","family":"Hui","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Zhenyu","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jun","family":"Li","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jian","family":"Yang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2025,5,23]]},"reference":[{"key":"2470_CR1","doi-asserted-by":"crossref","unstructured":"Albanis, G., Zioulis, N., Drakoulis, P., Gkitsas, V., Sterzentsenko, V., Alvarez, F., Zarpalas, D., & Daras, P. (2021). Pano3d: A holistic benchmark and a solid baseline for $$360^{\\circ }$$ depth estimation. In: CVPR Workshops, pp. 3722\u20133732. IEEE","DOI":"10.1109\/CVPRW53098.2021.00413"},{"issue":"1","key":"2470_CR2","doi-asserted-by":"publisher","first-page":"9","DOI":"10.1089\/cpb.2007.9935","volume":"11","author":"C Armbr\u00fcster","year":"2008","unstructured":"Armbr\u00fcster, C., Wolter, M., Kuhlen, T., Spijkers, W., & Fimm, B. (2008). Depth perception in virtual reality: distance estimations in peri-and extrapersonal space. Cyberpsychology & Behavior, 11(1), 9\u201315.","journal-title":"Cyberpsychology & Behavior"},{"key":"2470_CR3","doi-asserted-by":"crossref","unstructured":"Cai, Z., & Vasconcelos, N. (2018). Cascade r-cnn: Delving into high quality object detection. In: CVPR, pp. 6154\u20136162","DOI":"10.1109\/CVPR.2018.00644"},{"key":"2470_CR4","doi-asserted-by":"crossref","unstructured":"Chen, D., Huang, T., Song, Z., Deng, S., & Jia, T. (2023). Agg-net: Attention guided gated-convolutional network for depth image completion. In: ICCV, pp. 8853\u20138862","DOI":"10.1109\/ICCV51070.2023.00813"},{"key":"2470_CR5","doi-asserted-by":"crossref","unstructured":"Chen, T., Zhu, L., Deng, C., Cao, R., Wang, Y., Zhang, S., Li, Z., Sun, L., Zang, Y., & Mao, P. (2023). Sam-adapter: Adapting segment anything in underperformed scenes. In: ICCV, pp. 3367\u20133375","DOI":"10.1109\/ICCVW60793.2023.00361"},{"key":"2470_CR6","doi-asserted-by":"crossref","unstructured":"Chen, Y., Yang, B., Liang, M., & Urtasun, R. (2019). Learning joint 2d-3d representations for depth completion. In: ICCV, pp. 10023\u201310032","DOI":"10.1109\/ICCV.2019.01012"},{"key":"2470_CR7","doi-asserted-by":"crossref","unstructured":"Cheng, X., Wang, P., Guan, C., & Yang, R. (2020). Cspn++: Learning context and resource aware convolutional spatial propagation networks for depth completion. In: AAAI, pp. 10615\u201310622","DOI":"10.1609\/aaai.v34i07.6635"},{"key":"2470_CR8","doi-asserted-by":"crossref","unstructured":"Cheng, X., Wang, P., & Yang, R. (2018). Learning depth with convolutional spatial propagation network. In: ECCV, pp. 103\u2013119","DOI":"10.1007\/978-3-030-01270-0_7"},{"key":"2470_CR9","doi-asserted-by":"crossref","unstructured":"Chodosh, N., Wang, C., & Lucey, S. (2018). Deep convolutional compressed sensing for lidar depth completion. In: ACCV, pp. 499\u2013513","DOI":"10.1007\/978-3-030-20887-5_31"},{"key":"2470_CR10","doi-asserted-by":"crossref","unstructured":"Cui, Z., Heng, L., Yeo, Y.C., Geiger, A., Pollefeys, M., & Sattler, T. (2019). Real-time dense mapping for self-driving vehicles using fisheye cameras. In: ICRA, pp. 6087\u20136093","DOI":"10.1109\/ICRA.2019.8793884"},{"key":"2470_CR11","doi-asserted-by":"crossref","unstructured":"Dey, A., Jarvis, G., Sandor, C., & Reitmayr, G. (2012). Tablet versus phone: Depth perception in handheld augmented reality. In: ISMAR, pp. 187\u2013196","DOI":"10.1109\/ISMAR.2012.6402556"},{"key":"2470_CR12","unstructured":"Dosovitskiy, A., Beyer, L., Kolesnikov, A., Weissenborn, D., Zhai, X., Unterthiner, T., Dehghani, M., Minderer, M., Heigold, G., Gelly, S., et\u00a0al. (2020). An image is worth 16x16 words: Transformers for image recognition at scale. arXiv preprint arXiv:2010.11929"},{"key":"2470_CR13","unstructured":"Dosovitskiy, A., Ros, G., Codevilla, F., Lopez, A., & Koltun, V. (2017). Carla: An open urban driving simulator. In: CoRL, pp. 1\u201316. PMLR"},{"issue":"10","key":"2470_CR14","doi-asserted-by":"publisher","first-page":"2423","DOI":"10.1109\/TPAMI.2019.2929170","volume":"42","author":"A Eldesokey","year":"2020","unstructured":"Eldesokey, A., Felsberg, M., & Khan, F. S. (2020). Confidence propagation through cnns for guided sparse depth regression. IEEE Transactions on Pattern Analysis and Machine Intelligence, 42(10), 2423\u20132436.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"2470_CR15","doi-asserted-by":"crossref","unstructured":"Gaidon, A., Wang, Q., Cabon, Y., & Vig, E. (2016). Virtual worlds as proxy for multi-object tracking analysis. In: CVPR, pp. 4340\u20134349","DOI":"10.1109\/CVPR.2016.470"},{"key":"2470_CR16","doi-asserted-by":"crossref","unstructured":"Gao, R., Chen, C., Al-Halah, Z., Schissler, C., & Grauman, K. (2020). Visualechoes: Spatial image representation learning through echolocation. In: ECCV, pp. 658\u2013676. Springer","DOI":"10.1007\/978-3-030-58545-7_38"},{"key":"2470_CR17","doi-asserted-by":"crossref","unstructured":"Ghiasi, G., Lin, T.Y., & Le, Q.V. (2019). Nas-fpn: Learning scalable feature pyramid architecture for object detection. In: CVPR, pp. 7036\u20137045","DOI":"10.1109\/CVPR.2019.00720"},{"key":"2470_CR18","unstructured":"Glorot, X., Bordes, A., & Bengio, Y. (2011). Deep sparse rectifier neural networks. In: AISTATS, pp. 315\u2013323. JMLR Workshop and Conference Proceedings"},{"key":"2470_CR19","doi-asserted-by":"publisher","first-page":"14","DOI":"10.1016\/j.imavis.2017.07.003","volume":"68","author":"C H\u00e4ne","year":"2017","unstructured":"H\u00e4ne, C., Heng, L., Lee, G. H., Fraundorfer, F., Furgale, P., Sattler, T., & Pollefeys, M. (2017). 3d visual perception for self-driving cars using a multi-camera system: Calibration, mapping, localization, and obstacle detection. Image and Vision Computing, 68, 14\u201327.","journal-title":"Image and Vision Computing"},{"key":"2470_CR20","doi-asserted-by":"crossref","unstructured":"He, K., Chen, X., Xie, S., Li, Y., Doll\u00e1r, P., & Girshick, R. (2022). Masked autoencoders are scalable vision learners. In: CVPR, pp. 16000\u201316009","DOI":"10.1109\/CVPR52688.2022.01553"},{"key":"2470_CR21","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., & Sun, J. (2016). Deep residual learning for image recognition. In: CVPR, pp. 770\u2013778","DOI":"10.1109\/CVPR.2016.90"},{"key":"2470_CR22","doi-asserted-by":"crossref","unstructured":"He, L., Zhu, H., Li, F., Bai, H., Cong, R., Zhang, C., Lin, C., Liu, M., & Zhao, Y. (2021). Towards fast and accurate real-world depth super-resolution: Benchmark dataset and baseline. In: CVPR, pp. 9229\u20139238","DOI":"10.1109\/CVPR46437.2021.00911"},{"key":"2470_CR23","unstructured":"Howard, A.G., Zhu, M., Chen, B., Kalenichenko, D., Wang, W., Weyand, T., Andreetto, M., & Adam, H. (2017). Mobilenets: Efficient convolutional neural networks for mobile vision applications. arXiv preprint arXiv:1704.04861"},{"key":"2470_CR24","doi-asserted-by":"crossref","unstructured":"Hu, J., Shen, L., & Sun, G. (2018). Squeeze-and-excitation networks. In: CVPR, pp. 7132\u20137141","DOI":"10.1109\/CVPR.2018.00745"},{"key":"2470_CR25","doi-asserted-by":"crossref","unstructured":"Hu, M., Wang, S., Li, B., Ning, S., Fan, L., & Gong, X. (2021). Penet: Towards precise and efficient image guided depth completion. In: ICRA","DOI":"10.1109\/ICRA48506.2021.9561035"},{"key":"2470_CR26","doi-asserted-by":"crossref","unstructured":"Huang, G., Liu, Z., Van Der\u00a0Maaten, L., & Weinberger, K.Q. (2017). Densely connected convolutional networks. In: CVPR, pp. 4700\u20134708","DOI":"10.1109\/CVPR.2017.243"},{"key":"2470_CR27","doi-asserted-by":"crossref","unstructured":"Huang, Y., Yang, X., Liu, L., Zhou, H., Chang, A., Zhou, X., Chen, R., Yu, J., Chen, J., Chen, C., et\u00a0al. (2023). Segment anything model for medical images? arXiv preprint arXiv:2304.14660","DOI":"10.1016\/j.media.2023.103061"},{"key":"2470_CR28","doi-asserted-by":"crossref","unstructured":"Huang, Y.K., Wu, T.H., Liu, Y.C., & Hsu, W.H. (2019). Indoor depth completion with boundary consistency and self-attention. In: ICCV Workshops, pp. 0\u20130","DOI":"10.1109\/ICCVW.2019.00137"},{"key":"2470_CR29","doi-asserted-by":"crossref","unstructured":"Huynh, L., Nguyen, P., Matas, J., Rahtu, E., & Heikkil\u00e4, J. (2021). Boosting monocular depth estimation with lightweight 3d point fusion. In: ICCV, pp. 12767\u201312776","DOI":"10.1109\/ICCV48922.2021.01253"},{"key":"2470_CR30","doi-asserted-by":"crossref","unstructured":"Imran, S., Liu, X., & Morris, D. (2021). Depth completion with twin surface extrapolation at occlusion boundaries. In: CVPR, pp. 2583\u20132592","DOI":"10.1109\/CVPR46437.2021.00261"},{"key":"2470_CR31","unstructured":"Ioffe, S., & Szegedy, C. (2015). Batch normalization: Accelerating deep network training by reducing internal covariate shift. In: ICML, pp. 448\u2013456. PMLR"},{"key":"2470_CR32","doi-asserted-by":"crossref","unstructured":"Jaritz, M., De\u00a0Charette, R., Wirbel, E., Perrotton, X., & Nashashibi, F. (2018). Sparse and dense data with cnns: Depth completion and semantic segmentation. In: 3DV, pp. 52\u201360","DOI":"10.1109\/3DV.2018.00017"},{"key":"2470_CR33","doi-asserted-by":"crossref","unstructured":"Jaritz, M., De\u00a0Charette, R., Wirbel, E., Perrotton, X., & Nashashibi, F. (2018). Sparse and dense data with cnns: Depth completion and semantic segmentation. In: 3DV, pp. 52\u201360. IEEE","DOI":"10.1109\/3DV.2018.00017"},{"issue":"2","key":"2470_CR34","doi-asserted-by":"publisher","first-page":"1519","DOI":"10.1109\/LRA.2021.3058957","volume":"6","author":"H Jiang","year":"2021","unstructured":"Jiang, H., Sheng, Z., Zhu, S., Dong, Z., & Huang, R. (2021). Unifuse: Unidirectional fusion for 360 panorama depth estimation. IEEE Robotics and Automation Letters, 6(2), 1519\u20131526.","journal-title":"IEEE Robotics and Automation Letters"},{"key":"2470_CR35","unstructured":"Kingma, D.P., & Ba, J. (2014). Adam: A method for stochastic optimization. arXiv preprint arXiv:1412.6980"},{"key":"2470_CR36","doi-asserted-by":"crossref","unstructured":"Kirillov, A., Mintun, E., Ravi, N., Mao, H., Rolland, C., Gustafson, L., Xiao, T., Whitehead, S., Berg, A.C., Lo, W.Y., et\u00a0al. (2023). Segment anything. arXiv preprint arXiv:2304.02643","DOI":"10.1109\/ICCV51070.2023.00371"},{"key":"2470_CR37","doi-asserted-by":"crossref","unstructured":"Ku, J., Harakeh, A., & Waslander, S.L. (2018). In defense of classical image processing: Fast depth completion on the cpu. In: CRV, pp. 16\u201322","DOI":"10.1109\/CRV.2018.00013"},{"key":"2470_CR38","doi-asserted-by":"crossref","unstructured":"Lagos, J.P., & Rahtu, E. (2022). Semsegdepth: A combined model for semantic segmentation and depth completion. arXiv preprint arXiv:2209.00381","DOI":"10.5220\/0010838500003124"},{"key":"2470_CR39","doi-asserted-by":"crossref","unstructured":"Lee, B.U., Lee, K., & Kweon, I.S. (2021). Depth completion using plane-residual representation. In: CVPR, pp. 13916\u201313925","DOI":"10.1109\/CVPR46437.2021.01370"},{"key":"2470_CR40","doi-asserted-by":"crossref","unstructured":"Levin, A., Lischinski, D., & Weiss, Y. (2004). Colorization using optimization. In: SIGGRAPH, pp. 689\u2013694. ACM","DOI":"10.1145\/1186562.1015780"},{"key":"2470_CR41","doi-asserted-by":"crossref","unstructured":"Li, A., Yuan, Z., Ling, Y., Chi, W., Zhang, C., et\u00a0al. (2020). A multi-scale guided cascade hourglass network for depth completion. In: WACV, pp. 32\u201340","DOI":"10.1109\/WACV45572.2020.9093407"},{"key":"2470_CR42","doi-asserted-by":"crossref","unstructured":"Li, S., Liu, M., Zhang, Y., Chen, S., Li, H., Chen, H., & Dou, Z. (2023). Sam-deblur: Let segment anything boost image deblurring. arXiv preprint arXiv:2309.02270","DOI":"10.1109\/ICASSP48485.2024.10445844"},{"key":"2470_CR43","doi-asserted-by":"crossref","unstructured":"Li, X., Wang, W., Hu, X., & Yang, J. (2019). Selective kernel networks. In: CVPR, pp. 510\u2013519","DOI":"10.1109\/CVPR.2019.00060"},{"key":"2470_CR44","first-page":"12934","volume":"35","author":"Y Li","year":"2022","unstructured":"Li, Y., Hu, J., Wen, Y., Evangelidis, G., Salahi, K., Wang, Y., Tulyakov, S., & Ren, J. (2022). Rethinking vision transformers for mobilenet size and speed. NeurIPS, 35, 12934\u201312949.","journal-title":"NeurIPS"},{"key":"2470_CR45","doi-asserted-by":"crossref","unstructured":"Li, Y., Hu, J., Wen, Y., Evangelidis, G., Salahi, K., Wang, Y., Tulyakov, S., & Ren, J. (2023). Rethinking vision transformers for mobilenet size and speed. In: ICCV, pp. 16889\u201316900","DOI":"10.1109\/ICCV51070.2023.01549"},{"key":"2470_CR46","doi-asserted-by":"crossref","unstructured":"Lin, J., Liu, L., Lu, D., & Jia, K. (2023). Sam-6d: Segment anything model meets zero-shot 6d object pose estimation. arXiv preprint arXiv:2311.15707","DOI":"10.1109\/CVPR52733.2024.02636"},{"key":"2470_CR47","doi-asserted-by":"crossref","unstructured":"Lin, T.Y., Doll\u00e1r, P., Girshick, R., He, K., Hariharan, B., & Belongie, S. (2017). Feature pyramid networks for object detection. In: CVPR, pp. 2117\u20132125","DOI":"10.1109\/CVPR.2017.106"},{"key":"2470_CR48","doi-asserted-by":"publisher","first-page":"1638","DOI":"10.1609\/aaai.v36i2.20055","volume":"36","author":"Y Lin","year":"2022","unstructured":"Lin, Y., Cheng, T., Zhong, Q., Zhou, W., & Yang, H. (2022). Dynamic spatial propagation network for depth completion. AAAI, 36, 1638\u20131646.","journal-title":"AAAI"},{"key":"2470_CR49","doi-asserted-by":"crossref","unstructured":"Lin, Y., Yang, H., Cheng, T., Zhou, W., & Yin, Z. (2023). Dyspn: Learning dynamic affinity for image-guided depth completion. IEEE Transactions on Circuits and Systems for Video Technology","DOI":"10.1109\/TCSVT.2023.3321668"},{"key":"2470_CR50","doi-asserted-by":"publisher","first-page":"2136","DOI":"10.1609\/aaai.v35i3.16311","volume":"35","author":"L Liu","year":"2021","unstructured":"Liu, L., Song, X., Lyu, X., Diao, J., Wang, M., Liu, Y., & Zhang, L. (2021). Fcfr-net: Feature fusion based coarse-to-fine residual learning for depth completion. AAAI, 35, 2136\u20132144.","journal-title":"AAAI"},{"issue":"2","key":"2470_CR51","doi-asserted-by":"publisher","first-page":"920","DOI":"10.1109\/LRA.2023.3234776","volume":"8","author":"L Liu","year":"2023","unstructured":"Liu, L., Song, X., Sun, J., Lyu, X., Li, L., Liu, Y., & Zhang, L. (2023). Mff-net: Towards efficient monocular depth completion with multi-modal feature fusion. IEEE Robotics and Automation Letters, 8(2), 920\u2013927.","journal-title":"IEEE Robotics and Automation Letters"},{"key":"2470_CR52","unstructured":"Liu, S., De\u00a0Mello, S., Gu, J., Zhong, G., Yang, M.H., & Kautz, J. (2017). Learning affinity via spatial propagation networks. In: NeurIPS, vol.\u00a030"},{"key":"2470_CR53","doi-asserted-by":"crossref","unstructured":"Liu, S., Qi, L., Qin, H., Shi, J., Jia, & J. (2018). Path aggregation network for instance segmentation. In: CVPR, pp. 8759\u20138768","DOI":"10.1109\/CVPR.2018.00913"},{"key":"2470_CR54","doi-asserted-by":"crossref","unstructured":"Liu, X., Shao, X., Wang, B., Li, Y., & Wang, S. (2022). Graphcspn: Geometry-aware depth completion via dynamic gcns. In: ECCV, pp. 90\u2013107. Springer","DOI":"10.1007\/978-3-031-19827-4_6"},{"key":"2470_CR55","doi-asserted-by":"publisher","first-page":"11653","DOI":"10.1609\/aaai.v34i07.6834","volume":"34","author":"Y Liu","year":"2020","unstructured":"Liu, Y., Wang, Y., Wang, S., Liang, T., Zhao, Q., Tang, Z., & Ling, H. (2020). Cbnet: A novel composite backbone network architecture for object detection. AAAI, 34, 11653\u201311660.","journal-title":"AAAI"},{"key":"2470_CR56","doi-asserted-by":"crossref","unstructured":"Lu, K., Barnes, N., Anwar, S., & Zheng, L. (2020). From depth what can you see? depth completion via auxiliary image reconstruction. In: CVPR, pp. 11306\u201311315","DOI":"10.1109\/CVPR42600.2020.01132"},{"key":"2470_CR57","doi-asserted-by":"crossref","unstructured":"Ma, F., Cavalheiro, G.V., & Karaman, S. (2019). Self-supervised sparse-to-dense: Self-supervised depth completion from lidar and monocular camera. In: ICRA","DOI":"10.1109\/ICRA.2019.8793637"},{"key":"2470_CR58","doi-asserted-by":"crossref","unstructured":"Ma, F., & Karaman, S. (2018). Sparse-to-dense: Depth prediction from sparse depth samples and a single image. In: ICRA, pp. 4796\u20134803. IEEE","DOI":"10.1109\/ICRA.2018.8460184"},{"key":"2470_CR59","doi-asserted-by":"crossref","unstructured":"Ma, J., & Wang, B. (2023). Segment anything in medical images. arXiv preprint arXiv:2304.12306","DOI":"10.1038\/s41467-024-44824-z"},{"key":"2470_CR60","doi-asserted-by":"publisher","DOI":"10.1016\/j.media.2023.102918","volume":"89","author":"MA Mazurowski","year":"2023","unstructured":"Mazurowski, M. A., Dong, H., Gu, H., Yang, J., Konz, N., & Zhang, Y. (2023). Segment anything model for medical image analysis: an experimental study. Medical Image Analysis, 89, Article 102918.","journal-title":"Medical Image Analysis"},{"key":"2470_CR61","doi-asserted-by":"publisher","first-page":"120781","DOI":"10.1109\/ACCESS.2022.3214316","volume":"10","author":"D Nazir","year":"2022","unstructured":"Nazir, D., Pagani, A., Liwicki, M., Stricker, D., & Afzal, M. Z. (2022). Semattnet: Toward attention-based semantic aware guided depth completion. IEEE Access, 10, 120781\u2013120791.","journal-title":"IEEE Access"},{"key":"2470_CR62","doi-asserted-by":"crossref","unstructured":"Parida, K.K., Srivastava, S., & Sharma, G. (2021). Beyond image to depth: Improving depth prediction using echoes. In: CVPR, pp. 8268\u20138277","DOI":"10.1109\/CVPR46437.2021.00817"},{"key":"2470_CR63","doi-asserted-by":"crossref","unstructured":"Park, J., Joo, K., Hu, Z., Liu, C.K., & Kweon, I.S. (2020). Non-local spatial propagation network for depth completion. In: ECCV","DOI":"10.1007\/978-3-030-58601-0_8"},{"key":"2470_CR64","doi-asserted-by":"crossref","unstructured":"Qiao, S., Chen, L.C., & Yuille, A. (2021). Detectors: Detecting objects with recursive feature pyramid and switchable atrous convolution. In: CVPR, pp. 10213\u201310224","DOI":"10.1109\/CVPR46437.2021.01008"},{"key":"2470_CR65","doi-asserted-by":"crossref","unstructured":"Qiu, J., Cui, Z., Zhang, Y., Zhang, X., Liu, S., Zeng, B., & Pollefeys, M. (2019). Deeplidar: Deep surface normal guided depth prediction for outdoor scene from sparse lidar data and single color image. In: CVPR, pp. 3313\u20133322","DOI":"10.1109\/CVPR.2019.00343"},{"key":"2470_CR66","doi-asserted-by":"crossref","unstructured":"Qu, C., Liu, W., & Taylor, C.J. (2021). Bayesian deep basis fitting for depth completion with uncertainty. In: ICCV, pp. 16147\u201316157","DOI":"10.1109\/ICCV48922.2021.01584"},{"key":"2470_CR67","first-page":"91","volume":"28","author":"S Ren","year":"2015","unstructured":"Ren, S., He, K., Girshick, R., & Sun, J. (2015). Faster r-cnn: Towards real-time object detection with region proposal networks. NeurIPS, 28, 91\u201399.","journal-title":"NeurIPS"},{"key":"2470_CR68","doi-asserted-by":"crossref","unstructured":"Rey-Area, M., Yuan, M., & Richardt, C. (2022). 360monodepth: High-resolution 360deg monocular depth estimation. In: CVPR, pp. 3762\u20133772","DOI":"10.1109\/CVPR52688.2022.00374"},{"key":"2470_CR69","doi-asserted-by":"crossref","unstructured":"Rho, K., Ha, J., & Kim, Y. (2022). Guideformer: Transformers for image guided depth completion. In: CVPR, pp. 6250\u20136259","DOI":"10.1109\/CVPR52688.2022.00615"},{"key":"2470_CR70","doi-asserted-by":"crossref","unstructured":"Ronneberger, O., Fischer, P., & Brox, T. (2015). U-net: Convolutional networks for biomedical image segmentation. In: MICCAI, pp. 234\u2013241. Springer","DOI":"10.1007\/978-3-319-24574-4_28"},{"key":"2470_CR71","doi-asserted-by":"crossref","unstructured":"Schneider, N., Schneider, L., Pinggera, P., Franke, U., Pollefeys, M., & Stiller, C. (2016). Semantically guided depth upsampling. In: GCPR, pp. 37\u201348. Springer","DOI":"10.1007\/978-3-319-45886-1_4"},{"key":"2470_CR72","doi-asserted-by":"crossref","unstructured":"Shen, Z., Lin, C., Liao, K., Nie, L., Zheng, Z., & Zhao, Y. (2022). Panoformer: Panorama transformer for indoor 360 depth estimation. In: ECCV, pp. 195\u2013211. Springer","DOI":"10.1007\/978-3-031-19769-7_12"},{"key":"2470_CR73","doi-asserted-by":"crossref","unstructured":"Silberman, N., Hoiem, D., Kohli, P., & Fergus, R. (2012). Indoor segmentation and support inference from rgbd images. In: ECCV, pp. 746\u2013760. Springer","DOI":"10.1007\/978-3-642-33715-4_54"},{"key":"2470_CR74","doi-asserted-by":"crossref","unstructured":"Song, X., Dai, Y., Zhou, D., Liu, L., Li, W., Li, H., & Yang, R. (2020). Channel attention based iterative residual learning for depth map super-resolution. In: CVPR, pp. 5631\u20135640","DOI":"10.1109\/CVPR42600.2020.00567"},{"key":"2470_CR75","doi-asserted-by":"crossref","unstructured":"Sun, C., Sun, M., & Chen, H.T. (2021). Hohonet: 360 indoor holistic understanding with latent horizontal features. In: CVPR, pp. 2573\u20132582","DOI":"10.1109\/CVPR46437.2021.00260"},{"key":"2470_CR76","unstructured":"Tan, M., & Le, Q. (2019). Efficientnet: Rethinking model scaling for convolutional neural networks. In: ICML, pp. 6105\u20136114. PMLR"},{"key":"2470_CR77","unstructured":"Tan, M., & Le, Q. (2021). Efficientnetv2: Smaller models and faster training. In: ICML, pp. 10096\u201310106. PMLR"},{"key":"2470_CR78","doi-asserted-by":"crossref","unstructured":"Tan, M., Pang, R., & Le, Q.V. (2020). Efficientdet: Scalable and efficient object detection. In: CVPR, pp. 10781\u201310790","DOI":"10.1109\/CVPR42600.2020.01079"},{"key":"2470_CR79","doi-asserted-by":"publisher","first-page":"1116","DOI":"10.1109\/TIP.2020.3040528","volume":"30","author":"J Tang","year":"2020","unstructured":"Tang, J., Tian, F. P., Feng, W., Li, J., & Tan, P. (2020). Learning guided convolutional network for depth completion. IEEE Transactions on Image Processing, 30, 1116\u20131129.","journal-title":"IEEE Transactions on Image Processing"},{"key":"2470_CR80","doi-asserted-by":"crossref","unstructured":"Uhrig, J., Schneider, N., Schneider, L., Franke, U., Brox, T., & Geiger, A. (2017). Sparsity invariant cnns. In: 3DV, pp. 11\u201320","DOI":"10.1109\/3DV.2017.00012"},{"key":"2470_CR81","doi-asserted-by":"crossref","unstructured":"Van\u00a0Gansbeke, W., Neven, D., De\u00a0Brabandere, B., & Van\u00a0Gool, L. (2019). Sparse and noisy lidar completion with rgb guidance and uncertainty. In: MVA, pp. 1\u20136","DOI":"10.23919\/MVA.2019.8757939"},{"key":"2470_CR82","doi-asserted-by":"crossref","unstructured":"Wang, K., Zhang, Z., Yan, Z., Li, X., Xu, B., Li, J., & Yang, J. (2021). Regularizing nighttime weirdness: Efficient self-supervised monocular depth estimation in the dark. In: ICCV, pp. 16055\u201316064","DOI":"10.1109\/ICCV48922.2021.01575"},{"issue":"4","key":"2470_CR83","doi-asserted-by":"publisher","first-page":"11476","DOI":"10.1109\/LRA.2022.3201193","volume":"7","author":"Y Wang","year":"2022","unstructured":"Wang, Y., Dai, Y., Liu, Q., Yang, P., Sun, J., & Li, B. (2022). Cu-net: Lidar depth-only completion with coupled u-net. IEEE Robotics and Automation Letters, 7(4), 11476\u201311483.","journal-title":"IEEE Robotics and Automation Letters"},{"key":"2470_CR84","doi-asserted-by":"crossref","unstructured":"Wang, Y., Li, B., Zhang, G., Liu, Q., Gao, T., & Dai, Y. (2023). Lrru: Long-short range recurrent updating networks for depth completion. In: ICCV, pp. 9422\u20139432","DOI":"10.1109\/ICCV51070.2023.00864"},{"key":"2470_CR85","doi-asserted-by":"crossref","unstructured":"Xu, Y., Zhu, X., Shi, J., Zhang, G., Bao, H., & Li, H. (2019). Depth completion from sparse lidar data with depth-normal constraints. In: ICCV, pp. 2811\u20132820","DOI":"10.1109\/ICCV.2019.00290"},{"key":"2470_CR86","doi-asserted-by":"crossref","unstructured":"Xu, Z., Yin, H., & Yao, J. (2020). Deformable spatial propagation networks for depth completion. In: ICIP, pp. 913\u2013917. IEEE","DOI":"10.1109\/ICIP40778.2020.9191138"},{"key":"2470_CR87","doi-asserted-by":"publisher","first-page":"8779","DOI":"10.1609\/aaai.v36i8.20858","volume":"36","author":"F Xue","year":"2022","unstructured":"Xue, F., Shi, Z., Wei, F., Lou, Y., Liu, Y., & You, Y. (2022). Go wider instead of deeper. AAAI, 36, 8779\u20138787.","journal-title":"AAAI"},{"key":"2470_CR88","unstructured":"Yan, Z., Li, X., Wang, K., Chen, S., Li, J., & Yang, J. (2023). Distortion and uncertainty aware loss for panoramic depth completion. In: ICML, pp. 39099\u201339109. PMLR"},{"key":"2470_CR89","doi-asserted-by":"crossref","unstructured":"Yan, Z., Li, X., Wang, K., Zhang, Z., Li, J., & Yang, J. (2022). Multi-modal masked pre-training for monocular panoramic depth completion. In: ECCV, pp. 378\u2013395. Springer","DOI":"10.1007\/978-3-031-19769-7_22"},{"key":"2470_CR90","unstructured":"Yan, Z., Wang, K., Li, X., Zhang, Z., Li, G., Li, J., & Yang, J. (2022). Learning complementary correlations for depth super-resolution with incomplete data in real world. IEEE Transactions on Neural Networks and Learning Systems"},{"key":"2470_CR91","doi-asserted-by":"crossref","unstructured":"Yan, Z., Wang, K., Li, X., Zhang, Z., Li, J., & Yang, J. (2022). Rignet: Repetitive image guided network for depth completion. In: ECCV, pp. 214\u2013230. Springer","DOI":"10.1007\/978-3-031-19812-0_13"},{"key":"2470_CR92","doi-asserted-by":"publisher","first-page":"3109","DOI":"10.1609\/aaai.v37i3.25415","volume":"37","author":"Z Yan","year":"2023","unstructured":"Yan, Z., Wang, K., Li, X., Zhang, Z., Li, J., & Yang, J. (2023). Desnet: Decomposed scale-consistent network for unsupervised depth completion. AAAI, 37, 3109\u20133117.","journal-title":"AAAI"},{"key":"2470_CR93","doi-asserted-by":"crossref","unstructured":"Yan, Z., Zheng, Y., Wang, K., Li, X., Zhang, Z., Chen, S., Li, J., & Yang, J. (2023). Learnable differencing center for nighttime depth perception. arXiv preprint arXiv:2306.14538","DOI":"10.1007\/s44267-024-00048-9"},{"key":"2470_CR94","unstructured":"Yang, J., Gao, M., Li, Z., Gao, S., Wang, F., & Zheng, F. (2023). Track anything: Segment anything meets videos. arXiv preprint arXiv:2304.11968"},{"key":"2470_CR95","doi-asserted-by":"crossref","unstructured":"Yang, Y., Wong, A., & Soatto, S. (2020). Dense depth posterior (ddp) from single image and sparse range. In: CVPR, pp. 3353\u20133362","DOI":"10.1109\/CVPR.2019.00347"},{"key":"2470_CR96","doi-asserted-by":"crossref","unstructured":"Yu, Q., Du, H., Liu, C., & Yu, X. (2023). When 3d bounding-box meets sam: Point cloud instance segmentation with weak-and-noisy supervision. arXiv preprint arXiv:2309.00828","DOI":"10.1109\/WACV57701.2024.00368"},{"key":"2470_CR97","doi-asserted-by":"crossref","unstructured":"Yu, Z., Sheng, Z., Zhou, Z., Luo, L., Cao, S.Y., Gu, H., Zhang, H., & Shen, H.L. (2023). Aggregating feature point cloud for depth completion. In: ICCV, pp. 8732\u20138743","DOI":"10.1109\/ICCV51070.2023.00802"},{"key":"2470_CR98","doi-asserted-by":"crossref","unstructured":"Zeiler, M.D., & Fergus, R. (2014). Visualizing and understanding convolutional networks. In: ECCV, pp. 818\u2013833. Springer","DOI":"10.1007\/978-3-319-10590-1_53"},{"key":"2470_CR99","doi-asserted-by":"crossref","unstructured":"Zhang, H., Dana, K., Shi, J., Zhang, Z., Wang, X., Tyagi, A., & Agrawal, A. (2018). Context encoding for semantic segmentation. In: CVPR, pp. 7151\u20137160","DOI":"10.1109\/CVPR.2018.00747"},{"key":"2470_CR100","doi-asserted-by":"crossref","unstructured":"Zhang, K., & Liu, D. (2023). Customized segment anything model for medical image segmentation. arXiv preprint arXiv:2304.13785","DOI":"10.2139\/ssrn.4495221"},{"key":"2470_CR101","doi-asserted-by":"crossref","unstructured":"Zhang, Y., & Funkhouser, T. (2018). Deep depth completion of a single rgb-d image. In: CVPR, pp. 175\u2013185","DOI":"10.1109\/CVPR.2018.00026"},{"key":"2470_CR102","doi-asserted-by":"crossref","unstructured":"Zhang, Y., Guo, X., Poggi, M., Zhu, Z., Huang, G., & Mattoccia, S. (2023). Completionformer: Depth completion with convolutions and vision transformers. In: CVPR, pp. 18527\u201318536","DOI":"10.1109\/CVPR52729.2023.01777"},{"key":"2470_CR103","doi-asserted-by":"crossref","unstructured":"Zhang, Z., Cui, Z., Xu, C., Yan, Y., Sebe, N., & Yang, J. (2019). Pattern-affinitive propagation across depth, surface normal and semantic segmentation. In: CVPR, pp. 4106\u20134115","DOI":"10.1109\/CVPR.2019.00423"},{"key":"2470_CR104","doi-asserted-by":"crossref","unstructured":"Zhao, H., Shi, J., Qi, X., Wang, X., & Jia, J. (2017). Pyramid scene parsing network. In: CVPR, pp. 2881\u20132890","DOI":"10.1109\/CVPR.2017.660"},{"key":"2470_CR105","doi-asserted-by":"publisher","first-page":"5264","DOI":"10.1109\/TIP.2021.3079821","volume":"30","author":"S Zhao","year":"2021","unstructured":"Zhao, S., Gong, M., Fu, H., & Tao, D. (2021). Adaptive context-aware multi-modal network for depth completion. IEEE Transactions on Image Processing, 30, 5264\u20135276.","journal-title":"IEEE Transactions on Image Processing"},{"key":"2470_CR106","doi-asserted-by":"crossref","unstructured":"Zhou, W., Yan, X., Liao, Y., Lin, Y., Huang, J., Zhao, G., Cui, S., & Li, Z. (2023). Bev@ dc: Bird\u2019s-eye view assisted training for depth completion. In: CVPR, pp. 9233\u20139242","DOI":"10.1109\/CVPR52729.2023.00891"},{"key":"2470_CR107","doi-asserted-by":"crossref","unstructured":"Zhu, X., Hu, H., Lin, S., & Dai, J. (2019). Deformable convnets v2: More deformable, better results. In: CVPR, pp. 9308\u20139316","DOI":"10.1109\/CVPR.2019.00953"},{"key":"2470_CR108","doi-asserted-by":"crossref","unstructured":"Zioulis, N., Karakottas, A., Zarpalas, D., Alvarez, F., & Daras, P. (2019). Spherical view synthesis for self-supervised 360 depth estimation. In: 3DV, pp. 690\u2013699. IEEE","DOI":"10.1109\/3DV.2019.00081"}],"container-title":["International Journal of Computer Vision"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-025-02470-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11263-025-02470-y\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-025-02470-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,9,9]],"date-time":"2025-09-09T08:04:47Z","timestamp":1757405087000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11263-025-02470-y"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,5,23]]},"references-count":108,"journal-issue":{"issue":"9","published-print":{"date-parts":[[2025,9]]}},"alternative-id":["2470"],"URL":"https:\/\/doi.org\/10.1007\/s11263-025-02470-y","relation":{},"ISSN":["0920-5691","1573-1405"],"issn-type":[{"value":"0920-5691","type":"print"},{"value":"1573-1405","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,5,23]]},"assertion":[{"value":"22 December 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"2 May 2025","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"23 May 2025","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflicts of Interest"}}]}}