{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,8,2]],"date-time":"2026-08-02T06:45:56Z","timestamp":1785653156799,"version":"3.56.0"},"publisher-location":"Cham","reference-count":38,"publisher":"Springer Nature Switzerland","isbn-type":[{"value":"9783032316653","type":"print"},{"value":"9783032316660","type":"electronic"}],"license":[{"start":{"date-parts":[[2026,8,3]],"date-time":"2026-08-03T00:00:00Z","timestamp":1785715200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,8,3]],"date-time":"2026-08-03T00:00:00Z","timestamp":1785715200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2027]]},"DOI":"10.1007\/978-3-032-31666-0_23","type":"book-chapter","created":{"date-parts":[[2026,8,2]],"date-time":"2026-08-02T05:46:09Z","timestamp":1785649569000},"page":"344-359","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["FuseDPT: Multi-scale and\u00a0Multi-projection Model for\u00a0Learning Depth in\u00a0360$$^\\circ $$C"],"prefix":"10.1007","author":[{"given":"Matheus","family":"Paula","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Nevrez","family":"Imamoglu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Guillaume","family":"Caron","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Antoine","family":"Andr\u00e9","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,8,3]]},"reference":[{"key":"23_CR1","doi-asserted-by":"crossref","unstructured":"Albanis, G., et al.: Pano3D: a holistic benchmark and a solid baseline for 360$$^{\\circ }$$ depth estimation. In: Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops (2021)","DOI":"10.1109\/CVPRW53098.2021.00413"},{"key":"23_CR2","doi-asserted-by":"publisher","first-page":"2322","DOI":"10.1109\/TRO.2025.3547266","volume":"41","author":"AN Andr\u00e9","year":"2025","unstructured":"Andr\u00e9, A.N., Morbidi, F., Caron, G.: UniphorM: a new uniform spherical image representation for robotic vision. IEEE Trans. Robot. 41, 2322\u20132339 (2025)","journal-title":"IEEE Trans. Robot."},{"key":"23_CR3","unstructured":"Armeni, I., Sax, S., Zamir, A.R., Savarese, S.: Joint 2D-3D-semantic data for indoor scene understanding. In: arXiv\/1702.01105 (2017)"},{"key":"23_CR4","doi-asserted-by":"crossref","unstructured":"Berenguel-Baeta, B., Bermudez-Cameo, J., Guerrero, J.J.: FreDSNet: joint monocular depth and semantic segmentation with fast Fourier convolutions from single panoramas. In: Proceedings of IEEE International Conference on Robotics and Automation, pp. 6080\u20136086 (2023)","DOI":"10.1109\/ICRA48891.2023.10161142"},{"key":"23_CR5","unstructured":"Bhat, S.F., Birkl, R., Wofk, D., Wonka, P., M\u00fcller, M.: ZoeDepth: zero-shot transfer by combining relative and metric depth (2023). arXiv\/2302.12288"},{"key":"23_CR6","doi-asserted-by":"crossref","unstructured":"Cao, Z., et al.: Panda: towards panoramic depth anything with unlabeled panoramas and m\u00f6bius spatial augmentation. In: 2025 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR), pp. 982\u2013992 (2025)","DOI":"10.1109\/CVPR52734.2025.00100"},{"key":"23_CR7","doi-asserted-by":"crossref","unstructured":"Chang, A., et al.: MatterPort3D: learning from RGB-D data in indoor environments. In: Proceedings of International Conference on 3D Vision, pp. 667\u2013676 (2017)","DOI":"10.1109\/3DV.2017.00081"},{"key":"23_CR8","doi-asserted-by":"crossref","unstructured":"Chappellet, K., Caron, G., Kanehiro, F., Sakurada, K., Kheddar, A.: Benchmarking cameras for open VSLAM indoors. In: Proc. of Int. Conf. on Pattern Recognition, pp. 4857\u20134864 (2021)","DOI":"10.1109\/ICPR48806.2021.9413278"},{"issue":"23","key":"23_CR9","doi-asserted-by":"publisher","first-page":"26912","DOI":"10.1109\/JSEN.2021.3120753","volume":"21","author":"Z Cheng","year":"2021","unstructured":"Cheng, Z., Zhang, Y., Tang, C.: Swin-Depth: using transformers and multi-scale fusion for monocular-based depth estimation. IEEE Sens. J. 21(23), 26912\u201326920 (2021)","journal-title":"IEEE Sens. J."},{"key":"23_CR10","doi-asserted-by":"publisher","first-page":"133","DOI":"10.1016\/j.comcom.2021.06.029","volume":"177","author":"F Chiariotti","year":"2021","unstructured":"Chiariotti, F.: A survey on 360-degree video: coding, quality of experience and streaming. Comput. Commun. 177, 133\u2013155 (2021)","journal-title":"Comput. Commun."},{"issue":"3","key":"23_CR11","doi-asserted-by":"publisher","first-page":"577","DOI":"10.1145\/1073204.1073232","volume":"24","author":"D Hoiem","year":"2005","unstructured":"Hoiem, D., Efros, A.A., Hebert, M.: Automatic photo pop-up. ACM Trans. Graph. 24(3), 577\u2013584 (2005)","journal-title":"ACM Trans. Graph."},{"key":"23_CR12","doi-asserted-by":"crossref","unstructured":"Jiang, H., Sheng, Z., Zhu, S., Dong, Z., Huang, R.: UniFuse: unidirectional fusion for 360$$^{\\circ }$$ panorama depth estimation. IEEE Robot. Autom. Lett. 6(2), 1519\u20131526 (2021)","DOI":"10.1109\/LRA.2021.3058957"},{"issue":"11","key":"23_CR13","doi-asserted-by":"publisher","first-page":"2144","DOI":"10.1109\/TPAMI.2014.2316835","volume":"36","author":"K Karsch","year":"2014","unstructured":"Karsch, K., Liu, C., Kang, S.B.: Depth Transfer: depth extraction from video using non-parametric sampling. IEEE Trans. Pattern Anal. Mach. Intell. 36(11), 2144\u20132158 (2014). https:\/\/doi.org\/10.1109\/TPAMI.2014.2316835","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"issue":"10s","key":"23_CR14","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1145\/3505244","volume":"54","author":"S Khan","year":"2022","unstructured":"Khan, S., Naseer, M., Hayat, M., Zamir, S.W., Khan, F.S., Shah, M.: Transformers in vision: a survey. ACM Comput. Surv. 54(10s), 1\u201341 (2022)","journal-title":"ACM Comput. Surv."},{"key":"23_CR15","unstructured":"de\u00a0La\u00a0Garanderie, G.P., Abarghouei, A.A., Breckon, T.P.: Eliminating the blind spot: adapting 3D object detection and monocular depth estimation to 360$$^{\\circ } $$ panoramic imagery. In: Proceedings of the European Conference on Computer Vision, pp. 789\u2013807 (2018)"},{"key":"23_CR16","doi-asserted-by":"crossref","unstructured":"Laina, I., Rupprecht, C., Belagiannis, V., Tombari, F., Navab, N.: Deeper depth prediction with fully convolutional residual networks. In: Proceedings of International Conference on 3D Vision, pp. 239\u2013248 (2016)","DOI":"10.1109\/3DV.2016.32"},{"issue":"2","key":"23_CR17","doi-asserted-by":"publisher","first-page":"1053","DOI":"10.1109\/LRA.2023.3234820","volume":"8","author":"M Li","year":"2023","unstructured":"Li, M., Wang, S., Yuan, W., Shen, W., Sheng, Z., Dong, Z.: $$\\cal{S} ^{2}$$net: accurate panorama depth estimation on spherical surface. IEEE Robot. Autom. Lett. 8(2), 1053\u20131060 (2023)","journal-title":"IEEE Robot. Autom. Lett."},{"key":"23_CR18","doi-asserted-by":"crossref","unstructured":"Li, R., Xian, K., Shen, C., Cao, Z., Lu, H., Hang, L.: Deep attention-based classification network for robust depth prediction. In: Proceedings of Asian Conferences on Computer Vision, pp. 663\u2013678 (2019)","DOI":"10.1007\/978-3-030-20870-7_41"},{"key":"23_CR19","doi-asserted-by":"crossref","unstructured":"Li, Y., Guo, Y., Yan, Z., Huang, X., Duan, Y., Ren, L.: OmniFusion: 360 monocular depth estimation via geometry-aware fusion. In: Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 2791\u20132800 (2022)","DOI":"10.1109\/CVPR52688.2022.00282"},{"key":"23_CR20","doi-asserted-by":"crossref","unstructured":"Li, Z., Bhat, S.F., Wonka, P.: PatchFusion: an end-to-end tile-based framework for high-resolution monocular metric depth estimation. In: Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 10016\u201310025 (2024)","DOI":"10.1109\/CVPR52733.2024.00955"},{"key":"23_CR21","doi-asserted-by":"crossref","unstructured":"Liu, F., Shen, C., Lin, G.: Deep convolutional neural fields for depth estimation from a single image. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, pp. 5162\u20135170 (2015)","DOI":"10.1109\/CVPR.2015.7299152"},{"key":"23_CR22","doi-asserted-by":"crossref","unstructured":"Pintore, G., Agus, M., Almansa, E., Schneider, J., Gobbetti, E.: SliceNet: deep dense depth estimation from a single indoor panorama using a slice-based representation. In: Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 11531\u201311540 (2021)","DOI":"10.1109\/CVPR46437.2021.01137"},{"key":"23_CR23","doi-asserted-by":"crossref","unstructured":"Ranftl, R., Bochkovskiy, A., Koltun, V.: Vision transformers for dense prediction. In: Proceedings of IEEE\/CVF International Conference on Computer Vision, pp. 12159\u201312168 (2021)","DOI":"10.1109\/ICCV48922.2021.01196"},{"key":"23_CR24","doi-asserted-by":"crossref","unstructured":"Ranftl, R., Lasinger, K., Hafner, D., Schindler, K., Koltun, V.: Towards robust monocular depth estimation: mixing datasets for zero-shot cross-dataset transfer. IEEE Trans. Pattern Anal. Mach. Intell. 44(3), 1623\u20131637 (2022)","DOI":"10.1109\/TPAMI.2020.3019967"},{"key":"23_CR25","doi-asserted-by":"crossref","unstructured":"Rey\u2013Area, M., Yuan, M., Richardt, C.: 360MonoDepth: high-resolution 360$$^{\\circ } $$ monocular depth estimation. In: Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 3752\u20133762 (2022)","DOI":"10.1109\/CVPR52688.2022.00374"},{"key":"23_CR26","doi-asserted-by":"crossref","unstructured":"Saxena, A., Sun, M., Ng, A.Y.: Make3D: learning 3D scene structure from a single still image. IEEE Trans. Pattern Anal. Mach. Intel. 31(5), 824\u2013840 (2009)","DOI":"10.1109\/TPAMI.2008.132"},{"key":"23_CR27","doi-asserted-by":"crossref","unstructured":"Shen, Z., Lin, C., Liao, K., Nie, L., Zheng, Z., Zhao, Y.: PanoFormer: panorama transformer for indoor 360 degree depth estimation. In: Computer Vision - ECCV 2022, pp. 195\u2013211. Springer Nature Switzerland (2022)","DOI":"10.1007\/978-3-031-19769-7_12"},{"key":"23_CR28","unstructured":"Straub, J., et al.: The Replica dataset: a digital replica of indoor spaces. In: arXiv\/1906.05797 (2019)"},{"key":"23_CR29","doi-asserted-by":"crossref","unstructured":"Sun, C., Sun, M., Chen, H.T.: HoHoNet: 360 indoor holistic understanding with latent horizontal features. In: Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 2573\u20132582 (2021)","DOI":"10.1109\/CVPR46437.2021.00260"},{"key":"23_CR30","doi-asserted-by":"crossref","unstructured":"Wang, F.E., Yeh, Y.H., Sun, M., Chiu, W.C., Tsai, Y.H.: BiFuse: monocular 360 depth estimation via bi-projection fusion. In: Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 459\u2013468 (2020)","DOI":"10.1109\/CVPR42600.2020.00054"},{"key":"23_CR31","doi-asserted-by":"crossref","unstructured":"Wang, N.H., Liu, Y.L.: Depth anywhere: enhancing 360 monocular depth estimation via perspective distillation and unlabeled data augmentation. Adv. Neural. Inf. Process. Syst. 37 (2024)","DOI":"10.52202\/079017-4056"},{"key":"23_CR32","doi-asserted-by":"crossref","unstructured":"Yang, G., Tang, H., Ding, M., Sebe, N., Ricci, E.: Transformer-based attention networks for continuous pixel-wise prediction. In: Proceedings of IEEE\/CVF International Conference on Computer Vision, pp. 16249\u201316259 (2021)","DOI":"10.1109\/ICCV48922.2021.01596"},{"key":"23_CR33","unstructured":"Yang, J., An, L., Dixit, A., Koo, J., Park, S.I.: Depth estimation with simplified transformer. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition Workshops (2022)"},{"key":"23_CR34","doi-asserted-by":"crossref","unstructured":"Yang, L., Kang, B., Huang, Z., Xu, X., Feng, J., Zhao, H.: Depth anything: unleashing the power of large-scale unlabeled data. In: Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 10371\u201310381 (2024)","DOI":"10.1109\/CVPR52733.2024.00987"},{"key":"23_CR35","doi-asserted-by":"crossref","unstructured":"Yun, I., Lee, H., Rhee, C.: Improving 360 monocular depth estimation via non-local dense prediction transformer and joint supervised and self-supervised learning. In: Proceedings of the AAAI Conference on Artificial Intelligence, pp. 3224\u20133233 (2022)","DOI":"10.1609\/aaai.v36i3.20231"},{"key":"23_CR36","doi-asserted-by":"crossref","unstructured":"Yun, I., Shin, C., Lee, H., Lee, H.J., Rhee, C.E.: EGformer: equirectangular geometry-biased transformer for 360 depth estimation. In: Proceedings of IEEE\/CVF International Conference on Computer Vision, pp. 6078\u20136089 (2023)","DOI":"10.1109\/ICCV51070.2023.00561"},{"key":"23_CR37","doi-asserted-by":"crossref","unstructured":"Zhang, Z., Rebecq, H., Forster, C., Scaramuzza, D.: Benefit of large field-of-view cameras for visual odometry. In: Proceedings of IEEE International Conference on Robotics and Automation, pp. 801\u2013808 (2016)","DOI":"10.1109\/ICRA.2016.7487210"},{"key":"23_CR38","doi-asserted-by":"crossref","unstructured":"Zioulis, N., Karakottas, A., Zarpalas, D., Daras, P.: OmniDepth: dense depth estimation for indoors spherical panoramas. In: Proceedings of the European Conference on Computer Vision, pp. 448\u2013465 (2018)","DOI":"10.1007\/978-3-030-01231-1_28"}],"container-title":["Lecture Notes in Computer Science","Pattern Recognition"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/978-3-032-31666-0_23","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,8,2]],"date-time":"2026-08-02T05:46:11Z","timestamp":1785649571000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/978-3-032-31666-0_23"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,8,3]]},"ISBN":["9783032316653","9783032316660"],"references-count":38,"URL":"https:\/\/doi.org\/10.1007\/978-3-032-31666-0_23","relation":{},"ISSN":["0302-9743","1611-3349"],"issn-type":[{"value":"0302-9743","type":"print"},{"value":"1611-3349","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,8,3]]},"assertion":[{"value":"3 August 2026","order":1,"name":"first_online","label":"First Online","group":{"name":"ChapterHistory","label":"Chapter History"}},{"value":"The authors have no competing interests to declare that are relevant to the content of this article.","order":1,"name":"Ethics","label":"Disclosure of Interests","group":{"name":"EthicsHeading","label":"Ethics"}},{"value":"ICPR","order":1,"name":"conference_acronym","label":"Conference Acronym","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"International Conference on Pattern Recognition","order":2,"name":"conference_name","label":"Conference Name","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Lyon","order":3,"name":"conference_city","label":"Conference City","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"France","order":4,"name":"conference_country","label":"Conference Country","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"2026","order":5,"name":"conference_year","label":"Conference Year","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"17 August 2026","order":7,"name":"conference_start_date","label":"Conference Start Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"22 August 2026","order":8,"name":"conference_end_date","label":"Conference End Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"28","order":9,"name":"conference_number","label":"Conference Number","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"icpr2026","order":10,"name":"conference_id","label":"Conference ID","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"https:\/\/icpr2026.org\/","order":11,"name":"conference_url","label":"Conference URL","group":{"name":"ConferenceInfo","label":"Conference Information"}}]}}