{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,4,2]],"date-time":"2026-04-02T12:14:22Z","timestamp":1775132062108,"version":"3.50.1"},"reference-count":43,"publisher":"Springer Science and Business Media LLC","issue":"2","license":[{"start":{"date-parts":[[2026,2,3]],"date-time":"2026-02-03T00:00:00Z","timestamp":1770076800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,2,3]],"date-time":"2026-02-03T00:00:00Z","timestamp":1770076800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"Scientific Research Fund Project of Yunnan Education Department","award":["2025Y0006"],"award-info":[{"award-number":["2025Y0006"]}]},{"name":"Yunnan Major Scientific and Technological Special Project","award":["202002AD080001"],"award-info":[{"award-number":["202002AD080001"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimedia Systems"],"published-print":{"date-parts":[[2026,4]]},"DOI":"10.1007\/s00530-025-02193-7","type":"journal-article","created":{"date-parts":[[2026,2,3]],"date-time":"2026-02-03T07:44:43Z","timestamp":1770104683000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Enhancing neural radiance fields with geometry-aware transformers and depth fusion"],"prefix":"10.1007","volume":"32","author":[{"given":"Mingqiang","family":"Xu","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Zhengyao","family":"Bai","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Chenghao","family":"Cao","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Zenan","family":"Xu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Han","family":"Ma","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Tao","family":"Lin","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Qiqin","family":"Huang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2026,2,3]]},"reference":[{"issue":"1","key":"2193_CR1","doi-asserted-by":"publisher","first-page":"99","DOI":"10.1145\/3503250","volume":"65","author":"B Mildenhall","year":"2021","unstructured":"Mildenhall, B., Srinivasan, P.P., Tancik, M., Barron, J.T., Ramamoorthi, R., Ng, R.: Nerf: representing scenes as neural radiance fields for view synthesis. Commun. ACM 65(1), 99\u2013106 (2021)","journal-title":"Commun. ACM"},{"key":"2193_CR2","unstructured":"Gao, K., Gao, Y., He, H., Lu, D., Xu, L., Li, J.: Nerf: Neural radiance field in 3d vision, a comprehensive review. arXiv:2210.00379 (2022)"},{"key":"2193_CR3","doi-asserted-by":"crossref","unstructured":"Tretschk, E., Tewari, A., Golyanik, V., Zollh\u00f6fer, M., Lassner, C., Theobalt, C.: Non-rigid neural radiance fields: Reconstruction and novel view synthesis of a dynamic scene from monocular video. In:\u00a0Proceedings of the IEEE\/CVF International Conference on Computer Vision. pp. 12959\u201312970 (2021)","DOI":"10.1109\/ICCV48922.2021.01272"},{"issue":"4","key":"2193_CR4","first-page":"2295","volume":"16","author":"K Singla","year":"2024","unstructured":"Singla, K., Nand, P.: Optimizing deep learning architectures for novel view synthesis: investigating the impact of nerf mlp parameters on complex scenes. Int. J. Inf. Technol. 16(4), 2295\u20132305 (2024)","journal-title":"Int. J. Inf. Technol."},{"key":"2193_CR5","unstructured":"Wang, Z., Wu, S., Xie, W., Chen, M., Prisacariu, V.A.: Nerf\u2013: Neural radiance fields without known camera parameters. arXiv:2102.07064 (2021)"},{"key":"2193_CR6","doi-asserted-by":"crossref","unstructured":"Wang, Q., Wang, Z., Genova, K., Srinivasan, P.P., Zhou, H., Barron, J.T., Martin-Brualla, R., Snavely, N., Funkhouser, T.: Ibrnet: learning multi-view image-based rendering. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 4690\u20134699 (2021)","DOI":"10.1109\/CVPR46437.2021.00466"},{"key":"2193_CR7","doi-asserted-by":"crossref","unstructured":"Yu, A., Ye, V., Tancik, M., Kanazawa, A.: pixelnerf: Neural radiance fields from one or few images. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 4578\u20134587 (2021)","DOI":"10.1109\/CVPR46437.2021.00455"},{"issue":"4","key":"2193_CR8","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1145\/3528223.3530127","volume":"41","author":"T M\u00fcller","year":"2022","unstructured":"M\u00fcller, T., Evans, A., Schied, C., Keller, A.: Instant neural graphics primitives with a multiresolution hash encoding. ACM Trans. Graphics (TOG) 41(4), 1\u201315 (2022)","journal-title":"ACM Trans. Graphics (TOG)"},{"key":"2193_CR9","doi-asserted-by":"crossref","unstructured":"Fridovich-Keil, S., Yu, A., Tancik, M., Chen, Q., Recht, B., Kanazawa, A.: Plenoxels: radiance fields without neural networks. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 5501\u20135510 (2022)","DOI":"10.1109\/CVPR52688.2022.00542"},{"key":"2193_CR10","doi-asserted-by":"crossref","unstructured":"Liu, Y., Peng, S., Liu, L., Wang, Q., Wang, P., Theobalt, C., Zhou, X., Wang, W.: Neural rays for occlusion-aware image-based rendering. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 7824\u20137833 (2022)","DOI":"10.1109\/CVPR52688.2022.00767"},{"key":"2193_CR11","unstructured":"Wang, P., Chen, X., Chen, T., Venugopalan, S., Wang, Z., et al.: Is attention all that nerf needs? arXiv:2207.13298 (2022)"},{"key":"2193_CR12","doi-asserted-by":"crossref","unstructured":"Guo, Y.-C., Kang, D., Bao, L., He, Y., Zhang, S.-H.: Nerfren: neural radiance fields with reflections. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 18409\u201318418 (2022)","DOI":"10.1109\/CVPR52688.2022.01786"},{"key":"2193_CR13","doi-asserted-by":"crossref","unstructured":"Zhang, J., Yao, Y., Li, S., Liu, J., Fang, T., McKinnon, D., Tsin, Y., Quan, L.: Neilf++: inter-reflectable light fields for geometry and material estimation. In:\u00a0Proceedings of the IEEE\/CVF International Conference on Computer Vision. pp. 3601\u20133610 (2023)","DOI":"10.1109\/ICCV51070.2023.00333"},{"issue":"3","key":"2193_CR14","doi-asserted-by":"publisher","first-page":"1623","DOI":"10.1109\/TPAMI.2020.3019967","volume":"44","author":"R Ranftl","year":"2020","unstructured":"Ranftl, R., Lasinger, K., Hafner, D., Schindler, K., Koltun, V.: Towards robust monocular depth estimation: mixing datasets for zero-shot cross-dataset transfer. IEEE Trans. Pattern Anal. Mach. Intell. 44(3), 1623\u20131637 (2020)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"2193_CR15","doi-asserted-by":"crossref","unstructured":"Ranftl, R., Bochkovskiy, A., Koltun, V.: Vision transformers for dense prediction. In:\u00a0Proceedings of the IEEE\/CVF International Conference on Computer Vision. pp. 12179\u201312188 (2021)","DOI":"10.1109\/ICCV48922.2021.01196"},{"key":"2193_CR16","doi-asserted-by":"crossref","unstructured":"Gui, M., Schusterbauer, J., Prestel, U., Ma, P., Kotovenko, D., Grebenkova, O., Baumann, S.A., Hu, V.T., Ommer, B.: Depthfm: fast monocular depth estimation with flow matching. arXiv:2403.13788 (2024)","DOI":"10.1609\/aaai.v39i3.32330"},{"key":"2193_CR17","doi-asserted-by":"crossref","unstructured":"Patni, S., Agarwal, A., Arora, C.: Ecodepth: Effective conditioning of diffusion models for monocular depth estimation. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 28285\u201328295 (2024)","DOI":"10.1109\/CVPR52733.2024.02672"},{"key":"2193_CR18","doi-asserted-by":"crossref","unstructured":"Wu, Z., Li, X., Peng, J., Lu, H., Cao, Z., Zhong, W.: Dof-nerf: depth-of-field meets neural radiance fields. In:\u00a0Proceedings of the 30th ACM International Conference on Multimedia. pp. 1718\u20131729 (2022)","DOI":"10.1145\/3503161.3548088"},{"key":"2193_CR19","doi-asserted-by":"crossref","unstructured":"Deng, K., Liu, A., Zhu, J.-Y., Ramanan, D.: Depth-supervised nerf: Fewer views and faster training for free. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 12882\u201312891 (2022)","DOI":"10.1109\/CVPR52688.2022.01254"},{"key":"2193_CR20","doi-asserted-by":"publisher","first-page":"75","DOI":"10.1109\/LSP.2023.3240370","volume":"30","author":"D Lee","year":"2023","unstructured":"Lee, D., Lee, K.M.: Dense depth-guided generalizable nerf. IEEE Sig. Process. Lett. 30, 75\u201379 (2023)","journal-title":"IEEE Sig. Process. Lett."},{"key":"2193_CR21","unstructured":"Vaswani, A.: Attention is all you need. Adv. Neural Inform. Process. Sys. (2017)"},{"key":"2193_CR22","unstructured":"Dosovitskiy, A.: An image is worth 16x16 words: transformers for image recognition at scale. arXiv:2010.11929 (2020)"},{"key":"2193_CR23","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2024.3445463","author":"C Chen","year":"2024","unstructured":"Chen, C., Wu, Y., Dai, Q., Zhou, H.-Y., Xu, M., Yang, S., Han, X., Yu, Y.: A survey on graph neural networks and graph transformers in computer vision: a task-oriented perspective. IEEE Trans. Pattern Anal. Mach. Intell. (2024). https:\/\/doi.org\/10.1109\/TPAMI.2024.3445463","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"issue":"12","key":"2193_CR24","doi-asserted-by":"publisher","first-page":"7682","DOI":"10.1109\/TPAMI.2024.3392941","volume":"46","author":"L Papa","year":"2024","unstructured":"Papa, L., Russo, P., Amerini, I., Zhou, L.: A survey on efficient vision transformers: algorithms, techniques, and performance benchmarking. IEEE Trans. Pattern Anal. Mach. Intell. 46(12), 7682\u20137700 (2024)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"2193_CR25","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2024.3434373","author":"X Li","year":"2024","unstructured":"Li, X., Ding, H., Yuan, H., Zhang, W., Pang, J., Cheng, G., Chen, K., Liu, Z., Loy, C.C.: Transformer-based visual segmentation: a survey. IEEE Trans. Pattern Anal. Mach. Intell. (2024). https:\/\/doi.org\/10.1109\/TPAMI.2024.3434373","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"2193_CR26","doi-asserted-by":"crossref","unstructured":"Li, Q., Wen, C., Fu, R.: Improving few-shot neural radiance field with image based rendering. In:\u00a02024 IEEE International Conference on Multimedia and Expo (ICME), pp. 1\u20136 (2024). IEEE","DOI":"10.1109\/ICME57554.2024.10687977"},{"key":"2193_CR27","doi-asserted-by":"crossref","unstructured":"Suhail, M., Esteves, C., Sigal, L., Makadia, A.: Light field neural rendering. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 8269\u20138279 (2022)","DOI":"10.1109\/CVPR52688.2022.00809"},{"key":"2193_CR28","doi-asserted-by":"crossref","unstructured":"Min, Z., Luo, Y., Yang, W., Wang, Y., Yang, Y.: Entangled view-epipolar information aggregation for generalizable neural radiance fields. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 4906\u20134916 (2024)","DOI":"10.1109\/CVPR52733.2024.00469"},{"key":"2193_CR29","doi-asserted-by":"crossref","unstructured":"Johari, M.M., Lepoittevin, Y., Fleuret, F.: Geonerf: generalizing nerf with geometry priors. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 18365\u201318375 (2022)","DOI":"10.1109\/CVPR52688.2022.01782"},{"key":"2193_CR30","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2025.3598711","author":"Y Chen","year":"2025","unstructured":"Chen, Y., Xu, H., Wu, Q., Zheng, C., Cham, T.-J., Cai, J.: Explicit correspondence matching for generalizable neural radiance fields. IEEE Trans. Pattern Anal. Mach. Intell. (2025). https:\/\/doi.org\/10.1109\/TPAMI.2025.3598711","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"2193_CR31","doi-asserted-by":"crossref","unstructured":"Roessle, B., Barron, J.T., Mildenhall, B., Srinivasan, P.P., Nie\u00dfner, M.: Dense depth priors for neural radiance fields from sparse input views. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 12892\u201312901 (2022)","DOI":"10.1109\/CVPR52688.2022.01255"},{"key":"2193_CR32","doi-asserted-by":"publisher","DOI":"10.1016\/j.displa.2025.102996","volume":"88","author":"Y Shi","year":"2025","unstructured":"Shi, Y., Rong, D., Chen, C., Ma, C., Ni, B., Zhang, W.: Darf: depth-aware generalizable neural radiance field. Displays 88, 102996 (2025)","journal-title":"Displays"},{"key":"2193_CR33","doi-asserted-by":"crossref","unstructured":"Wu, H., Hu, Z., Li, L., Zhang, Y., Fan, C., Yu, X.: Nefii: Inverse rendering for reflectance decomposition with near-field indirect illumination. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 4295\u20134304 (2023)","DOI":"10.1109\/CVPR52729.2023.00418"},{"key":"2193_CR34","doi-asserted-by":"crossref","unstructured":"Chen, A., Xu, Z., Geiger, A., Yu, J., Su, H.: Tensorf: tensorial radiance fields. In:\u00a0European Conference on Computer Vision. pp. 333\u2013350. Springer\u00a0(2022).","DOI":"10.1007\/978-3-031-19824-3_20"},{"key":"2193_CR35","unstructured":"Chen, X., Liu, J., Zhao, H., Zhou, G., Zhang, Y.-Q.: Nerrf: 3d reconstruction and view synthesis for transparent and specular objects with neural refractive-reflective fields. arXiv:2309.13039 (2023)"},{"key":"2193_CR36","doi-asserted-by":"crossref","unstructured":"Li, J., Li, Y., Sun, C., Wang, C., Xiang, J.: Spec-nerf: Multi-spectral neural radiance fields. In:\u00a0ICASSP 2024-2024 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), pp. 2485\u20132489. IEEE\u00a0(2024)","DOI":"10.1109\/ICASSP48485.2024.10446015"},{"key":"2193_CR37","doi-asserted-by":"crossref","unstructured":"Woo, S., Park, J., Lee, J.-Y., Kweon, I.S.: Cbam: convolutional block attention module. In:\u00a0Proceedings of the European Conference on Computer Vision (ECCV), pp. 3\u201319 (2018)","DOI":"10.1007\/978-3-030-01234-2_1"},{"key":"2193_CR38","doi-asserted-by":"crossref","unstructured":"Wang, Q., Wu, B., Zhu, P., Li, P., Zuo, W., Hu, Q.: Eca-net: efficient channel attention for deep convolutional neural networks. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 11534\u201311542 (2020)","DOI":"10.1109\/CVPR42600.2020.01155"},{"key":"2193_CR39","doi-asserted-by":"crossref","unstructured":"Yang, Z., Zhu, L., Wu, Y., Yang, Y.: Gated channel transformation for visual recognition. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 11794\u201311803 (2020)","DOI":"10.1109\/CVPR42600.2020.01181"},{"key":"2193_CR40","doi-asserted-by":"crossref","unstructured":"Barron, J.T., Mildenhall, B., Tancik, M., Hedman, P., Martin-Brualla, R., Srinivasan, P.P.: Mip-nerf: a multiscale representation for anti-aliasing neural radiance fields. In:\u00a0Proceedings of the IEEE\/CVF International Conference on Computer Vision. pp. 5855\u20135864 (2021)","DOI":"10.1109\/ICCV48922.2021.00580"},{"key":"2193_CR41","doi-asserted-by":"crossref","unstructured":"Barron, J.T., Mildenhall, B., Verbin, D., Srinivasan, P.P., Hedman, P.: Mip-nerf 360: Unbounded anti-aliased neural radiance fields. In:\u00a0Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp. 5470\u20135479 (2022)","DOI":"10.1109\/CVPR52688.2022.00539"},{"key":"2193_CR42","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2023.127063","volume":"568","author":"J Su","year":"2024","unstructured":"Su, J., Ahmed, M., Lu, Y., Pan, S., Bo, W., Liu, Y.: Roformer: enhanced transformer with rotary position embedding. Neurocomputing 568, 127063 (2024)","journal-title":"Neurocomputing"},{"key":"2193_CR43","unstructured":"Xu, M., Men, X., Wang, B., Zhang, Q., Lin, H., Han, X., : Base of rope bounds context length. In:\u00a0The Thirty-eighth Annual Conference on Neural Information Processing Systems\u00a0(2024)"}],"container-title":["Multimedia Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-025-02193-7.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00530-025-02193-7","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-025-02193-7.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,4,2]],"date-time":"2026-04-02T11:37:17Z","timestamp":1775129837000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00530-025-02193-7"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,2,3]]},"references-count":43,"journal-issue":{"issue":"2","published-print":{"date-parts":[[2026,4]]}},"alternative-id":["2193"],"URL":"https:\/\/doi.org\/10.1007\/s00530-025-02193-7","relation":{},"ISSN":["0942-4962","1432-1882"],"issn-type":[{"value":"0942-4962","type":"print"},{"value":"1432-1882","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,2,3]]},"assertion":[{"value":"8 September 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"26 December 2025","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"3 February 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no Conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"134"}}