{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,3,13]],"date-time":"2026-03-13T17:13:41Z","timestamp":1773422021431,"version":"3.50.1"},"reference-count":62,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2026,2,1]],"date-time":"2026-02-01T00:00:00Z","timestamp":1769904000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,2,1]],"date-time":"2026-02-01T00:00:00Z","timestamp":1769904000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["4242016"],"award-info":[{"award-number":["4242016"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100004826","name":"Natural Science Foundation of Beijing Municipality","doi-asserted-by":"publisher","award":["62572018"],"award-info":[{"award-number":["62572018"]}],"id":[{"id":"10.13039\/501100004826","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Vis Comput"],"published-print":{"date-parts":[[2026,2]]},"DOI":"10.1007\/s00371-025-04349-y","type":"journal-article","created":{"date-parts":[[2026,2,2]],"date-time":"2026-02-02T06:12:30Z","timestamp":1770012750000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Language-guided semantic editing in single-view 3D reconstruction"],"prefix":"10.1007","volume":"42","author":[{"given":"Maoyang","family":"Xu","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Guanglei","family":"Qi","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Nana","family":"He","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Fujiao","family":"Ju","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2026,2,2]]},"reference":[{"key":"4349_CR1","doi-asserted-by":"publisher","unstructured":"Choy, C.B., Xu, D., Gwak, J., Chen, K., Savarese, S.: 3d-r2n2: A unified approach for single and multi-view 3d object reconstruction. In: Computer Vision\u2013ECCV 2016: 14th European Conference, Amsterdam, The Netherlands, October 11\u201314, 2016, Proceedings, Part V. Lecture Notes in Computer Science, vol. 9912, pp. 628\u2013644. Springer, Amsterdam (2016). https:\/\/doi.org\/10.1007\/978-3-319-46484-8_38","DOI":"10.1007\/978-3-319-46484-8_38"},{"key":"4349_CR2","doi-asserted-by":"publisher","unstructured":"Barron, J.T., Mildenhall, B., Tancik, M., Hedman, P., Martin-Brualla, R., Srinivasan, P.P.: MIP-NERF: a multiscale representation for anti-aliasing neural radiance fields. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 5855\u20135864. IEEE, Montreal, QC (2021). https:\/\/doi.org\/10.1109\/ICCV48922.2021.00582","DOI":"10.1109\/ICCV48922.2021.00582"},{"issue":"1","key":"4349_CR3","doi-asserted-by":"publisher","first-page":"71","DOI":"10.1162\/jocn.1991.3.1.71","volume":"3","author":"M Turk","year":"1991","unstructured":"Turk, M., Pentland, A.: Eigenfaces for recognition. J. Cogn. Neurosci. 3(1), 71\u201386 (1991). https:\/\/doi.org\/10.1162\/jocn.1991.3.1.71","journal-title":"J. Cogn. Neurosci."},{"key":"4349_CR4","unstructured":"Li, Y., Bu, R., Sun, M., Wu, W., Di, X., Chen, B.: POINTCNN: convolution on x-transformed points. In: Advances in Neural Information Processing Systems, vol. 31. Montreal, Canada, pp. 88\u201398 (2018)"},{"key":"4349_CR5","doi-asserted-by":"publisher","unstructured":"Zhou, Y., Tuzel, O.: Voxelnet: End-to-end learning for point cloud based 3D object detection. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 4490\u20134499. IEEE, Salt Lake City, UT (2018). https:\/\/doi.org\/10.1109\/CVPR.2018.00472","DOI":"10.1109\/CVPR.2018.00472"},{"key":"4349_CR6","doi-asserted-by":"publisher","unstructured":"Fan, H., Su, H., Guibas, L.J.: A point set generation network for 3d object reconstruction from a single image. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 2463\u20132471. IEEE, Honolulu, HI (2017). https:\/\/doi.org\/10.1109\/CVPR.2017.264","DOI":"10.1109\/CVPR.2017.264"},{"issue":"10","key":"4349_CR7","doi-asserted-by":"publisher","first-page":"3600","DOI":"10.1109\/TPAMI.2020.2984232","volume":"43","author":"N Wang","year":"2021","unstructured":"Wang, N., Zhang, Y., Li, Z., Fu, Y., Liu, W., Jiang, Y.-G.: Pixel2Mesh: 3D mesh model generation via image guided deformation. IEEE Trans. Pattern Anal. Mach. Intell. 43(10), 3600\u20133613 (2021). https:\/\/doi.org\/10.1109\/TPAMI.2020.2984232","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"4349_CR8","doi-asserted-by":"publisher","unstructured":"Mildenhall, B., Srinivasan, P.P., Tancik, M., Barron, J.T., Ramamoorthi, R., Ng, R.: Nerf: representing scenes as neural radiance fields for view synthesis. In: Computer Vision\u2013ECCV 2020: 16th European Conference, Glasgow, UK, August 23\u201328, 2020, Proceedings, Part I. Lecture Notes in Computer Science, vol. 12346, pp. 405\u2013421. Springer, Glasgow (2020). https:\/\/doi.org\/10.1007\/978-3-030-58452-8_24","DOI":"10.1007\/978-3-030-58452-8_24"},{"key":"4349_CR9","doi-asserted-by":"publisher","unstructured":"Zhang, K., Riegler, G., Snavely, N., Koltun, V.: Nerf++: analyzing and improving neural radiance fields. (2020) https:\/\/doi.org\/10.48550\/arXiv.2010.07492","DOI":"10.48550\/arXiv.2010.07492"},{"issue":"2","key":"4349_CR10","doi-asserted-by":"publisher","first-page":"2166","DOI":"10.1109\/TPAMI.2022.3169735","volume":"45","author":"C Wen","year":"2023","unstructured":"Wen, C., Zhang, Y., Cao, C., Li, Z., Xue, X., Fu, Y.: Pixel2Mesh++: 3D mesh generation and refinement from multi-view images. IEEE Trans. Pattern Anal. Mach. Intell. 45(2), 2166\u20132180 (2023). https:\/\/doi.org\/10.1109\/TPAMI.2022.3169735","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"4349_CR11","unstructured":"Fu, R., Zhan, X., Chen, Y., Ritchie, D., Sridhar, S.: ShapeCrafter: a recursive text-conditioned 3D shape generation model. In: Advances in Neural Information Processing Systems, vol. 35. New Orleans, LA, pp. 14707\u201314721 (2022)"},{"key":"4349_CR12","doi-asserted-by":"publisher","unstructured":"He, K., Zhang, X., Ren, S., Sun, J.: Deep residual learning for image recognition. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 770\u2013778. IEEE, Las Vegas, NV (2016). https:\/\/doi.org\/10.1109\/CVPR.2016.90","DOI":"10.1109\/CVPR.2016.90"},{"key":"4349_CR13","unstructured":"Xu, K., Ba, J., Kiros, R., Cho, K., Courville, A., Salakhutdinov, R., Zemel, R., Bengio, Y.: Show, attend and tell: neural image caption generation with visual attention. In: Proceedings of the 32nd International Conference on Machine Learning, pp. 2048\u20132057. PMLR, Lille (2015)"},{"key":"4349_CR14","doi-asserted-by":"publisher","unstructured":"Wang, C., Chai, M., He, M., Chen, D., Liao, J.: CLIP-NeRF: text- and-image driven manipulation of neural radiance fields. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 3825\u20133834. IEEE, New Orleans, LA (2022). https:\/\/doi.org\/10.1109\/CVPR52688.2022.00381","DOI":"10.1109\/CVPR52688.2022.00381"},{"key":"4349_CR15","doi-asserted-by":"publisher","unstructured":"Michel, O., Bar-On, R., Liu, R., Benaim, S., Hanocka, R.: Text2mesh: text-driven neural stylization for meshes. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 13482\u201313492. IEEE, New Orleans, LA (2022). https:\/\/doi.org\/10.1109\/CVPR52688.2022.01313","DOI":"10.1109\/CVPR52688.2022.01313"},{"key":"4349_CR16","unstructured":"Poole, B., Jain, A., Barron, J.T., Mildenhall, B.: Dreamfusion: text-to-3D using 2D diffusion. In: International Conference on Learning Representations, Kigali, Rwanda (2023)"},{"key":"4349_CR17","doi-asserted-by":"publisher","unstructured":"Lin, C.-H., Gao, J., Tang, L., Takikawa, T., Zeng, X., Huang, X., Munkberg, J., Fidler, S., Kautz, J., Liu, T.-Y.: Magic3d: High-resolution text-to-3d content creation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 300\u2013309. IEEE, Vancouver, BC (2023). https:\/\/doi.org\/10.1109\/CVPR52729.2023.00037","DOI":"10.1109\/CVPR52729.2023.00037"},{"key":"4349_CR18","doi-asserted-by":"publisher","unstructured":"Zhou, S., Tang, T., Zhou, B.: CADParser: a learning approach of sequence modeling for b-rep cad. In: Proceedings of the Thirty-Second International Joint Conference on Artificial Intelligence, pp. 1797\u20131805. IJCAI, Macao (2023). https:\/\/doi.org\/10.24963\/ijcai.2023\/200","DOI":"10.24963\/ijcai.2023\/200"},{"issue":"4","key":"4349_CR19","doi-asserted-by":"publisher","first-page":"68","DOI":"10.1007\/s00138-024-01544-0","volume":"35","author":"D Sharma","year":"2024","unstructured":"Sharma, D., Dhiman, C., Kumar, D.: Fdt-dr$$^2$$t: a unified dense radiology report generation transformer framework for X-ray images. Mach. Vis. Appl. 35(4), 68 (2024). https:\/\/doi.org\/10.1007\/s00138-024-01544-0","journal-title":"Mach. Vis. Appl."},{"key":"4349_CR20","doi-asserted-by":"publisher","first-page":"119773","DOI":"10.1016\/j.eswa.2023.119773","volume":"221","author":"D Sharma","year":"2023","unstructured":"Sharma, D., Dhiman, C., Kumar, D.: Evolution of visual data captioning methods, datasets, and evaluation metrics: a comprehensive survey. Expert Syst. Appl. 221, 119773 (2023). https:\/\/doi.org\/10.1016\/j.eswa.2023.119773","journal-title":"Expert Syst. Appl."},{"key":"4349_CR21","doi-asserted-by":"publisher","unstructured":"Radford, A., Kim, J.W., Hallacy, C., Ramesh, A., Goh, G., Agarwal, S., Sastry, G., Askell, A., Mishkin, P., Clark, J., Krueger, G., Sutskever, I.: Learning transferable visual models from natural language supervision (2021) https:\/\/doi.org\/10.48550\/arXiv.2103.00020","DOI":"10.48550\/arXiv.2103.00020"},{"key":"4349_CR22","doi-asserted-by":"publisher","DOI":"10.1109\/TNNLS.2025.3633075","author":"Y Liu","year":"2025","unstructured":"Liu, Y., Dai, D., Xia, S., Wang, G.: FDSRM: a feature-driven style-agnostic foundation model for sketch-less facial image retrieval. IEEE Trans. Neural Netw. Learn. Syst. (2025). https:\/\/doi.org\/10.1109\/TNNLS.2025.3633075","journal-title":"IEEE Trans. Neural Netw. Learn. Syst."},{"key":"4349_CR23","doi-asserted-by":"publisher","unstructured":"Liu, Y., Dai, D., Hou, X., Zhao, S., Wang, G.: From sparse to complete: semantic understanding based on stroke evolution in on-the-fly sketch-based image retrieval. In: Proceedings of the Thirty-Fourth International Joint Conference on Artificial Intelligence. IJCAI\u201925 (2025). https:\/\/doi.org\/10.24963\/ijcai.2025\/183","DOI":"10.24963\/ijcai.2025\/183"},{"key":"4349_CR24","doi-asserted-by":"publisher","unstructured":"Sanghi, A., Chu, H., Lambourne, J.G., Wang, Y., Cheng, C.-Y., Fouhey, D.: Clip-forge: towards zero-shot text-to-shape generation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 18582\u201318592. IEEE, New Orleans, LA (2022). https:\/\/doi.org\/10.1109\/CVPR52688.2022.01805","DOI":"10.1109\/CVPR52688.2022.01805"},{"key":"4349_CR25","doi-asserted-by":"publisher","unstructured":"Qi, C.R., Su, H., Mo, K., Guibas, L.J.: PointNet: deep learning on point sets for 3d classification and segmentation. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 77\u201385. IEEE, Honolulu, HI, USA (2017). https:\/\/doi.org\/10.1109\/CVPR.2017.16","DOI":"10.1109\/CVPR.2017.16"},{"key":"4349_CR26","doi-asserted-by":"publisher","unstructured":"Qi, C.R., Yi, L., Su, H., Guibas, L.J.: Pointnet++: deep hierarchical feature learning on point sets in a metric space (2017) https:\/\/doi.org\/10.48550\/arXiv.1706.02413","DOI":"10.48550\/arXiv.1706.02413"},{"key":"4349_CR27","doi-asserted-by":"publisher","unstructured":"Yu, X., Tang, L., Rao, Y., Huang, T., Zhou, J., Lu, J.: Point-BERT: pre-training 3D point cloud transformers with masked point modeling. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 19313\u201319322. IEEE, New Orleans, LA (2022). https:\/\/doi.org\/10.1109\/CVPR52688.2022.01871","DOI":"10.1109\/CVPR52688.2022.01871"},{"key":"4349_CR28","doi-asserted-by":"publisher","unstructured":"Yu, X., Rao, Y., Wang, Z., Liu, Z., Lu, J., Zhou, J.: PoinTr: diverse point cloud completion with geometry-aware transformers (2021) https:\/\/doi.org\/10.48550\/arXiv.2108.08839","DOI":"10.48550\/arXiv.2108.08839"},{"key":"4349_CR29","doi-asserted-by":"publisher","unstructured":"Li, X., Bi, L., Kim, J., Li, T., Li, P., Tian, Y., Sheng, B., Feng, D.D.: Malocclusion treatment planning via pointnet based spatial transformation network. In: Medical Image Computing and Computer Assisted Intervention\u2014MICCAI 2020: 23rd International Conference, Lima, Peru, October 4\u20138, 2020, Proceedings, Part VI, pp. 105\u2013114. Springer, Lima (2020). https:\/\/doi.org\/10.1007\/978-3-030-59716-0_11","DOI":"10.1007\/978-3-030-59716-0_11"},{"key":"4349_CR30","doi-asserted-by":"publisher","first-page":"5062","DOI":"10.1109\/TMM.2025.3543001","volume":"27","author":"J Zhu","year":"2025","unstructured":"Zhu, J., Yan, J., Huang, J., Nie, Y., Sheng, B., Lee, T.-Y.: SGG-Nets: generic rotation-invariant plugin networks for point cloud analysis. IEEE Trans. Multimedia 27, 5062\u20135076 (2025). https:\/\/doi.org\/10.1109\/TMM.2025.3543001","journal-title":"IEEE Trans. Multimedia"},{"key":"4349_CR31","doi-asserted-by":"publisher","DOI":"10.1007\/s00371-025-04168-1","author":"Y Luo","year":"2025","unstructured":"Luo, Y., Chen, J., Yao, Z.: Efficient semantic segmentation across domains: enhancing generalization with multi-scale and simple attention modules. Vis. Comput. (2025). https:\/\/doi.org\/10.1007\/s00371-025-04168-1","journal-title":"Vis. Comput."},{"issue":"6","key":"4349_CR32","doi-asserted-by":"publisher","first-page":"490","DOI":"10.1016\/j.vrih.2023.06.004","volume":"5","author":"S Elezovikj","year":"2023","unstructured":"Elezovikj, S., Jia, J., Tan, C.C., Ling, H.: Partlabeling: a label management framework in 3D space. Virtual Real. Intell. Hardware 5(6), 490\u2013508 (2023). https:\/\/doi.org\/10.1016\/j.vrih.2023.06.004","journal-title":"Virtual Real. Intell. Hardware"},{"key":"4349_CR33","doi-asserted-by":"publisher","first-page":"104600","DOI":"10.1016\/j.jvcir.2025.104600","volume":"113","author":"D Sharma","year":"2025","unstructured":"Sharma, D., Dhiman, C., Kumar, D.: UnMA-CapSumT: unified and multi-head attention-driven caption summarization transformer. J. Visual Commun. Image Represent 113, 104600 (2025). https:\/\/doi.org\/10.1016\/j.jvcir.2025.104600","journal-title":"J. Visual Commun. Image Represent"},{"key":"4349_CR34","doi-asserted-by":"publisher","unstructured":"Kato, H., Ushiku, Y., Harada, T.: Neural 3D mesh renderer. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 3907\u20133916. IEEE, Salt Lake City, UT (2018). https:\/\/doi.org\/10.1109\/CVPR.2018.00411","DOI":"10.1109\/CVPR.2018.00411"},{"key":"4349_CR35","doi-asserted-by":"publisher","unstructured":"Qin, X., Yang, X., Zheng, H.: Implicit surface Boolean operations based cut-and-paste algorithm for mesh models. In: Advances in Artificial Reality and Tele-Existence. Lecture Notes in Computer Science, vol. 4282, pp. 709\u2013720. Springer, Hangzhou (2006). https:\/\/doi.org\/10.1007\/11941354_87","DOI":"10.1007\/11941354_87"},{"key":"4349_CR36","doi-asserted-by":"publisher","unstructured":"Turk, G., O\u2019Brien, J.F.: Shape transformation using variational implicit functions. In: Proceedings of the 26th ACM SIGGRAPH Conference on Computer Graphics and Interactive Techniques, pp. 335\u2013342. Association for Computing Machinery, Los Angeles, CA (1999). https:\/\/doi.org\/10.1145\/311535.311580","DOI":"10.1145\/311535.311580"},{"issue":"7","key":"4349_CR37","doi-asserted-by":"publisher","first-page":"2303","DOI":"10.1007\/s00371-021-02112-7","volume":"38","author":"T Liu","year":"2022","unstructured":"Liu, T., Cai, Y., Zheng, J., Thalmann, N.M.: BEACon: a boundary embedded attentional convolution network for point cloud instance segmentation. Vis. Comput. 38(7), 2303\u20132313 (2022). https:\/\/doi.org\/10.1007\/s00371-021-02112-7","journal-title":"Vis. Comput."},{"key":"4349_CR38","doi-asserted-by":"publisher","unstructured":"Sharma, D., Dhiman, C., Kumar, D.: Automated image caption generation framework using adaptive attention and Bi-LSTM. In: 2022 IEEE Delhi Section Conference (DELCON), New Delhi, India, pp. 1\u20135 (2022). https:\/\/doi.org\/10.1109\/DELCON54057.2022.9752859","DOI":"10.1109\/DELCON54057.2022.9752859"},{"key":"4349_CR39","unstructured":"Dosovitskiy, A., Beyer, L., Kolesnikov, A., Weissenborn, D., Zhai, X., Unterthiner, T., Dehghani, M., Minderer, M., Heigold, G., Gelly, S., Uszkoreit, J., Houlsby, N.: An image is worth $$16\\times 16$$ words: transformers for image recognition at scale. In: International Conference on Learning Representations, Virtual (2021)"},{"key":"4349_CR40","doi-asserted-by":"publisher","unstructured":"Sharma, D., Dingliwal, R., Dhiman, C., Kumar, D.: Lightweight transformer with GRU integrated decoder for image captioning. In: 2022 16th International Conference on Signal Image Technology & Internet Based Systems (SITIS), pp. 434\u2013438. Dijon, France (2022). https:\/\/doi.org\/10.1109\/SITIS57111.2022.00072","DOI":"10.1109\/SITIS57111.2022.00072"},{"key":"4349_CR41","unstructured":"Vaswani, A., Shazeer, N., Parmar, N., Uszkoreit, J., Jones, L., Gomez, A.N., Kaiser, \u0141., Polosukhin, I.: Attention is all you need. In: Advances in Neural Information Processing Systems, vol. 30. Long Beach, CA (2017)"},{"issue":"2","key":"4349_CR42","doi-asserted-by":"publisher","first-page":"4219","DOI":"10.1007\/s11042-023-15291-3","volume":"83","author":"D Sharma","year":"2024","unstructured":"Sharma, D., Dhiman, C., Kumar, D.: XGL-T transformer model for intelligent image captioning. Multimedia Tools Appl 83(2), 4219\u20134240 (2024). https:\/\/doi.org\/10.1007\/s11042-023-15291-3","journal-title":"Multimedia Tools Appl"},{"issue":"1","key":"4349_CR43","doi-asserted-by":"publisher","first-page":"2201","DOI":"10.1002\/cav.2201","volume":"35","author":"X Zhu","year":"2024","unstructured":"Zhu, X., Yao, X., Zhang, J., Zhu, M., You, L., Yang, X., Zhang, J., Zhao, H., Zeng, D.: TMSDNet: transformer with multi-scale dense network for single and multi-view 3D reconstruction. Comput. Anim. Virtual Worlds 35(1), 2201 (2024). https:\/\/doi.org\/10.1002\/cav.2201","journal-title":"Comput. Anim. Virtual Worlds"},{"issue":"6","key":"4349_CR44","doi-asserted-by":"publisher","first-page":"2032","DOI":"10.1109\/TCDS.2024.3405573","volume":"16","author":"D Sharma","year":"2024","unstructured":"Sharma, D., Dhiman, C., Kumar, D.: Control with style: style embedding-based variational autoencoder for controlled stylized caption generation framework. IEEE Trans. Cogn. Dev. Syst. 16(6), 2032\u20132042 (2024). https:\/\/doi.org\/10.1109\/TCDS.2024.3405573","journal-title":"IEEE Trans. Cogn. Dev. Syst."},{"key":"4349_CR45","doi-asserted-by":"publisher","DOI":"10.1007\/s00371-025-04195-y","author":"X Huang","year":"2025","unstructured":"Huang, X., Li, X., Tan, S., Chen, W., Li, G.: GaPTalk: precision-controlled 3d gaussian rendering for personalized talking-head synthesis. Vis. Comput. (2025). https:\/\/doi.org\/10.1007\/s00371-025-04195-y","journal-title":"Vis. Comput."},{"issue":"10","key":"4349_CR46","doi-asserted-by":"publisher","first-page":"6625","DOI":"10.1109\/TVCG.2023.3322416","volume":"30","author":"T Xu","year":"2024","unstructured":"Xu, T., Ren, X., Yang, J., Sheng, B., Wu, E.: Efficient binocular rendering of volumetric density fields with coupled adaptive cube-map ray marching for virtual reality. IEEE Trans. Visual Comput. Graph. 30(10), 6625\u20136638 (2024). https:\/\/doi.org\/10.1109\/TVCG.2023.3322416","journal-title":"IEEE Trans. Visual Comput. Graph."},{"key":"4349_CR47","doi-asserted-by":"publisher","unstructured":"Rusu, R.B., Blodow, N., Beetz, M.: Fast point feature histograms (FPFH) for 3D registration. In: Proceedings of the 2009 IEEE International Conference on Robotics and Automation, pp. 3212\u20133217. IEEE, Kobe, Japan (2009). https:\/\/doi.org\/10.1109\/ROBOT.2009.5152473","DOI":"10.1109\/ROBOT.2009.5152473"},{"key":"4349_CR48","doi-asserted-by":"publisher","unstructured":"Deng, H., Birdal, T., Ilic, S.: PPF-FoldNet: unsupervised learning of rotation invariant 3D local descriptors. In: Computer Vision\u2013ECCV 2018: 15th European Conference, Munich, Germany, September 8\u201314, 2018, Proceedings, Part V. Lecture Notes in Computer Science, vol. 11209, pp. 602\u2013618. Springer, Munich, Germany (2018). https:\/\/doi.org\/10.1007\/978-3-030-01228-1_37","DOI":"10.1007\/978-3-030-01228-1_37"},{"issue":"9","key":"4349_CR49","doi-asserted-by":"publisher","first-page":"7898","DOI":"10.1109\/TPAMI.2025.3572584","volume":"47","author":"J Wang","year":"2025","unstructured":"Wang, J., Lu, X., Bennamoun, M., Sheng, B.: Non-rigid point cloud registration via anisotropic hybrid field harmonization. IEEE Trans. Pattern Anal. Mach. Intell. 47(9), 7898\u20137915 (2025). https:\/\/doi.org\/10.1109\/TPAMI.2025.3572584","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"issue":"3","key":"4349_CR50","doi-asserted-by":"publisher","first-page":"2277","DOI":"10.1002\/cav.2277","volume":"35","author":"J Feng","year":"2024","unstructured":"Feng, J., He, C., Wang, G., Wang, M.: S-LASSIE: structure and smoothness enhanced learning from sparse image ensemble for 3D articulated shape reconstruction. Comput. Anim. Virtual Worlds 35(3), 2277 (2024). https:\/\/doi.org\/10.1002\/cav.2277","journal-title":"Comput. Anim. Virtual Worlds"},{"key":"4349_CR51","doi-asserted-by":"publisher","unstructured":"Jensen, R., Dahl, A., Vogiatzis, G., Tola, E., Aan\u00e6s, H.: Large scale multi-view stereopsis evaluation. In: Proceedings of the 2014 IEEE Conference on Computer Vision and Pattern Recognition, pp. 406\u2013413. IEEE, Columbus, OH (2014). https:\/\/doi.org\/10.1109\/CVPR.2014.59","DOI":"10.1109\/CVPR.2014.59"},{"key":"4349_CR52","doi-asserted-by":"publisher","unstructured":"Deitke, M., Schwenk, D., Salvador, J., Weihs, L., Michel, O., VanderBilt, E., Kembhavi, A., Ehsani, K., Schmidt, L., Mottaghi, R.: Objaverse: a universe of annotated 3D objects. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 13142\u201313153. IEEE, Vancouver, BC (2023). https:\/\/doi.org\/10.1109\/CVPR52729.2023.01263","DOI":"10.1109\/CVPR52729.2023.01263"},{"key":"4349_CR53","doi-asserted-by":"publisher","unstructured":"Chang, A.X., Funkhouser, T., Guibas, L., Hanrahan, P., Huang, Q., Li, Z., Savarese, S., Savva, M., Song, S., Su, H., Xiao, J., Yi, L., Yu, F.: ShapeNet: an information-rich 3D model repository. arXiv preprint (2015) https:\/\/doi.org\/10.48550\/arXiv.1512.03012","DOI":"10.48550\/arXiv.1512.03012"},{"key":"4349_CR54","doi-asserted-by":"publisher","unstructured":"Groueix, T., Fisher, M., Kim, V.G., Russell, B.C., Aubry, M.: A papier-m\u00e2ch\u00e9 approach to learning 3D surface generation. In: Proceedings of the 2018 IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 216\u2013224. IEEE, Salt Lake City, UT (2018). https:\/\/doi.org\/10.1109\/CVPR.2018.00030","DOI":"10.1109\/CVPR.2018.00030"},{"key":"4349_CR55","doi-asserted-by":"publisher","unstructured":"Yang, H., Zhu, H., Wang, Y., Huang, M., Shen, Q., Yang, R., Cao, X.: FaceScape: a large-scale high-quality 3D face dataset and detailed Riggable 3D face prediction. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 6011\u20136020. IEEE, Seattle, WA, USA (2020). https:\/\/doi.org\/10.1109\/CVPR42600.2020.00605","DOI":"10.1109\/CVPR42600.2020.00605"},{"issue":"11","key":"4349_CR56","doi-asserted-by":"publisher","first-page":"2274","DOI":"10.1109\/TPAMI.2012.120","volume":"34","author":"R Achanta","year":"2012","unstructured":"Achanta, R., Shaji, A., Smith, K., Lucchi, A., Fua, P., S\u00fcsstrunk, S.: SLIC superpixels compared to state-of-the-art superpixel methods. IEEE Trans. Pattern Anal. Mach. Intell. 34(11), 2274\u20132282 (2012). https:\/\/doi.org\/10.1109\/TPAMI.2012.120","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"4349_CR57","doi-asserted-by":"publisher","unstructured":"Zhou, Q.-Y., Park, J., Koltun, V.: Open3D: a modern library for 3D data processing (2018) https:\/\/doi.org\/10.48550\/arXiv.1801.09847","DOI":"10.48550\/arXiv.1801.09847"},{"key":"4349_CR58","doi-asserted-by":"publisher","unstructured":"Zhang, R., Isola, P., Efros, A.A., Shechtman, E., Wang, O.: The unreasonable effectiveness of deep features as a perceptual metric. In: Proceedings of the 2018 IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 586\u2013595. IEEE, Salt Lake City, UT (2018). https:\/\/doi.org\/10.1109\/CVPR.2018.00068","DOI":"10.1109\/CVPR.2018.00068"},{"key":"4349_CR59","doi-asserted-by":"publisher","DOI":"10.1007\/s00371-025-04186-z","author":"B Cong","year":"2025","unstructured":"Cong, B., Wang, X., Zhao, X., et al.: 3D-SDIS: enhanced 3d instance segmentation through frequency fusion and dual-sphere sampling. Vis. Comput. (2025). https:\/\/doi.org\/10.1007\/s00371-025-04186-z","journal-title":"Vis. Comput."},{"key":"4349_CR60","doi-asserted-by":"publisher","unstructured":"Hou, J., Chau, L.-P., He, Y., Magnenat-Thalmann, N.: A novel compression framework for 3D time-varying meshes. In: 2014 IEEE International Symposium on Circuits and Systems (ISCAS), pp. 2161\u20132164, Melbourne, Australia (2014). https:\/\/doi.org\/10.1109\/ISCAS.2014.6865596","DOI":"10.1109\/ISCAS.2014.6865596"},{"key":"4349_CR61","doi-asserted-by":"publisher","DOI":"10.1109\/TCSVT.2025.3597210","author":"M Xiong","year":"2025","unstructured":"Xiong, M., Ge, L., Hu, R., Muhammad, K., Bakshi, S., Del Ser, J., Yang, X., Sheng, B.: HPRnet: human parsing reconstruction with non-local multi-scale perception network for cloth-changing person re-identification. IEEE Trans. Circuits Syst. Video Technol. (2025). https:\/\/doi.org\/10.1109\/TCSVT.2025.3597210","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"4349_CR62","doi-asserted-by":"publisher","unstructured":"Zhao, H., Shi, J., Qi, X., Wang, X., Jia, J.: Pyramid scene parsing network. In: Proceedings of the 2017 IEEE Conference on Computer Vision and Pattern Recognition, pp. 6230\u20136239. IEEE, Honolulu, HI (2017). https:\/\/doi.org\/10.1109\/CVPR.2017.660","DOI":"10.1109\/CVPR.2017.660"}],"container-title":["The Visual Computer"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00371-025-04349-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00371-025-04349-y","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00371-025-04349-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,3,13]],"date-time":"2026-03-13T16:28:23Z","timestamp":1773419303000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00371-025-04349-y"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,2]]},"references-count":62,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2026,2]]}},"alternative-id":["4349"],"URL":"https:\/\/doi.org\/10.1007\/s00371-025-04349-y","relation":{},"ISSN":["0178-2789","1432-2315"],"issn-type":[{"value":"0178-2789","type":"print"},{"value":"1432-2315","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,2]]},"assertion":[{"value":"16 October 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"29 December 2025","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"2 February 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no Conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"153"}}