{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,21]],"date-time":"2026-05-21T08:08:07Z","timestamp":1779350887943,"version":"3.51.4"},"reference-count":62,"publisher":"Springer Science and Business Media LLC","issue":"4","license":[{"start":{"date-parts":[[2026,3,6]],"date-time":"2026-03-06T00:00:00Z","timestamp":1772755200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,3,6]],"date-time":"2026-03-06T00:00:00Z","timestamp":1772755200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int J Comput Vis"],"published-print":{"date-parts":[[2026,4]]},"DOI":"10.1007\/s11263-026-02747-w","type":"journal-article","created":{"date-parts":[[2026,3,6]],"date-time":"2026-03-06T09:27:42Z","timestamp":1772789262000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["CylindFormer: Image-to-Point Cloud Registration with Cylindrical Transformer"],"prefix":"10.1007","volume":"134","author":[{"given":"Jingtao","family":"Wang","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Hao","family":"Tang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yanpeng","family":"Sun","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Shengfeng","family":"He","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-5341-5985","authenticated-orcid":false,"given":"Zechao","family":"Li","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2026,3,6]]},"reference":[{"key":"2747_CR1","unstructured":"Bradski, G. (2000). The opencv library. Dr. Dobb\u2019s Journal of Software Tools"},{"issue":"2","key":"2747_CR2","doi-asserted-by":"publisher","first-page":"910","DOI":"10.1007\/s11263-024-02207-3","volume":"133","author":"MA Butt","year":"2025","unstructured":"Butt, M. A., Ali, H., Qayyum, A., Sultani, W., Al-Fuqaha, A., & Qadir, J. (2025). R 2 s100k: Road-region segmentation dataset for semi-supervised autonomous driving in the wild. International Journal of Computer Vision, 133(2), 910\u2013928.","journal-title":"International Journal of Computer Vision"},{"key":"2747_CR3","doi-asserted-by":"crossref","unstructured":"Caesar, H., Bankiti, V., Lang, A. H., Vora, S., Liong, V. E., Xu, Q., Krishnan, A., Pan, Y., Baldan, G., & Beijbom, O. (2020). nuscenes: A multimodal dataset for autonomous driving. In Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 11,621\u201311,631.","DOI":"10.1109\/CVPR42600.2020.01164"},{"key":"2747_CR4","doi-asserted-by":"crossref","unstructured":"Cattaneo, D., Vaghi, M., Fontana, S., Ballardini, A. L., & Sorrenti, D. G. (2020). Global visual localization in lidar-maps through shared 2d\u20133d embedding space. In 2020 IEEE International Conference on Robotics and Automation (ICRA), 4365\u20134371.","DOI":"10.1109\/ICRA40945.2020.9196859"},{"key":"2747_CR5","doi-asserted-by":"crossref","unstructured":"Chen, H., Luo, Z., Zhou, L., Tian, Y., Zhen, M., Fang, T., Mckinnon, D., Tsin, Y., & Quan, L. (2022). Aspanformer: Detector-free image matching with adaptive span transformer. In: European Conference on Computer Vision, 20\u201336. Springer","DOI":"10.1007\/978-3-031-19824-3_2"},{"key":"2747_CR6","doi-asserted-by":"crossref","unstructured":"Chen, H., Yan, P., Xiang, S., & Tan, Y. (2024). Dynamic cues-assisted transformer for robust point cloud registration. In Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 21698\u201321707.","DOI":"10.1109\/CVPR52733.2024.02050"},{"key":"2747_CR7","doi-asserted-by":"crossref","unstructured":"Choy, C., Park, J., & Koltun, V. (2019). Fully convolutional geometric features. In Proceedings of the IEEE\/CVF International Conference on Computer Vision, 8958\u20138966.","DOI":"10.1109\/ICCV.2019.00905"},{"issue":"2","key":"2747_CR8","doi-asserted-by":"publisher","first-page":"99","DOI":"10.1109\/MRA.2006.1638022","volume":"13","author":"H Durrant-Whyte","year":"2006","unstructured":"Durrant-Whyte, H., & Bailey, T. (2006). Simultaneous localization and mapping: part i. IEEE Robotics & Automation Magazine, 13(2), 99\u2013110.","journal-title":"IEEE Robotics & Automation Magazine"},{"key":"2747_CR9","doi-asserted-by":"crossref","unstructured":"Feng, M., Hu, S., Ang, M.H., & Lee, G.H. (2019). 2d3d-matchnet: Learning to match keypoints across 2d image and 3d point cloud. In: 2019 International Conference on Robotics and Automation, 4790\u20134796. IEEE","DOI":"10.1109\/ICRA.2019.8794415"},{"issue":"6","key":"2747_CR10","doi-asserted-by":"publisher","first-page":"381","DOI":"10.1145\/358669.358692","volume":"24","author":"MA Fischler","year":"1981","unstructured":"Fischler, M. A., & Bolles, R. C. (1981). Random sample consensus: A paradigm for model fitting with applications to image analysis and automated cartography. Communications of the ACM, 24(6), 381\u2013395.","journal-title":"Communications of the ACM"},{"issue":"11","key":"2747_CR11","doi-asserted-by":"publisher","first-page":"1231","DOI":"10.1177\/0278364913491297","volume":"32","author":"A Geiger","year":"2013","unstructured":"Geiger, A., Lenz, P., Stiller, C., & Urtasun, R. (2013). Vision meets robotics: The kitti dataset. The International Journal of Robotics Research, 32(11), 1231\u20131237.","journal-title":"The International Journal of Robotics Research"},{"key":"2747_CR12","doi-asserted-by":"crossref","unstructured":"Giang, K. T., Song, S., & Jo, S. (2023). Topicfm: Robust and interpretable topic-assisted feature matching. In: Proceedings of the AAAI Conference on Artificial Intelligence, 37, 2447\u20132455.","DOI":"10.1609\/aaai.v37i2.25341"},{"key":"2747_CR13","doi-asserted-by":"crossref","unstructured":"Glocker, B., Izadi, S., Shotton, J., & Criminisi, A. (2013). Real-time rgb-d camera relocalization. In: 2013 IEEE International Symposium on Mixed and Augmented Reality, 173\u2013179. IEEE","DOI":"10.1109\/ISMAR.2013.6671777"},{"key":"2747_CR14","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., & Sun, J. (2016). Deep residual learning for image recognition. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 770\u2013778.","DOI":"10.1109\/CVPR.2016.90"},{"key":"2747_CR15","doi-asserted-by":"crossref","unstructured":"Huang, S., Gojcic, Z., Usvyatsov, M., Wieser, A., & Schindler, K. (2021). Predator: Registration of 3d point clouds with low overlap. In Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 4267\u20134276.","DOI":"10.1109\/CVPR46437.2021.00425"},{"key":"2747_CR16","doi-asserted-by":"crossref","unstructured":"Kang, S., Liao, Y., Li, J., Liang, F., Li, Y., Zou, X., Li, F., Chen, X., Dong, Z., & Yang, B. (2024). Cofii2p: Coarse-to-fine correspondences-based image to point cloud registration. IEEE Robotics and Automation Letters","DOI":"10.1109\/LRA.2024.3466068"},{"key":"2747_CR17","doi-asserted-by":"crossref","unstructured":"Kim, M., Koo, J., & Kim, G. (2023). Ep2p-loc: End-to-end 3d point to 2d pixel localization for large-scale visual localization. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, 21527\u201321537.","DOI":"10.1109\/ICCV51070.2023.01968"},{"key":"2747_CR18","doi-asserted-by":"crossref","unstructured":"Lai, K., Bo, L., & Fox, D. (2014). Unsupervised feature learning for 3d scene labeling. In: 2014 IEEE International Conference on Robotics and Automation, 3050\u20133057. IEEE","DOI":"10.1109\/ICRA.2014.6907298"},{"key":"2747_CR19","doi-asserted-by":"publisher","first-page":"155","DOI":"10.1007\/s11263-008-0152-6","volume":"81","author":"V Lepetit","year":"2009","unstructured":"Lepetit, V., Moreno-Noguer, F., & Fua, P. (2009). Ep n p: An accurate o (n) solution to the p n p problem. International Journal of Computer Vision, 81, 155\u2013166.","journal-title":"International Journal of Computer Vision"},{"key":"2747_CR20","doi-asserted-by":"crossref","unstructured":"Li, J., & Lee, G.H. (2021). Deepi2p: Image-to-point cloud registration via deep classification. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 15,960\u201315,969","DOI":"10.1109\/CVPR46437.2021.01570"},{"key":"2747_CR21","doi-asserted-by":"crossref","unstructured":"Li, M., Qin, Z., Gao, Z., Yi, R., Zhu, C., Guo, Y., & Xu, K. (2023). 2d3d-matr: 2d-3d matching transformer for detection-free registration between images and point clouds. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, 14,128\u201314,138","DOI":"10.1109\/ICCV51070.2023.01299"},{"key":"2747_CR22","doi-asserted-by":"crossref","unstructured":"Lin, T.Y., Doll\u00e1r, P., Girshick, R., He, K., Hariharan, B., & Belongie, S. (2017). Feature pyramid networks for object detection. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 2117\u20132125","DOI":"10.1109\/CVPR.2017.106"},{"key":"2747_CR23","doi-asserted-by":"crossref","unstructured":"Liu, J., Zhuo, D., Feng, Z., Zhu, S., Peng, C., Liu, Z., & Wang, H. (2025). Dvlo: Deep visual-lidar odometry with local-to-global feature fusion and bi-directional structure alignment. In: European Conference on Computer Vision, 475\u2013493. Springer","DOI":"10.1007\/978-3-031-72684-2_27"},{"key":"2747_CR24","doi-asserted-by":"publisher","first-page":"1367","DOI":"10.1109\/TIP.2023.3242598","volume":"32","author":"X Liu","year":"2023","unstructured":"Liu, X., Xiao, G., Chen, R., & Ma, J. (2023). Pgfnet: Preference-guided filtering network for two-view correspondence learning. IEEE Transactions on Image Processing, 32, 1367\u20131378.","journal-title":"IEEE Transactions on Image Processing"},{"key":"2747_CR25","doi-asserted-by":"crossref","unstructured":"Liu, X., & Yang, J. (2023). Progressive neighbor consistency mining for correspondence pruning. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 9527\u20139537","DOI":"10.1109\/CVPR52729.2023.00919"},{"key":"2747_CR26","doi-asserted-by":"publisher","first-page":"512","DOI":"10.1007\/s11263-018-1117-z","volume":"127","author":"J Ma","year":"2019","unstructured":"Ma, J., Zhao, J., Jiang, J., Zhou, H., & Guo, X. (2019). Locality preserving matching. International Journal of Computer Vision, 127, 512\u2013531.","journal-title":"International Journal of Computer Vision"},{"key":"2747_CR27","doi-asserted-by":"publisher","first-page":"155","DOI":"10.1016\/j.isprsjprs.2025.08.016","volume":"229","author":"W Ma","year":"2025","unstructured":"Ma, W., Huang, Y., Tang, S., Zheng, X., Dong, Z., Ge, L., Pan, J., Li, Q., & Wang, B. (2025). Cross-modal 2d\u20133d feature matching: Simultaneous local feature description and detection across images and point clouds. ISPRS Journal of Photogrammetry and Remote Sensing, 229, 155\u2013169.","journal-title":"ISPRS Journal of Photogrammetry and Remote Sensing"},{"issue":"1","key":"2747_CR28","doi-asserted-by":"publisher","first-page":"99","DOI":"10.1145\/3503250","volume":"65","author":"B Mildenhall","year":"2021","unstructured":"Mildenhall, B., Srinivasan, P. P., Tancik, M., Barron, J. T., Ramamoorthi, R., & Ng, R. (2021). Nerf: Representing scenes as neural radiance fields for view synthesis. Communications of the ACM, 65(1), 99\u2013106.","journal-title":"Communications of the ACM"},{"key":"2747_CR29","doi-asserted-by":"crossref","unstructured":"Mu, J., Bie, L., Du, S., & Gao, Y. (2024). Colorpcr: Color point cloud registration with multi-stage geometric-color fusion. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 21,061\u201321,070","DOI":"10.1109\/CVPR52733.2024.01990"},{"key":"2747_CR30","unstructured":"Oquab, M., Darcet, T., Moutakanni, T., Vo, H., Szafraniec, M., Khalidov, V., Fernandez, P., Haziza, D., Massa, F., & El-Nouby, A., et al. (2023). Dinov2: Learning robust visual features without supervision. arXiv preprint arXiv:2304.07193"},{"key":"2747_CR31","doi-asserted-by":"crossref","unstructured":"Peng, C., Wang, G., Lo, X.W., Wu, X., Xu, C., Tomizuka, M., Zhan, W., & Wang, H. (2023). Delflow: Dense efficient learning of scene flow for large-scale point clouds. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, 16,901\u201316,910","DOI":"10.1109\/ICCV51070.2023.01550"},{"key":"2747_CR32","doi-asserted-by":"crossref","unstructured":"Pham, Q. H., Uy, M. A., Hua, B. S., Nguyen, D. T., Roig, G., & Yeung, S. K. (2020). Lcd: Learned cross-domain descriptors for 2d\u20133d matching. In: Proceedings of the AAAI conference on artificial intelligence, 34, 11856\u201311864.","DOI":"10.1609\/aaai.v34i07.6859"},{"key":"2747_CR33","doi-asserted-by":"crossref","unstructured":"Qin, Z., Yu, H., Wang, C., Guo, Y., Peng, Y., & Xu, K. (2022). Geometric transformer for fast and robust point cloud registration. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 11,143\u201311,152","DOI":"10.1109\/CVPR52688.2022.01086"},{"issue":"3","key":"2747_CR34","doi-asserted-by":"publisher","first-page":"1198","DOI":"10.1109\/TCSVT.2022.3208859","volume":"33","author":"S Ren","year":"2022","unstructured":"Ren, S., Zeng, Y., Hou, J., & Chen, X. (2022). Corri2p: Deep image-to-point cloud registration via dense correspondence. IEEE Transactions on Circuits and Systems for Video Technology, 33(3), 1198\u20131208.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"key":"2747_CR35","doi-asserted-by":"crossref","unstructured":"Rosten, E., & Drummond, T. (2006). Machine learning for high-speed corner detection. In: Proceedings of the 9th European Conference on Computer Vision, pp. 430\u2013443","DOI":"10.1007\/11744023_34"},{"key":"2747_CR36","doi-asserted-by":"publisher","first-page":"963","DOI":"10.1007\/s00371-011-0610-y","volume":"27","author":"I Sipiran","year":"2011","unstructured":"Sipiran, I., & Bustos, B. (2011). Harris 3d: a robust extension of the harris operator for interest point detection on 3d meshes. The Visual Computer, 27, 963\u2013976.","journal-title":"The Visual Computer"},{"key":"2747_CR37","doi-asserted-by":"crossref","unstructured":"Sturm, J., Engelhard, N., Endres, F., Burgard, W., & Cremers, D. (2012). A benchmark for the evaluation of rgb-d slam systems. In: 2012 IEEE\/RSJ international conference on intelligent robots and systems, pp. 573\u2013580","DOI":"10.1109\/IROS.2012.6385773"},{"key":"2747_CR38","doi-asserted-by":"crossref","unstructured":"Sun, J., Shen, Z., Wang, Y., Bao, H., & Zhou, X. (2021). Loftr: Detector-free local feature matching with transformers. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 8922\u20138931","DOI":"10.1109\/CVPR46437.2021.00881"},{"key":"2747_CR39","doi-asserted-by":"crossref","unstructured":"Sun, Y., Cheng, C., Zhang, Y., Zhang, C., Zheng, L., Wang, Z., & Wei, Y. (2020). Circle loss: A unified perspective of pair similarity optimization. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 6398\u20136407","DOI":"10.1109\/CVPR42600.2020.00643"},{"key":"2747_CR40","unstructured":"Tang, S., Zhang, J., Zhu, S., & Tan, P. (2022). Quadtree attention for vision transformers. International Conference on Learning Representations"},{"key":"2747_CR41","doi-asserted-by":"crossref","unstructured":"Thomas, H., Qi, C.R., Deschaud, J.E., Marcotegui, B., Goulette, F., & Guibas, L.J. (2019). Kpconv: Flexible and deformable convolution for point clouds. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, 6411\u20136420","DOI":"10.1109\/ICCV.2019.00651"},{"key":"2747_CR42","unstructured":"Vaswani, A., Shazeer, N., Parmar, N., Uszkoreit, J., Jones, L., Gomez, A.N., Kaiser, \u0141., & Polosukhin, I. (2017). Attention is all you need. Advances in Neural Information Processing Systems 30"},{"key":"2747_CR43","doi-asserted-by":"crossref","unstructured":"Wang, B., Chen, C., Cui, Z., Qin, J., Lu, C.X., Yu, Z., Zhao, P., Dong, Z., Zhu, F., & Trigoni, N., et\u00a0al. (2021). P2-net: Joint description and detection of local features for pixel and point matching. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 16,004\u201316,013","DOI":"10.1109\/ICCV48922.2021.01570"},{"issue":"5","key":"2747_CR44","first-page":"5749","volume":"45","author":"G Wang","year":"2022","unstructured":"Wang, G., Wu, X., Jiang, S., Liu, Z., & Wang, H. (2022). Efficient 3d deep lidar odometry. IEEE Transactions on Pattern Analysis and Machine Intelligence, 45(5), 5749\u20135765.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"2747_CR45","unstructured":"Wang, H., Liu, Y., Wang, B., Sun, Y., Dong, Z., Wang, W., & Yang, B. (2024). Freereg: Image-to-point cloud registration leveraging pretrained diffusion models and monocular depth estimators. International Conference on Learning Representations"},{"key":"2747_CR46","doi-asserted-by":"crossref","unstructured":"Wang, J., & Li, Z. (2024). 3dpcp-net: A lightweight progressive 3d correspondence pruning network for accurate and efficient point cloud registration. In: Proceedings of the 32nd ACM International Conference on Multimedia, pp. 1885\u20131894","DOI":"10.1145\/3664647.3681320"},{"key":"2747_CR47","first-page":"1","volume":"61","author":"J Wang","year":"2023","unstructured":"Wang, J., Liu, X., Dai, L., Ma, J., Wei, L., Yang, C., & Chen, R. (2023). Pg-net: Progressive guidance network via robust contextual embedding for efficient point cloud registration. IEEE Transactions on Geoscience and Remote Sensing, 61, 1\u201312.","journal-title":"IEEE Transactions on Geoscience and Remote Sensing"},{"issue":"22","key":"2747_CR48","doi-asserted-by":"publisher","first-page":"5751","DOI":"10.3390\/rs14225751","volume":"14","author":"J Wang","year":"2022","unstructured":"Wang, J., Yang, C., Wei, L., & Chen, R. (2022). Csce-net: Channel-spatial contextual enhancement network for robust point cloud registration. Remote Sensing, 14(22), 5751.","journal-title":"Remote Sensing"},{"issue":"11","key":"2747_CR49","doi-asserted-by":"publisher","first-page":"11981","DOI":"10.1109\/TITS.2023.3285651","volume":"24","author":"L Wang","year":"2023","unstructured":"Wang, L., Zhang, X., Qin, W., Li, X., Gao, J., Yang, L., Li, Z., Li, J., Zhu, L., Wang, H., et al. (2023). Camo-mot: Combined appearance-motion optimization for 3d multi-object tracking with camera-lidar fusion. IEEE Transactions on Intelligent Transportation Systems, 24(11), 11981\u201311996.","journal-title":"IEEE Transactions on Intelligent Transportation Systems"},{"key":"2747_CR50","doi-asserted-by":"crossref","unstructured":"Wang, Q., Zhang, J., Yang, K., Peng, K., & Stiefelhagen, R. (2022). Matchformer: Interleaving attention in transformers for feature matching. In: Proceedings of the Asian Conference on Computer Vision, pp. 2746\u20132762","DOI":"10.1007\/978-3-031-26313-2_16"},{"key":"2747_CR51","doi-asserted-by":"crossref","unstructured":"Wang, Y., He, X., Peng, S., Tan, D., & Zhou, X. (2024). Efficient loftr: Semi-dense local feature matching with sparse-like speed. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 21,666\u201321,675","DOI":"10.1109\/CVPR52733.2024.02047"},{"issue":"5","key":"2747_CR52","doi-asserted-by":"publisher","first-page":"2620","DOI":"10.1007\/s11263-024-02317-y","volume":"133","author":"Y Wang","year":"2025","unstructured":"Wang, Y., Li, H., & Luo, C. (2025). Object pose estimation based on multi-precision vectors and seg-driven pnp. International Journal of Computer Vision, 133(5), 2620\u20132634.","journal-title":"International Journal of Computer Vision"},{"key":"2747_CR53","doi-asserted-by":"crossref","unstructured":"Wu, Q., Jiang, H., Luo, L., Li, J., Ding, Y., Xie, J., & Yang, J. (2025). Diff-reg: Diffusion model in doubly stochastic matrix space for registration problem. In: European Conference on Computer Vision, 160\u2013178. Springer","DOI":"10.1007\/978-3-031-73650-6_10"},{"key":"2747_CR54","doi-asserted-by":"crossref","unstructured":"Yao, G., Xuan, Y., Chen, Y., & Pan, Y. (2024). Quantity-aware coarse-to-fine correspondence for image-to-point cloud registration. IEEE Sensors Journal","DOI":"10.1109\/JSEN.2024.3454822"},{"key":"2747_CR55","doi-asserted-by":"crossref","unstructured":"Yeshwanth, C., Liu, Y.C., Nie\u00dfner, M., & Dai, A. (2023). Scannet++: A high-fidelity dataset of 3d indoor scenes. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 12\u201322","DOI":"10.1109\/ICCV51070.2023.00008"},{"key":"2747_CR56","first-page":"23872","volume":"34","author":"H Yu","year":"2021","unstructured":"Yu, H., Li, F., Saleh, M., Busam, B., & Ilic, S. (2021). Cofinet: Reliable coarse-to-fine correspondences for robust pointcloud registration. Advances in Neural Information Processing Systems, 34, 23872\u201323884.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"2747_CR57","doi-asserted-by":"crossref","unstructured":"Yu, H., Qin, Z., Hou, J., Saleh, M., Li, D., Busam, B., & Ilic, S. (2023a). Rotation-invariant transformer for point cloud matching. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 5384\u20135393","DOI":"10.1109\/CVPR52729.2023.00521"},{"key":"2747_CR58","doi-asserted-by":"crossref","unstructured":"Yu, J., Ren, L., Zhang, Y., Zhou, W., Lin, L., & Dai, G. (2023b). Peal: Prior-embedded explicit attention learning for low-overlap point cloud registration. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 17,702\u201317,711","DOI":"10.1109\/CVPR52729.2023.01698"},{"key":"2747_CR59","doi-asserted-by":"crossref","unstructured":"Yu, Z., Qin, Z., Zheng, L., & Xu, K. (2024). Learning instance-aware correspondences for robust multi-instance point cloud registration in cluttered scenes. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 19,605\u201319,614","DOI":"10.1109\/CVPR52733.2024.01854"},{"issue":"10","key":"2747_CR60","doi-asserted-by":"publisher","first-page":"2425","DOI":"10.1007\/s11263-022-01657-x","volume":"130","author":"\u00c9 Zablocki","year":"2022","unstructured":"Zablocki, \u00c9., Ben-Younes, H., P\u00e9rez, P., & Cord, M. (2022). Explainability of deep vision-based autonomous driving systems: Review and challenges. International Journal of Computer Vision, 130(10), 2425\u20132452.","journal-title":"International Journal of Computer Vision"},{"key":"2747_CR61","doi-asserted-by":"crossref","unstructured":"Zheng, Q., Liu, D., Wang, C., Zhang, J., Wang, D., & Tao, D. (2025). Esceme: Vision-and-language navigation with episodic scene memory. International Journal of Computer Vision, 133(1), 254\u2013274.","DOI":"10.1007\/s11263-024-02159-8"},{"key":"2747_CR62","first-page":"51166","volume":"36","author":"J Zhou","year":"2023","unstructured":"Zhou, J., Ma, B., Zhang, W., Fang, Y., Liu, Y. S., & Han, Z. (2023). Differentiable registration of images and lidar point clouds with voxelpoint-to-pixel matching. Advances in Neural Information Processing Systems, 36, 51166\u201351177.","journal-title":"Advances in Neural Information Processing Systems"}],"container-title":["International Journal of Computer Vision"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-026-02747-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11263-026-02747-w","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-026-02747-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,5,21]],"date-time":"2026-05-21T07:32:09Z","timestamp":1779348729000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11263-026-02747-w"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,3,6]]},"references-count":62,"journal-issue":{"issue":"4","published-print":{"date-parts":[[2026,4]]}},"alternative-id":["2747"],"URL":"https:\/\/doi.org\/10.1007\/s11263-026-02747-w","relation":{},"ISSN":["0920-5691","1573-1405"],"issn-type":[{"value":"0920-5691","type":"print"},{"value":"1573-1405","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,3,6]]},"assertion":[{"value":"10 July 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"4 January 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"6 March 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflicts of Interest"}}],"article-number":"145"}}