{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,1,18]],"date-time":"2026-01-18T03:17:23Z","timestamp":1768706243187,"version":"3.49.0"},"reference-count":71,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2024,5,31]],"date-time":"2024-05-31T00:00:00Z","timestamp":1717113600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,5,31]],"date-time":"2024-05-31T00:00:00Z","timestamp":1717113600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"crossref","award":["62111530300"],"award-info":[{"award-number":["62111530300"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"crossref"}]},{"DOI":"10.13039\/100022955","name":"Fundamental Research Funds for the Provincial Universities of Zhejiang","doi-asserted-by":"publisher","award":["JRK22003"],"award-info":[{"award-number":["JRK22003"]}],"id":[{"id":"10.13039\/100022955","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100004731","name":"Natural Science Foundation of Zhejiang Province","doi-asserted-by":"crossref","award":["Z24F020002"],"award-info":[{"award-number":["Z24F020002"]}],"id":[{"id":"10.13039\/501100004731","id-type":"DOI","asserted-by":"crossref"}]},{"name":"Opening Foundation of State Key Laboratory of Virtual Reality Technology and System of Beihang University","award":["VRLAB2023B02"],"award-info":[{"award-number":["VRLAB2023B02"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimedia Systems"],"published-print":{"date-parts":[[2024,6]]},"DOI":"10.1007\/s00530-024-01369-x","type":"journal-article","created":{"date-parts":[[2024,5,31]],"date-time":"2024-05-31T13:02:08Z","timestamp":1717160528000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":2,"title":["Multiscale geometric window transformer for orthodontic teeth point cloud registration"],"prefix":"10.1007","volume":"30","author":[{"given":"Hao","family":"Wang","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yan","family":"Tian","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yongchuan","family":"Xu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Jiahui","family":"Xu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Tao","family":"Yang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yan","family":"Lu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Hong","family":"Chen","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2024,5,31]]},"reference":[{"key":"1369_CR1","doi-asserted-by":"crossref","unstructured":"Tian, Y., Xu, Z., Ma, Y., Ding, W., Wang, R., Gao, Z., Cheng, G., He, L., Zhao, X.: Survey on deep learning in multimodal medical imaging for cancer detection. Neural Computing and Applications, 22071\u201322085 (2023)","DOI":"10.1007\/s00521-023-09214-4"},{"key":"1369_CR2","doi-asserted-by":"crossref","unstructured":"Tian, Y., Jian, G., Wang, J., Chen, H., Pan, L., Xu, Z., Li, J., Wang, R.: A revised approach to orthodontic treatment monitoring from oralscan video. IEEE Journal of Biomedical and Health Informatics 27(12) (2023)","DOI":"10.1109\/JBHI.2023.3319361"},{"issue":"1","key":"1369_CR3","doi-asserted-by":"publisher","DOI":"10.1007\/s11432-023-3847-x","volume":"67","author":"Y Tian","year":"2024","unstructured":"Tian, Y., Fu, H., Wang, H., Liu, Y., Xu, Z., Chen, H., Li, J., Wang, R.: Rgb oralscan video-based orthodontic treatment monitoring. Sci. China Inform. Sci. 67(1), 112107 (2024)","journal-title":"Sci. China Inform. Sci."},{"key":"1369_CR4","doi-asserted-by":"crossref","unstructured":"Zeng, A., Song, S., Nie\u00dfner, M., Fisher, M., Xiao, J., Funkhouser, T.: 3dmatch: Learning local geometric descriptors from rgb-d reconstructions. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 1802\u20131811 (2017)","DOI":"10.1109\/CVPR.2017.29"},{"key":"1369_CR5","doi-asserted-by":"crossref","unstructured":"Geiger, A., Lenz, P., Urtasun, R.: Are we ready for autonomous driving? the kitti vision benchmark suite. In: IEEE Conference on Computer Vision and Pattern Recognition, pp. 3354\u20133361 (2012)","DOI":"10.1109\/CVPR.2012.6248074"},{"key":"1369_CR6","first-page":"23872","volume":"34","author":"H Yu","year":"2021","unstructured":"Yu, H., Li, F., Saleh, M., Busam, B., Ilic, S.: Cofinet: Reliable coarse-to-fine correspondences for robust pointcloud registration. Adv. Neural Inform. Process. Syst. 34, 23872\u201323884 (2021)","journal-title":"Adv. Neural Inform. Process. Syst."},{"key":"1369_CR7","doi-asserted-by":"crossref","unstructured":"Li, Y., Harada, T.: Lepard: Learning partial point cloud matching in rigid and deformable scenes. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 5554\u20135564 (2022)","DOI":"10.1109\/CVPR52688.2022.00547"},{"key":"1369_CR8","doi-asserted-by":"crossref","unstructured":"Qin, Z., Yu, H., Wang, C., Guo, Y., Peng, Y., Xu, K.: Geometric transformer for fast and robust point cloud registration. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 11143\u201311152 (2022)","DOI":"10.1109\/CVPR52688.2022.01086"},{"key":"1369_CR9","unstructured":"Vaswani, A., Shazeer, N., Parmar, N., Uszkoreit, J., Jones, L., Gomez, A.N., Kaiser, \u0141., Polosukhin, I.: Attention is all you need. Advances in Neural Information Processing Systems 30 (2017)"},{"key":"1369_CR10","doi-asserted-by":"crossref","unstructured":"Chen, Z., Chen, H., Gong, L., Yan, X., Wang, J., Guo, Y., Qin, J., Wei, M.: Utopic: Uncertainty-aware overlap prediction network for partial point cloud registration. arXiv preprint arXiv:2208.02712 (2022)","DOI":"10.1111\/cgf.14659"},{"key":"1369_CR11","doi-asserted-by":"crossref","unstructured":"Tu, Z., Talebi, H., Zhang, H., Yang, F., Milanfar, P., Bovik, A., Li, Y.: Maxvit: Multi-axis vision transformer. In: European Conference on Computer Vision, pp. 459\u2013479 (2022)","DOI":"10.1007\/978-3-031-20053-3_27"},{"key":"1369_CR12","doi-asserted-by":"crossref","unstructured":"Liu, Z., Lin, Y., Cao, Y., Hu, H., Wei, Y., Zhang, Z., Lin, S., Guo, B.: Swin transformer: Hierarchical vision transformer using shifted windows. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 10012\u201310022 (2021)","DOI":"10.1109\/ICCV48922.2021.00986"},{"key":"1369_CR13","unstructured":"Besl, P.J., McKay, N.D.: Method for registration of 3-d shapes. In: Sensor Fusion IV: Control Paradigms and Data Structures, vol. 1611, pp. 586\u2013606 (1992)"},{"key":"1369_CR14","doi-asserted-by":"crossref","unstructured":"Segal, A., Haehnel, D., Thrun, S.: Generalized-icp. In: Robotics: Science and Systems, p. 435 (2009)","DOI":"10.15607\/RSS.2009.V.021"},{"key":"1369_CR15","doi-asserted-by":"crossref","unstructured":"Bouaziz, S., Tagliasacchi, A., Pauly, M.: Sparse iterative closest point. In: Computer Graphics Forum, vol. 32, pp. 113\u2013123 (2013)","DOI":"10.1111\/cgf.12178"},{"key":"1369_CR16","doi-asserted-by":"crossref","unstructured":"Yang, J., Li, H., Campbell, D., Jia, Y.: Go-icp: A globally optimal solution to 3d icp point-set registration. IEEE Transactions on Pattern Analysis and Machine Intelligence 38(11) (2015)","DOI":"10.1109\/TPAMI.2015.2513405"},{"key":"1369_CR17","doi-asserted-by":"crossref","unstructured":"Mellado, N., Aiger, D., Mitra, N.J.: Super 4pcs fast global pointcloud registration via smart indexing. In: Computer Graphics Forum, vol. 33, pp. 205\u2013215 (2014)","DOI":"10.1111\/cgf.12446"},{"key":"1369_CR18","doi-asserted-by":"crossref","unstructured":"Zhou, Q.-Y., Park, J., Koltun, V.: Fast global registration. In: European Conference on Computer Vision, pp. 766\u2013782 (2016)","DOI":"10.1007\/978-3-319-46475-6_47"},{"key":"1369_CR19","doi-asserted-by":"crossref","unstructured":"Choy, C., Park, J., Koltun, V.: Fully convolutional geometric features. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 8958\u20138966 (2019)","DOI":"10.1109\/ICCV.2019.00905"},{"key":"1369_CR20","doi-asserted-by":"crossref","unstructured":"Gojcic, Z., Zhou, C., Wegner, J.D., Wieser, A.: The perfect match: 3d point cloud matching with smoothed densities. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 5545\u20135554 (2019)","DOI":"10.1109\/CVPR.2019.00569"},{"key":"1369_CR21","doi-asserted-by":"crossref","unstructured":"Yao, Y., Deng, B., Xu, W., Zhang, J.: Quasi-newton solver for robust non-rigid registration. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 7600\u20137609 (2020)","DOI":"10.1109\/CVPR42600.2020.00762"},{"key":"1369_CR22","doi-asserted-by":"crossref","unstructured":"Bai, X., Luo, Z., Zhou, L., Chen, H., Li, L., Hu, Z., Fu, H., Tai, C.-L.: Pointdsc: Robust point cloud registration using deep spatial consistency. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 15859\u201315869 (2021)","DOI":"10.1109\/CVPR46437.2021.01560"},{"key":"1369_CR23","doi-asserted-by":"crossref","unstructured":"Chen, Z., Sun, K., Yang, F., Tao, W.: Sc2-pcr: A second order spatial compatibility for efficient and robust point cloud registration. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 13221\u201313231 (2022)","DOI":"10.1109\/CVPR52688.2022.01287"},{"key":"1369_CR24","doi-asserted-by":"crossref","unstructured":"Choy, C., Dong, W., Koltun, V.: Deep global registration. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 2514\u20132523 (2020)","DOI":"10.1109\/CVPR42600.2020.00259"},{"key":"1369_CR25","doi-asserted-by":"crossref","unstructured":"Pais, G.D., Ramalingam, S., Govindu, V.M., Nascimento, J.C., Chellappa, R., Miraldo, P.: 3dregnet: A deep neural network for 3d point registration. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 7193\u20137203 (2020)","DOI":"10.1109\/CVPR42600.2020.00722"},{"key":"1369_CR26","doi-asserted-by":"crossref","unstructured":"Yew, Z.J., Lee, G.H.: 3dfeat-net: Weakly supervised local 3d features for point cloud registration. In: European Conference on Computer Vision, pp. 607\u2013623 (2018)","DOI":"10.1007\/978-3-030-01267-0_37"},{"key":"1369_CR27","doi-asserted-by":"crossref","unstructured":"Bai, X., Luo, Z., Zhou, L., Fu, H., Quan, L., Tai, C.-L.: D3feat: Joint learning of dense detection and description of 3d local features. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 6359\u20136367 (2020)","DOI":"10.1109\/CVPR42600.2020.00639"},{"key":"1369_CR28","unstructured":"Steder, B., Rusu, R.B., Konolige, K., Burgard, W.: Narf: 3d range image features for object recognition. In: Workshop on Defining and Solving Realistic Perception Problems in Personal Robotics at the IEEE\/RSJ Int. Conf. on Intelligent Robots and Systems, vol. 44, p. 2 (2010)"},{"key":"1369_CR29","unstructured":"Wang, Y., Solomon, J.M.: Prnet: Self-supervised learning for partial-to-partial registration. Advances in Neural Information Processing Systems 32 (2019)"},{"key":"1369_CR30","doi-asserted-by":"crossref","unstructured":"Li, J., Zhang, C., Xu, Z., Zhou, H., Zhang, C.: Iterative distance-aware similarity matrix convolution with mutual-supervised point elimination for efficient point cloud registration. In: European Conference on Computer Vision, pp. 378\u2013394 (2020)","DOI":"10.1007\/978-3-030-58586-0_23"},{"key":"1369_CR31","doi-asserted-by":"crossref","unstructured":"Tombari, F., Salti, S., Di\u00a0Stefano, L.: Unique shape context for 3d data description. In: Proceedings of the ACM Workshop on 3D Object Retrieval, pp. 57\u201362 (2010)","DOI":"10.1145\/1877808.1877821"},{"key":"1369_CR32","doi-asserted-by":"crossref","unstructured":"Rusu, R.B., Blodow, N., Beetz, M.: Fast point feature histograms (fpfh) for 3d registration. In: IEEE International Conference on Robotics and Automation, pp. 3212\u20133217 (2009)","DOI":"10.1109\/ROBOT.2009.5152473"},{"key":"1369_CR33","doi-asserted-by":"crossref","unstructured":"Huang, S., Gojcic, Z., Usvyatsov, M., Wieser, A., Schindler, K.: Predator: Registration of 3d point clouds with low overlap. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 4267\u20134276 (2021)","DOI":"10.1109\/CVPR46437.2021.00425"},{"key":"1369_CR34","doi-asserted-by":"crossref","unstructured":"Ao, S., Hu, Q., Yang, B., Markham, A., Guo, Y.: Spinnet: Learning a general surface descriptor for 3d point cloud registration. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 11753\u201311762 (2021)","DOI":"10.1109\/CVPR46437.2021.01158"},{"key":"1369_CR35","doi-asserted-by":"crossref","unstructured":"Wang, H., Liu, Y., Dong, Z., Wang, W.: You only hypothesize once: Point cloud registration with rotation-equivariant descriptors. In: Proceedings of the 30th ACM International Conference on Multimedia, pp. 1630\u20131641 (2022)","DOI":"10.1145\/3503161.3548023"},{"key":"1369_CR36","doi-asserted-by":"crossref","unstructured":"Yew, Z.J., Lee, G.H.: Regtr: End-to-end point cloud correspondences with transformers. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 6677\u20136686 (2022)","DOI":"10.1109\/CVPR52688.2022.00656"},{"key":"1369_CR37","doi-asserted-by":"crossref","unstructured":"Yu, J., Ren, L., Zhang, Y., Zhou, W., Lin, L., Dai, G.: Peal: Prior-embedded explicit attention learning for low-overlap point cloud registration. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 17702\u201317711 (2023)","DOI":"10.1109\/CVPR52729.2023.01698"},{"key":"1369_CR38","doi-asserted-by":"crossref","unstructured":"Yu, H., Qin, Z., Hou, J., Saleh, M., Li, D., Busam, B., Ilic, S.: Rotation-invariant transformer for point cloud matching. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 5384\u20135393 (2023)","DOI":"10.1109\/CVPR52729.2023.00521"},{"key":"1369_CR39","doi-asserted-by":"crossref","unstructured":"Gao, J., Dong, Q., Wang, R., Chen, S., Xin, S., Tu, C., Wang, W.: Oaaformer: Robust and efficient point cloud registration through overlapping-aware attention in transformer. arXiv preprint arXiv:2310.09817 (2023)","DOI":"10.1007\/s11390-024-4165-6"},{"key":"1369_CR40","doi-asserted-by":"crossref","unstructured":"Wang, Y., Solomon, J.M.: Deep closest point: Learning representations for point cloud registration. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 3523\u20133532 (2019)","DOI":"10.1109\/ICCV.2019.00362"},{"key":"1369_CR41","doi-asserted-by":"crossref","unstructured":"Yew, Z.J., Lee, G.H.: Rpm-net: Robust point matching using learned features. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 11824\u201311833 (2020)","DOI":"10.1109\/CVPR42600.2020.01184"},{"key":"1369_CR42","doi-asserted-by":"crossref","unstructured":"Fu, K., Liu, S., Luo, X., Wang, M.: Robust point cloud registration framework based on deep graph matching. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 8893\u20138902 (2021)","DOI":"10.1109\/CVPR46437.2021.00878"},{"key":"1369_CR43","doi-asserted-by":"crossref","unstructured":"Aoki, Y., Goforth, H., Srivatsan, R.A., Lucey, S.: Pointnetlk: Robust & efficient point cloud registration using pointnet. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 7163\u20137172 (2019)","DOI":"10.1109\/CVPR.2019.00733"},{"key":"1369_CR44","doi-asserted-by":"crossref","unstructured":"Xu, H., Liu, S., Wang, G., Liu, G., Zeng, B.: Omnet: Learning overlapping mask for partial-to-partial point cloud registration. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 3132\u20133141 (2021)","DOI":"10.1109\/ICCV48922.2021.00312"},{"key":"1369_CR45","doi-asserted-by":"crossref","unstructured":"Huang, X., Mei, G., Zhang, J.: Feature-metric registration: A fast semi-supervised approach for robust point cloud registration without correspondences. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 11366\u201311374 (2020)","DOI":"10.1109\/CVPR42600.2020.01138"},{"key":"1369_CR46","doi-asserted-by":"crossref","unstructured":"Misra, I., Girdhar, R., Joulin, A.: An end-to-end transformer model for 3d object detection. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 2906\u20132917 (2021)","DOI":"10.1109\/ICCV48922.2021.00290"},{"key":"1369_CR47","doi-asserted-by":"crossref","unstructured":"Carion, N., Massa, F., Synnaeve, G., Usunier, N., Kirillov, A., Zagoruyko, S.: End-to-end object detection with transformers. In: European Conference on Computer Vision, pp. 213\u2013229 (2020)","DOI":"10.1007\/978-3-030-58452-8_13"},{"key":"1369_CR48","unstructured":"Dosovitskiy, A., Beyer, L., Kolesnikov, A., Weissenborn, D., Zhai, X., Unterthiner, T., Dehghani, M., Minderer, M., Heigold, G., Gelly, S., et al.: An image is worth 16x16 words: Transformers for image recognition at scale. arXiv preprint arXiv:2010.11929 (2020)"},{"key":"1369_CR49","doi-asserted-by":"crossref","unstructured":"Yu, X., Rao, Y., Wang, Z., Liu, Z., Lu, J., Zhou, J.: Pointr: Diverse point cloud completion with geometry-aware transformers. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 12498\u201312507 (2021)","DOI":"10.1109\/ICCV48922.2021.01227"},{"issue":"12","key":"1369_CR50","doi-asserted-by":"publisher","first-page":"15949","DOI":"10.1109\/TPAMI.2023.3311447","volume":"45","author":"J Gao","year":"2023","unstructured":"Gao, J., Chen, M., Xu, C.: Vectorized evidential learning for weakly-supervised temporal action localization. IEEE Trans Pattern Anal Mach Intell 45(12), 15949\u201315963 (2023)","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"1369_CR51","unstructured":"Yang, J., Li, C., Zhang, P., Dai, X., Xiao, B., Yuan, L., Gao, J.: Focal self-attention for local-global interactions in vision transformers. arXiv preprint arXiv:2107.00641 (2021)"},{"key":"1369_CR52","unstructured":"Zhu, X., Su, W., Lu, L., Li, B., Wang, X., Dai, J.: Deformable detr: Deformable transformers for end-to-end object detection. arXiv preprint arXiv:2010.04159 (2020)"},{"key":"1369_CR53","doi-asserted-by":"crossref","unstructured":"Zhou, H., Yu, J., Yang, W.: Dual memory units with uncertainty regulation for weakly supervised video anomaly detection. arXiv preprint arXiv:2302.05160 (2023)","DOI":"10.1609\/aaai.v37i3.25489"},{"key":"1369_CR54","doi-asserted-by":"crossref","unstructured":"Lin, T.-Y., Doll\u00e1r, P., Girshick, R., He, K., Hariharan, B., Belongie, S.: Feature pyramid networks for object detection. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 2117\u20132125 (2017)","DOI":"10.1109\/CVPR.2017.106"},{"key":"1369_CR55","doi-asserted-by":"crossref","unstructured":"Thomas, H., Qi, C.R., Deschaud, J.-E., Marcotegui, B., Goulette, F., Guibas, L.J.: Kpconv: Flexible and deformable convolution for point clouds. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 6411\u20136420 (2019)","DOI":"10.1109\/ICCV.2019.00651"},{"key":"1369_CR56","doi-asserted-by":"crossref","unstructured":"Zhao, H., Jiang, L., Jia, J., Torr, P.H., Koltun, V.: Point transformer. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 16259\u201316268 (2021)","DOI":"10.1109\/ICCV48922.2021.01595"},{"key":"1369_CR57","doi-asserted-by":"crossref","unstructured":"Sun, J., Shen, Z., Wang, Y., Bao, H., Zhou, X.: Loftr: Detector-free local feature matching with transformers. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 8922\u20138931 (2021)","DOI":"10.1109\/CVPR46437.2021.00881"},{"key":"1369_CR58","doi-asserted-by":"crossref","unstructured":"Peyr\u00e9, G., Cuturi, M., et al.: Computational optimal transport: With applications to data science. Foundations and Trends\u00ae in Machine Learning 11(5-6), 355\u2013607 (2019)","DOI":"10.1561\/2200000073"},{"key":"1369_CR59","doi-asserted-by":"crossref","unstructured":"Sarlin, P.-E., DeTone, D., Malisiewicz, T., Rabinovich, A.: Superglue: Learning feature matching with graph neural networks. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 4938\u20134947 (2020)","DOI":"10.1109\/CVPR42600.2020.00499"},{"key":"1369_CR60","doi-asserted-by":"crossref","unstructured":"Sun, Y., Cheng, C., Zhang, Y., Zhang, C., Zheng, L., Wang, Z., Wei, Y.: Circle loss: A unified perspective of pair similarity optimization. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 6398\u20136407 (2020)","DOI":"10.1109\/CVPR42600.2020.00643"},{"key":"1369_CR61","unstructured":"Kingma, D.P., Ba, J.: Adam: A method for stochastic optimization. arXiv preprint arXiv:1412.6980 (2014)"},{"key":"1369_CR62","unstructured":"Zhu, L., Guan, H., Lin, C., Han, R.: Leveraging inlier correspondences proportion for point cloud registration. arXiv preprint arXiv:2201.12094 (2022)"},{"key":"1369_CR63","unstructured":"Yu, H., Hou, J., Qin, Z., Saleh, M., Shugurov, I., Wang, K., Busam, B., Ilic, S.: Riga: Rotation-invariant and globally-aware descriptors for point cloud registration. arXiv preprint arXiv:2209.13252 (2022)"},{"key":"1369_CR64","doi-asserted-by":"crossref","unstructured":"Lu, F., Chen, G., Liu, Y., Zhang, L., Qu, S., Liu, S., Gu, R.: Hregnet: A hierarchical network for large-scale outdoor lidar point cloud registration. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 16014\u201316023 (2021)","DOI":"10.1109\/ICCV48922.2021.01571"},{"key":"1369_CR65","doi-asserted-by":"crossref","unstructured":"Ghiasi, G., Lin, T.-Y., Le, Q.V.: Nas-fpn: Learning scalable feature pyramid architecture for object detection. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 7036\u20137045 (2019)","DOI":"10.1109\/CVPR.2019.00720"},{"key":"1369_CR66","doi-asserted-by":"crossref","unstructured":"Tan, M., Pang, R., Le, Q.V.: Efficientdet: Scalable and efficient object detection. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 10781\u201310790 (2020)","DOI":"10.1109\/CVPR42600.2020.01079"},{"key":"1369_CR67","doi-asserted-by":"crossref","unstructured":"Liu, S., Qi, L., Qin, H., Shi, J., Jia, J.: Path aggregation network for instance segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 8759\u20138768 (2018)","DOI":"10.1109\/CVPR.2018.00913"},{"key":"1369_CR68","doi-asserted-by":"crossref","unstructured":"Hu, M., Li, Y., Fang, L., Wang, S.: A2-fpn: Attention aggregation based feature pyramid network for instance segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 15343\u201315352 (2021)","DOI":"10.1109\/CVPR46437.2021.01509"},{"key":"1369_CR69","doi-asserted-by":"crossref","unstructured":"Hu, Y., Gao, J., Dong, J., Fan, B., Liu, H.: Exploring rich semantics for open-set action recognition. IEEE Transactions on Multimedia 26 (2023)","DOI":"10.1109\/TMM.2023.3333206"},{"issue":"10","key":"1369_CR70","doi-asserted-by":"publisher","first-page":"3476","DOI":"10.1109\/TPAMI.2020.2985708","volume":"43","author":"J Gao","year":"2020","unstructured":"Gao, J., Zhang, T., Xu, C.: Learning to model relationships for zero-shot video classification. IEEE Trans. Pattern Anal. Mach. Intell. 43(10), 3476\u20133491 (2020)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"issue":"3","key":"1369_CR71","doi-asserted-by":"publisher","first-page":"1646","DOI":"10.1109\/TCSVT.2021.3075470","volume":"32","author":"J Gao","year":"2021","unstructured":"Gao, J., Xu, C.: Learning video moment retrieval without a single annotated video. IEEE Trans. Circ. Syst. Video Technol. 32(3), 1646\u20131657 (2021)","journal-title":"IEEE Trans. Circ. Syst. Video Technol."}],"container-title":["Multimedia Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-024-01369-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00530-024-01369-x\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-024-01369-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,11,20]],"date-time":"2024-11-20T21:11:07Z","timestamp":1732137067000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00530-024-01369-x"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,5,31]]},"references-count":71,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2024,6]]}},"alternative-id":["1369"],"URL":"https:\/\/doi.org\/10.1007\/s00530-024-01369-x","relation":{},"ISSN":["0942-4962","1432-1882"],"issn-type":[{"value":"0942-4962","type":"print"},{"value":"1432-1882","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,5,31]]},"assertion":[{"value":"27 November 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"21 May 2024","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"31 May 2024","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"164"}}