{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,8,1]],"date-time":"2026-08-01T16:49:01Z","timestamp":1785602941790,"version":"3.56.0"},"publisher-location":"New York, NY, USA","reference-count":84,"publisher":"ACM","license":[{"start":{"date-parts":[[2024,10,28]],"date-time":"2024-10-28T00:00:00Z","timestamp":1730073600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.acm.org\/publications\/policies\/copyright_policy#Background"}],"funder":[{"name":"the National Natural Science Foundation of China Grant","award":["12071478,61972404"],"award-info":[{"award-number":["12071478,61972404"]}]},{"name":"Public Computing Cloud, Renmin University of China"},{"name":"the Blockchain Lab, School of Information, Renmin University of China"}],"content-domain":{"domain":["dl.acm.org"],"crossmark-restriction":true},"short-container-title":[],"published-print":{"date-parts":[[2024,10,28]]},"DOI":"10.1145\/3664647.3681209","type":"proceedings-article","created":{"date-parts":[[2024,10,26]],"date-time":"2024-10-26T06:59:41Z","timestamp":1729925981000},"page":"5250-5259","update-policy":"https:\/\/doi.org\/10.1145\/crossmark-policy","source":"Crossref","is-referenced-by-count":6,"title":["PRISM: PRogressive dependency maxImization for Scale-invariant image Matching"],"prefix":"10.1145","author":[{"ORCID":"https:\/\/orcid.org\/0009-0009-8051-1081","authenticated-orcid":false,"given":"Xudong","family":"Cai","sequence":"first","affiliation":[{"name":"Renmin University of China &amp; HAOMO.AI, Beijing, China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-4197-2258","authenticated-orcid":false,"given":"Yongcai","family":"Wang","sequence":"additional","affiliation":[{"name":"Renmin University of China, Beijing, China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-0531-9171","authenticated-orcid":false,"given":"Lun","family":"Luo","sequence":"additional","affiliation":[{"name":"HAOMO.AI, Beijing, China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-5772-1042","authenticated-orcid":false,"given":"Minhang","family":"Wang","sequence":"additional","affiliation":[{"name":"HAOMO.AI, Beijing, China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-7748-5427","authenticated-orcid":false,"given":"Deying","family":"Li","sequence":"additional","affiliation":[{"name":"Renmin University of China, Beijing, China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0001-9365-9562","authenticated-orcid":false,"given":"Jintao","family":"Xu","sequence":"additional","affiliation":[{"name":"HAOMO.AI, Beijing, China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-8718-0977","authenticated-orcid":false,"given":"Weihao","family":"Gu","sequence":"additional","affiliation":[{"name":"HAOMO.AI, Beijing, China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-3224-008X","authenticated-orcid":false,"given":"Rui","family":"Ai","sequence":"additional","affiliation":[{"name":"HAOMO.AI, Beijing, China"}],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"320","published-online":{"date-parts":[[2024,10,28]]},"reference":[{"key":"e_1_3_2_1_1_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2017.410"},{"key":"e_1_3_2_1_2_1","doi-asserted-by":"publisher","DOI":"10.1109\/72.298224"},{"key":"e_1_3_2_1_3_1","doi-asserted-by":"publisher","DOI":"10.1016\/j.cviu.2007.09.014"},{"key":"e_1_3_2_1_4_1","volume-title":"International conference on machine learning. PMLR, 531--540","author":"Belghazi Mohamed Ishmael","year":"2018","unstructured":"Mohamed Ishmael Belghazi, Aristide Baratin, Sai Rajeshwar, Sherjil Ozair, Yoshua Bengio, Aaron Courville, and Devon Hjelm. 2018. Mutual information neural estimation. In International conference on machine learning. PMLR, 531--540."},{"key":"e_1_3_2_1_5_1","doi-asserted-by":"publisher","DOI":"10.5555\/2503308.2188387"},{"key":"e_1_3_2_1_6_1","doi-asserted-by":"publisher","DOI":"10.48550\/ARXIV.2205.14756"},{"key":"e_1_3_2_1_7_1","volume-title":"VOLoc: Visual Place Recognition by Querying Compressed Lidar Map. arXiv preprint arXiv:2402.15961","author":"Cai Xudong","year":"2024","unstructured":"Xudong Cai, Yongcai Wang, Zhe Huang, Yu Shao, and Deying Li. 2024. VOLoc: Visual Place Recognition by Querying Compressed Lidar Map. arXiv preprint arXiv:2402.15961 (2024)."},{"key":"e_1_3_2_1_8_1","volume-title":"Improving Transformer-based Image Matching by Cascaded Capturing Spatially Informative Keypoints. ArXiv","author":"Cao Chenjie","year":"2023","unstructured":"Chenjie Cao and Yanwei Fu. 2023. Improving Transformer-based Image Matching by Cascaded Capturing Spatially Informative Keypoints. ArXiv, Vol. abs\/2303.02885 (2023). https:\/\/api.semanticscholar.org\/CorpusID:257365749"},{"key":"e_1_3_2_1_9_1","volume-title":"End-to-End Object Detection with Transformers. ArXiv","author":"Carion Nicolas","year":"2020","unstructured":"Nicolas Carion, Francisco Massa, Gabriel Synnaeve, Nicolas Usunier, Alexander Kirillov, and Sergey Zagoruyko. 2020. End-to-End Object Detection with Transformers. ArXiv, Vol. abs\/2005.12872 (2020). https:\/\/api.semanticscholar.org\/CorpusID:218889832"},{"key":"e_1_3_2_1_10_1","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV48922.2021.00624"},{"key":"e_1_3_2_1_11_1","volume-title":"ASpanFormer: Detector-Free Image Matching with Adaptive Span Transformer. In European Conference on Computer Vision.","author":"Chen Hongkai","year":"2022","unstructured":"Hongkai Chen, Zixin Luo, Lei Zhou, Yurun Tian, Mingmin Zhen, Tian Fang, David N. R. McKinnon, Yanghai Tsin, and Long Quan. 2022. ASpanFormer: Detector-Free Image Matching with Adaptive Span Transformer. In European Conference on Computer Vision."},{"key":"e_1_3_2_1_12_1","volume-title":"HCPM: Hierarchical Candidates Pruning for Efficient Detector-Free Matching. arXiv preprint arXiv:2403.12543","author":"Chen Ying","year":"2024","unstructured":"Ying Chen, Yong Liu, Kai Wu, Qiang Nie, Shang Xu, Huifang Ma, Bing Wang, and Chengjie Wang. 2024. HCPM: Hierarchical Candidates Pruning for Efficient Detector-Free Matching. arXiv preprint arXiv:2403.12543 (2024)."},{"key":"e_1_3_2_1_13_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2017.261"},{"key":"e_1_3_2_1_14_1","unstructured":"Tri Dao Daniel Y. Fu Stefano Ermon Atri Rudra and Christopher R\u00e9. 2022. FlashAttention: Fast and Memory-Efficient Exact Attention with IO-Awareness. In Advances in Neural Information Processing Systems."},{"key":"e_1_3_2_1_15_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPRW.2018.00060"},{"key":"e_1_3_2_1_16_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2019.00828"},{"key":"e_1_3_2_1_17_1","doi-asserted-by":"publisher","DOI":"10.1109\/TNN.2008.2005601"},{"key":"e_1_3_2_1_18_1","volume-title":"Chen Hu, and Shuchang Zhou.","author":"Fan Miao","year":"2023","unstructured":"Miao Fan, Ming lei Chen, Chen Hu, and Shuchang Zhou. 2023. Occ2Net: Robust Image Matching Based on 3D Occupancy Estimation for Occluded Regions. ArXiv, Vol. abs\/2308.16160 (2023). https:\/\/api.semanticscholar.org\/CorpusID:261339755"},{"key":"e_1_3_2_1_19_1","doi-asserted-by":"publisher","DOI":"10.1016\/J.PATREC.2018.06.005"},{"key":"e_1_3_2_1_20_1","unstructured":"Khang Truong Giang Soohwan Song and Sung-Guk Jo. 2022. TopicFM: Robust and Interpretable Feature Matching with Topic-assisted. ArXiv Vol. abs\/2207.00328 (2022). https:\/\/api.semanticscholar.org\/CorpusID:250243816"},{"key":"e_1_3_2_1_21_1","doi-asserted-by":"publisher","DOI":"10.1007\/S11063-019--10144--3"},{"key":"e_1_3_2_1_22_1","volume-title":"Alvey Vision Conference. https:\/\/api.semanticscholar.org\/CorpusID:1694378","author":"Harris Christopher G.","unstructured":"Christopher G. Harris and M. J. Stephens. 1988. A Combined Corner and Edge Detector. In Alvey Vision Conference. https:\/\/api.semanticscholar.org\/CorpusID:1694378"},{"key":"e_1_3_2_1_23_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2016.90"},{"key":"e_1_3_2_1_24_1","volume-title":"Adaptive Assignment for Geometry Aware Local Feature Matching. 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","author":"Huang Dihe","year":"2022","unstructured":"Dihe Huang, Ying Chen, Shang Xu, Yong Liu, Wen-Qi Wu, Yikang Ding, Chengjie Wang, and Fan Tang. 2022. Adaptive Assignment for Geometry Aware Local Feature Matching. 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2022), 5425--5434."},{"key":"e_1_3_2_1_25_1","doi-asserted-by":"publisher","DOI":"10.1109\/IJCNN54540.2023.10191728"},{"key":"e_1_3_2_1_26_1","doi-asserted-by":"publisher","DOI":"10.1016\/j.aei.2023.101971"},{"key":"e_1_3_2_1_27_1","volume-title":"COTR: Correspondence Transformer for Matching Across Images. 2021 IEEE\/CVF International Conference on Computer Vision (ICCV)","author":"Jiang Wei","year":"2021","unstructured":"Wei Jiang, Eduard Trulls, Jan Hendrik Hosang, Andrea Tagliasacchi, and Kwang Moo Yi. 2021. COTR: Correspondence Transformer for Matching Across Images. 2021 IEEE\/CVF International Conference on Computer Vision (ICCV) (2021), 6187--6197. https:\/\/api.semanticscholar.org\/CorpusID:232380104"},{"key":"e_1_3_2_1_28_1","volume-title":"International Conference on Machine Learning.","author":"Katharopoulos Angelos","year":"2020","unstructured":"Angelos Katharopoulos, Apoorv Vyas, Nikolaos Pappas, and Franccois Fleuret. 2020. Transformers are RNNs: Fast Autoregressive Transformers with Linear Attention. In International Conference on Machine Learning."},{"key":"e_1_3_2_1_29_1","volume-title":"Information theory and statistics","author":"Kullback Solomon","unstructured":"Solomon Kullback. 1997. Information theory and statistics. Courier Corporation."},{"key":"e_1_3_2_1_30_1","unstructured":"Viktor Larsson and contributors. 2020. PoseLib - Minimal Solvers for Camera Pose Estimation. https:\/\/github.com\/vlarsson\/PoseLib"},{"key":"e_1_3_2_1_31_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR52688.2022.01538"},{"key":"e_1_3_2_1_32_1","volume-title":"Dual-Resolution Correspondence Networks. ArXiv","author":"Li Xinghui","year":"2020","unstructured":"Xinghui Li, K. Han, Shuda Li, and Victor Adrian Prisacariu. 2020. Dual-Resolution Correspondence Networks. ArXiv, Vol. abs\/2006.08844 (2020). https:\/\/api.semanticscholar.org\/CorpusID:219708544"},{"key":"e_1_3_2_1_33_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2018.00218"},{"key":"e_1_3_2_1_34_1","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV51070.2023.00991"},{"key":"e_1_3_2_1_35_1","volume-title":"Feature Pyramid Networks for Object Detection. 2017 IEEE Conference on Computer Vision and Pattern Recognition (CVPR)","author":"Lin Tsung-Yi","year":"2016","unstructured":"Tsung-Yi Lin, Piotr Doll\u00e1r, Ross B. Girshick, Kaiming He, Bharath Hariharan, and Serge J. Belongie. 2016. Feature Pyramid Networks for Object Detection. 2017 IEEE Conference on Computer Vision and Pattern Recognition (CVPR) (2016), 936--944. https:\/\/api.semanticscholar.org\/CorpusID:10716717"},{"key":"e_1_3_2_1_36_1","doi-asserted-by":"crossref","unstructured":"Philipp Lindenberger Paul-Edouard Sarlin and Marc Pollefeys. 2023. LightGlue: Local Feature Matching at Light Speed. In ICCV.","DOI":"10.1109\/ICCV51070.2023.01616"},{"key":"e_1_3_2_1_37_1","volume-title":"Decoupled Weight Decay Regularization. In 7th International Conference on Learning Representations, ICLR 2019","author":"Loshchilov Ilya","year":"2019","unstructured":"Ilya Loshchilov and Frank Hutter. 2019. Decoupled Weight Decay Regularization. In 7th International Conference on Learning Representations, ICLR 2019, New Orleans, LA, USA, May 6--9, 2019. OpenReview.net. https:\/\/openreview.net\/forum?id=Bkg6RiCqY7"},{"key":"e_1_3_2_1_38_1","doi-asserted-by":"publisher","DOI":"10.1023\/B:VISI.0000029664.99615.94"},{"key":"e_1_3_2_1_39_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2019.00263"},{"key":"e_1_3_2_1_40_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR42600.2020.00662"},{"key":"e_1_3_2_1_41_1","doi-asserted-by":"publisher","DOI":"10.1017\/CBO9780511809071"},{"key":"e_1_3_2_1_42_1","doi-asserted-by":"publisher","DOI":"10.1007\/978--3-031--19815--1_8"},{"key":"e_1_3_2_1_43_1","doi-asserted-by":"publisher","DOI":"10.1109\/TIT.1954.1057469"},{"key":"e_1_3_2_1_44_1","volume-title":"PATS: Patch Area Transportation with Subdivision for Local Feature Matching. 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","author":"Ni Junjie","year":"2023","unstructured":"Junjie Ni, Yijin Li, Zhaoyang Huang, Hongsheng Li, Hujun Bao, Zhaopeng Cui, and Guofeng Zhang. 2023. PATS: Patch Area Transportation with Subdivision for Local Feature Matching. 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2023), 17776--17786."},{"key":"e_1_3_2_1_45_1","doi-asserted-by":"publisher","DOI":"10.1109\/TIP.2021.3092828"},{"key":"e_1_3_2_1_46_1","volume-title":"Estimation of entropy and mutual information. Neural computation","author":"Paninski Liam","year":"2003","unstructured":"Liam Paninski. 2003. Estimation of entropy and mutual information. Neural computation, Vol. 15, 6 (2003), 1191--1253."},{"key":"e_1_3_2_1_47_1","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2005.159"},{"key":"e_1_3_2_1_48_1","volume-title":"No'e Pion, Gabriela Csurka, Yohann Cabon, and M. Humenberger.","author":"Revaud J\u00e9r\u00f4me","year":"2019","unstructured":"J\u00e9r\u00f4me Revaud, Philippe Weinzaepfel, C\u00e9sar Roberto de Souza, No'e Pion, Gabriela Csurka, Yohann Cabon, and M. Humenberger. 2019. R2D2: Repeatable and Reliable Detector and Descriptor. ArXiv, Vol. abs\/1906.06195 (2019). https:\/\/api.semanticscholar.org\/CorpusID:189927786"},{"key":"e_1_3_2_1_49_1","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-030-58545-7_35"},{"key":"e_1_3_2_1_50_1","volume-title":"Neighbourhood consensus networks. Advances in neural information processing systems","author":"Rocco Ignacio","year":"2018","unstructured":"Ignacio Rocco, Mircea Cimpoi, Relja Arandjelovi\u0107, Akihiko Torii, Tomas Pajdla, and Josef Sivic. 2018. Neighbourhood consensus networks. Advances in neural information processing systems, Vol. 31 (2018)."},{"key":"e_1_3_2_1_51_1","doi-asserted-by":"publisher","DOI":"10.1007\/11744023_34"},{"key":"e_1_3_2_1_52_1","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV.2011.6126544"},{"key":"e_1_3_2_1_53_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2019.01300"},{"key":"e_1_3_2_1_54_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2019.01300"},{"key":"e_1_3_2_1_55_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR46437.2021.00326"},{"key":"e_1_3_2_1_56_1","volume-title":"SuperGlue: Learning Feature Matching With Graph Neural Networks. 2020 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2019","author":"Sarlin Paul-Edouard","year":"2019","unstructured":"Paul-Edouard Sarlin, Daniel DeTone, Tomasz Malisiewicz, and Andrew Rabinovich. 2019. SuperGlue: Learning Feature Matching With Graph Neural Networks. 2020 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2019), 4937--4946. https:\/\/api.semanticscholar.org\/CorpusID:208291327"},{"key":"e_1_3_2_1_57_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2018.00897"},{"key":"e_1_3_2_1_58_1","volume-title":"Structure-from-Motion Revisited. In Conference on Computer Vision and Pattern Recognition (CVPR).","author":"Sch\u00f6nberger Johannes Lutz","year":"2016","unstructured":"Johannes Lutz Sch\u00f6nberger and Jan-Michael Frahm. 2016. Structure-from-Motion Revisited. In Conference on Computer Vision and Pattern Recognition (CVPR)."},{"key":"e_1_3_2_1_59_1","volume-title":"Pixelwise View Selection for Unstructured Multi-View Stereo. In European Conference on Computer Vision (ECCV).","author":"Sch\u00f6nberger Johannes Lutz","year":"2016","unstructured":"Johannes Lutz Sch\u00f6nberger, Enliang Zheng, Marc Pollefeys, and Jan-Michael Frahm. 2016. Pixelwise View Selection for Unstructured Multi-View Stereo. In European Conference on Computer Vision (ECCV)."},{"key":"e_1_3_2_1_60_1","volume-title":"ClusterGNN: Cluster-based Coarse-to-Fine Graph Neural Network for Efficient Feature Matching. 2022 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","author":"Shi Yanxing","year":"2022","unstructured":"Yanxing Shi, Junxiong Cai, Yoli Shavit, Tai-Jiang Mu, Wensen Feng, and Kai Zhang. 2022. ClusterGNN: Cluster-based Coarse-to-Fine Graph Neural Network for Efficient Feature Matching. 2022 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2022)."},{"key":"e_1_3_2_1_61_1","volume-title":"Very deep convolutional networks for large-scale image recognition. arXiv preprint arXiv:1409.1556","author":"Simonyan Karen","year":"2014","unstructured":"Karen Simonyan and Andrew Zisserman. 2014. Very deep convolutional networks for large-scale image recognition. arXiv preprint arXiv:1409.1556 (2014)."},{"key":"e_1_3_2_1_62_1","volume-title":"RoFormer: Enhanced Transformer with Rotary Position Embedding. ArXiv","author":"Su Jianlin","year":"2021","unstructured":"Jianlin Su, Yu Lu, Shengfeng Pan, Bo Wen, and Yunfeng Liu. 2021. RoFormer: Enhanced Transformer with Rotary Position Embedding. ArXiv, Vol. abs\/2104.09864 (2021). https:\/\/api.semanticscholar.org\/CorpusID:233307138"},{"key":"e_1_3_2_1_63_1","volume-title":"LoFTR: Detector-Free Local Feature Matching with Transformers. 2021 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","author":"Sun Jiaming","year":"2021","unstructured":"Jiaming Sun, Zehong Shen, Yuang Wang, Hujun Bao, and Xiaowei Zhou. 2021. LoFTR: Detector-Free Local Feature Matching with Transformers. 2021 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2021), 8918--8927."},{"key":"e_1_3_2_1_64_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2015.7298594"},{"key":"e_1_3_2_1_65_1","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2019.2952114"},{"key":"e_1_3_2_1_66_1","volume-title":"ECO-TR: Efficient Correspondences Finding Via Coarse-to-Fine Refinement. ArXiv","author":"Tan Donglin","year":"2022","unstructured":"Donglin Tan, Jiangjiang Liu, Xingyu Chen, Chao Chen, Ruixin Zhang, Yunhang Shen, Shouhong Ding, and Rongrong Ji. 2022. ECO-TR: Efficient Correspondences Finding Via Coarse-to-Fine Refinement. ArXiv, Vol. abs\/2209.12213 (2022). https:\/\/api.semanticscholar.org\/CorpusID:252531344"},{"key":"e_1_3_2_1_67_1","volume-title":"QuadTree Attention for Vision Transformers. ArXiv","author":"Tang Shitao","year":"2022","unstructured":"Shitao Tang, Jiahui Zhang, Siyu Zhu, and Ping Tan. 2022. QuadTree Attention for Vision Transformers. ArXiv, Vol. abs\/2201.02767 (2022)."},{"key":"e_1_3_2_1_68_1","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2020.3032010"},{"key":"e_1_3_2_1_69_1","first-page":"14254","article-title":"DISK: Learning local features with policy gradient","volume":"33","author":"Tyszkiewicz Micha\u0142","year":"2020","unstructured":"Micha\u0142 Tyszkiewicz, Pascal Fua, and Eduard Trulls. 2020. DISK: Learning local features with policy gradient. Advances in Neural Information Processing Systems, Vol. 33 (2020), 14254--14265.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"e_1_3_2_1_70_1","unstructured":"Ashish Vaswani Noam M. Shazeer Niki Parmar Jakob Uszkoreit Llion Jones Aidan N. Gomez Lukasz Kaiser and Illia Polosukhin. 2017. Attention is All you Need. In Neural Information Processing Systems."},{"key":"e_1_3_2_1_71_1","volume-title":"A review of feature selection methods based on mutual information. Neural computing and applications","author":"Vergara Jorge R","year":"2014","unstructured":"Jorge R Vergara and Pablo A Est\u00e9vez. 2014. A review of feature selection methods based on mutual information. Neural computing and applications, Vol. 24 (2014), 175--186."},{"key":"e_1_3_2_1_72_1","doi-asserted-by":"publisher","DOI":"10.1109\/TKDE.2017.2650906"},{"key":"e_1_3_2_1_73_1","volume-title":"Proceedings, Part I 16","author":"Wang Qianqian","year":"2020","unstructured":"Qianqian Wang, Xiaowei Zhou, Bharath Hariharan, and Noah Snavely. 2020. Learning feature descriptors using camera pose supervision. In Computer Vision--ECCV 2020: 16th European Conference, Glasgow, UK, August 23--28, 2020, Proceedings, Part I 16. Springer, 757--774."},{"key":"e_1_3_2_1_74_1","doi-asserted-by":"publisher","DOI":"10.1109\/JSAC.2023.3242708"},{"key":"e_1_3_2_1_75_1","doi-asserted-by":"publisher","DOI":"10.3390\/s23052399"},{"key":"e_1_3_2_1_76_1","doi-asserted-by":"crossref","unstructured":"Yifan Wang Xingyi He Sida Peng Dongli Tan and Xiaowei Zhou. 2024. Efficient LoFTR: Semi-Dense Local Feature Matching with Sparse-Like Speed. In CVPR.","DOI":"10.1109\/CVPR52733.2024.02047"},{"key":"e_1_3_2_1_77_1","volume-title":"Proceedings, Part VI 14","author":"Yi Kwang Moo","year":"2016","unstructured":"Kwang Moo Yi, Eduard Trulls, Vincent Lepetit, and Pascal Fua. 2016. Lift: Learned invariant feature transform. In Computer Vision--ECCV 2016: 14th European Conference, Amsterdam, The Netherlands, October 11--14, 2016, Proceedings, Part VI 14. Springer, 467--483."},{"key":"e_1_3_2_1_78_1","volume-title":"Adaptive Spot-Guided Transformer for Consistent Local Feature Matching. 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2023","author":"Yu Jiahuan","year":"2023","unstructured":"Jiahuan Yu, Jiahao Chang, Jianfeng He, Tianzhu Zhang, and Feng Wu. 2023. Adaptive Spot-Guided Transformer for Consistent Local Feature Matching. 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2023), 21898--21908."},{"key":"e_1_3_2_1_79_1","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV.2019.00594"},{"key":"e_1_3_2_1_80_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR52733.2024.01911"},{"key":"e_1_3_2_1_81_1","volume-title":"Searching from area to point: A hierarchical framework for semantic-geometric combined feature matching. arXiv preprint arXiv:2305.00194","author":"Zhang Yesheng","year":"2023","unstructured":"Yesheng Zhang, Xu Zhao, and Dahong Qian. 2023. Searching from area to point: A hierarchical framework for semantic-geometric combined feature matching. arXiv preprint arXiv:2305.00194 (2023)."},{"key":"e_1_3_2_1_82_1","doi-asserted-by":"publisher","DOI":"10.1007\/S11263-020-01399--8"},{"key":"e_1_3_2_1_83_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR46437.2021.00464"},{"key":"e_1_3_2_1_84_1","volume-title":"PMatch: Paired Masked Image Modeling for Dense Geometric Matching. 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2023","author":"Zhu Shengjie","year":"2023","unstructured":"Shengjie Zhu and Xiaoming Liu. 2023. PMatch: Paired Masked Image Modeling for Dense Geometric Matching. 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2023), 21909--21918. https:\/\/api.semanticscholar.org\/CorpusID:257833673"}],"event":{"name":"MM '24: The 32nd ACM International Conference on Multimedia","location":"Melbourne VIC Australia","acronym":"MM '24","sponsor":["SIGMM ACM Special Interest Group on Multimedia"]},"container-title":["Proceedings of the 32nd ACM International Conference on Multimedia"],"original-title":[],"link":[{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/3664647.3681209","content-type":"unspecified","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/3664647.3681209","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,6,19]],"date-time":"2025-06-19T01:18:02Z","timestamp":1750295882000},"score":1,"resource":{"primary":{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/3664647.3681209"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,10,28]]},"references-count":84,"alternative-id":["10.1145\/3664647.3681209","10.1145\/3664647"],"URL":"https:\/\/doi.org\/10.1145\/3664647.3681209","relation":{},"subject":[],"published":{"date-parts":[[2024,10,28]]},"assertion":[{"value":"2024-10-28","order":3,"name":"published","label":"Published","group":{"name":"publication_history","label":"Publication History"}}]}}