{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,11,3]],"date-time":"2025-11-03T21:26:13Z","timestamp":1762205173478,"version":"build-2065373602"},"publisher-location":"New York, NY, USA","reference-count":32,"publisher":"ACM","funder":[{"name":"National Natural Science Foundation of China &#x28;NSFC&#x29;","award":["62225207, 62436008, 62422609 and 62276243"],"award-info":[{"award-number":["62225207, 62436008, 62422609 and 62276243"]}]}],"content-domain":{"domain":["dl.acm.org"],"crossmark-restriction":true},"short-container-title":[],"published-print":{"date-parts":[[2025,10,27]]},"DOI":"10.1145\/3728482.3757387","type":"proceedings-article","created":{"date-parts":[[2025,11,3]],"date-time":"2025-11-03T21:21:35Z","timestamp":1762204895000},"page":"15-20","update-policy":"https:\/\/doi.org\/10.1145\/crossmark-policy","source":"Crossref","is-referenced-by-count":0,"title":["Geometry-Aware Enhancement and Data Augmentation for Street-to-Satellite Geo-localization"],"prefix":"10.1145","author":[{"ORCID":"https:\/\/orcid.org\/0009-0007-5098-7447","authenticated-orcid":false,"given":"Xingbo","family":"Wang","sequence":"first","affiliation":[{"name":"University of Science and Technology of China, Anhui, China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0009-0006-1878-604X","authenticated-orcid":false,"given":"Yongchao","family":"Xu","sequence":"additional","affiliation":[{"name":"University of Science and Technology of China, Anhui, China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0009-0003-9404-0440","authenticated-orcid":false,"given":"Yuanfei","family":"Bao","sequence":"additional","affiliation":[{"name":"University of Science and Technology of China, Anhui, China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-8036-4071","authenticated-orcid":false,"given":"Xueyang","family":"Fu","sequence":"additional","affiliation":[{"name":"University of Science and Technology of China, Anhui, China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-2510-8993","authenticated-orcid":false,"given":"Zheng-Jun","family":"Zha","sequence":"additional","affiliation":[{"name":"University of Science and Technology of China, Anhui, China"}],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"320","published-online":{"date-parts":[[2025,11,3]]},"reference":[{"key":"e_1_3_2_1_1_1","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2017.2711011"},{"key":"e_1_3_2_1_2_1","doi-asserted-by":"publisher","DOI":"10.1145\/2072298.2071954"},{"key":"e_1_3_2_1_3_1","doi-asserted-by":"publisher","DOI":"10.1016\/j.cviu.2007.09.014"},{"key":"e_1_3_2_1_4_1","doi-asserted-by":"publisher","DOI":"10.1109\/ICCVW.2015.137"},{"key":"e_1_3_2_1_5_1","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV51070.2023.01545"},{"key":"e_1_3_2_1_6_1","doi-asserted-by":"publisher","DOI":"10.1177\/0278364913491297"},{"key":"e_1_3_2_1_7_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR52733.2024.01582"},{"key":"e_1_3_2_1_8_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR52733.2024.02619"},{"key":"e_1_3_2_1_9_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2013.120"},{"key":"e_1_3_2_1_10_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2015.7299135"},{"key":"e_1_3_2_1_11_1","unstructured":"Tsung-Yi Lin Piotr Doll\u00e1r Ross Girshick Kaiming He Bharath Hariharan and Serge Belongie. 2017. Feature Pyramid Networks for Object Detection. arXiv:1612.03144 [cs.CV] https:\/\/arxiv.org\/abs\/1612.03144"},{"key":"e_1_3_2_1_12_1","volume-title":"Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition. IEEE Computer Society","author":"Liu Liu","year":"2019","unstructured":"Liu Liu and Hongdong Li. 2019. Lending orientation to neural networks for crossview geo-localization. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition. IEEE Computer Society, Long Beach, CA, USA, 5624--5633."},{"key":"e_1_3_2_1_13_1","doi-asserted-by":"publisher","DOI":"10.1023\/B:VISI.0000029664.99615.94"},{"key":"e_1_3_2_1_14_1","doi-asserted-by":"crossref","unstructured":"Krishna Regmi and Ali Borji. 2018. Cross-View Image Synthesis using Conditional GANs. arXiv:1803.03396 [cs.CV] https:\/\/arxiv.org\/abs\/1803.03396","DOI":"10.1109\/CVPR.2018.00369"},{"key":"e_1_3_2_1_15_1","doi-asserted-by":"publisher","DOI":"10.1109\/ICCVW.2011.6130498"},{"key":"e_1_3_2_1_16_1","doi-asserted-by":"publisher","DOI":"10.1109\/TCSVT.2023.3296074"},{"key":"e_1_3_2_1_17_1","volume-title":"Advances in Neural Information Processing Systems, H. Wallach, H. Larochelle, A. Beygelzimer, F. d'Alch\u00e9-Buc","author":"Shi Yujiao","year":"2019","unstructured":"Yujiao Shi, Liu Liu, Xin Yu, and Hongdong Li. 2019. Spatial-Aware Feature Aggregation for Image based Cross-View Geo-Localization. In Advances in Neural Information Processing Systems, H. Wallach, H. Larochelle, A. Beygelzimer, F. d'Alch\u00e9-Buc, E. Fox, and R. Garnett (Eds.), Vol. 32. Curran Associates, Inc., Red Hook, NY, USA. https:\/\/proceedings.neurips.cc\/paper_files\/paper\/2019\/file\/ba2f0015122a5955f8b3a50240fb91b2-Paper.pdf"},{"key":"e_1_3_2_1_18_1","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2022.3189702"},{"key":"e_1_3_2_1_19_1","unstructured":"Yujiao Shi Xin Yu Liu Liu Dylan Campbell Piotr Koniusz and Hongdong Li. 2022. Accurate 3-DoF Camera Geo-Localization via Ground-to-Satellite Image Matching. arXiv:2203.14148 [cs.CV] https:\/\/arxiv.org\/abs\/2203.14148"},{"key":"e_1_3_2_1_20_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2017.216"},{"key":"e_1_3_2_1_21_1","volume-title":"The 3rd Workshop on UAVs in Multimedia: Capturing the World from a New Perspective. In Proceedings of the 33rd ACM International Conference on Multimedia Workshop.","author":"Wang Tingyu","year":"2025","unstructured":"Tingyu Wang, Yujiao Shi, Fabian Deuser, Shaofei Huang, Guosheng Hu, Si Liu, Zhedong Zheng, and Roger Zimmermann. 2025. The 3rd Workshop on UAVs in Multimedia: Capturing the World from a New Perspective. In Proceedings of the 33rd ACM International Conference on Multimedia Workshop."},{"key":"e_1_3_2_1_22_1","volume-title":"The 3rd Workshop on UAVs in Multimedia: Capturing the World from a New Perspective. In Proceedings of the 33rd ACM International Conference on Multimedia Workshop.","author":"Wang Tingyu","year":"2025","unstructured":"Tingyu Wang, Yujiao Shi, Fabian Deuser, Shaofei Huang, Guosheng Hu, Si Liu, Zhedong Zheng, and Roger Zimmermann. 2025. The 3rd Workshop on UAVs in Multimedia: Capturing the World from a New Perspective. In Proceedings of the 33rd ACM International Conference on Multimedia Workshop."},{"key":"e_1_3_2_1_23_1","doi-asserted-by":"publisher","DOI":"10.1109\/TCSVT.2021.3061265"},{"key":"e_1_3_2_1_24_1","unstructured":"TingyuWang Zhedong Zheng Zunjie Zhu Yuhan Gao Yi Yang and Chenggang Yan. 2022. Learning Cross-view Geo-localization Embeddings via Dynamic Weighted Decorrelation Regularization. arXiv:2211.05296 [cs.CV] https:\/\/arxiv.org\/abs\/2211.05296"},{"key":"e_1_3_2_1_25_1","unstructured":"Xiaolong Wang Runsen Xu Zuofan Cui Zeyu Wan and Yu Zhang. 2023. Fine-Grained Cross-View Geo-Localization Using a Correlation-Aware Homography Estimator. arXiv:2308.16906 [cs.CV] https:\/\/arxiv.org\/abs\/2308.16906"},{"key":"e_1_3_2_1_26_1","volume-title":"Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition Workshops. IEEE Computer Society","author":"Nathan Jacobs ScottWorkman","year":"2015","unstructured":"ScottWorkman and Nathan Jacobs. 2015. On the location dependence of convolutional neural network features. In Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition Workshops. IEEE Computer Society, Boston, MA, USA, 70--78."},{"key":"e_1_3_2_1_27_1","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV.2015.451"},{"key":"e_1_3_2_1_28_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2017.440"},{"key":"e_1_3_2_1_29_1","doi-asserted-by":"publisher","DOI":"10.1145\/3394171.3413896"},{"key":"e_1_3_2_1_30_1","unstructured":"Pengfei Zhu LongyinWen Xiao Bian Haibin Ling and Qinghua Hu. 2018. Vision Meets Drones: A Challenge. arXiv:1804.07437 [cs.CV]"},{"key":"e_1_3_2_1_31_1","doi-asserted-by":"publisher","DOI":"10.1109\/TCSVT.2023.3249204"},{"key":"e_1_3_2_1_32_1","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR52688.2022.00123"}],"event":{"name":"MM '25:The 33rd ACM International Conference on Multimedia","sponsor":["SIGMM ACM Special Interest Group on Multimedia"],"location":"Dublin Ireland"},"container-title":["Proceedings of the 3rd International Workshop on UAVs in Multimedia: Capturing the World from a New Perspective"],"original-title":[],"link":[{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/3728482.3757387","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,11,3]],"date-time":"2025-11-03T21:22:17Z","timestamp":1762204937000},"score":1,"resource":{"primary":{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/3728482.3757387"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,10,27]]},"references-count":32,"alternative-id":["10.1145\/3728482.3757387","10.1145\/3728482"],"URL":"https:\/\/doi.org\/10.1145\/3728482.3757387","relation":{},"subject":[],"published":{"date-parts":[[2025,10,27]]},"assertion":[{"value":"2025-11-03","order":3,"name":"published","label":"Published","group":{"name":"publication_history","label":"Publication History"}}]}}