{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,22]],"date-time":"2026-06-22T03:47:55Z","timestamp":1782100075085,"version":"3.54.5"},"reference-count":48,"publisher":"Springer Science and Business Media LLC","issue":"1","license":[{"start":{"date-parts":[[2026,5,2]],"date-time":"2026-05-02T00:00:00Z","timestamp":1777680000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,5,2]],"date-time":"2026-05-02T00:00:00Z","timestamp":1777680000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Geoinformatica"],"published-print":{"date-parts":[[2026,6]]},"DOI":"10.1007\/s10707-026-00575-1","type":"journal-article","created":{"date-parts":[[2026,5,2]],"date-time":"2026-05-02T09:03:15Z","timestamp":1777712595000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["From satellite to street: a hybrid framework integrating stable diffusion and PanoGAN for consistent cross\u2013view synthesis"],"prefix":"10.1007","volume":"30","author":[{"given":"Khawlah","family":"Bajbaa","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Abbas","family":"Anwar","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Muhammad","family":"Saqib","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Hafeez","family":"Anwar","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Nabin","family":"Sharma","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Muhammad","family":"Usman","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,5,2]]},"reference":[{"key":"575_CR1","doi-asserted-by":"crossref","unstructured":"Lu X, Li Z, Cui Z, Oswald M.R, Pollefeys M, Qin R (2020) Geometry-aware satellite-to-ground image synthesis for urban areas. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 859\u2013867","DOI":"10.1109\/CVPR42600.2020.00094"},{"key":"575_CR2","doi-asserted-by":"publisher","DOI":"10.1016\/j.landurbplan.2021.104217","volume":"215","author":"F Biljecki","year":"2021","unstructured":"Biljecki F (2021) Ito K Street view imagery in urban analytics and gis: A review. Landsc Urban Plan 215:104217","journal-title":"Landsc Urban Plan"},{"issue":"12","key":"575_CR3","doi-asserted-by":"publisher","first-page":"10009","DOI":"10.1109\/TPAMI.2022.3140750","volume":"44","author":"Y Shi","year":"2022","unstructured":"Shi Y, Campbell D, Yu X (2022) Li H Geometry-guided street-view panorama synthesis from satellite imagery. IEEE Trans Pattern Anal Mach Intell 44(12):10009\u201310022","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"575_CR4","doi-asserted-by":"publisher","first-page":"3546","DOI":"10.1109\/TMM.2022.3162474","volume":"25","author":"S Wu","year":"2022","unstructured":"Wu S, Tang H, Jing X-Y, Zhao H, Qian J, Sebe N (2022) Yan Y Cross-view panorama image synthesis. IEEE Trans Multimedia 25:3546\u20133559","journal-title":"IEEE Trans Multimedia"},{"key":"575_CR5","doi-asserted-by":"crossref","unstructured":"Toker A, Zhou Q, Maximov M, Leal-Taix\u00e9 L (2021) Coming down to earth: Satellite-to-street view synthesis for geo-localization. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 6488\u20136497","DOI":"10.1109\/CVPR46437.2021.00642"},{"issue":"2","key":"575_CR6","doi-asserted-by":"publisher","first-page":"867","DOI":"10.1109\/TCSVT.2021.3061265","volume":"32","author":"T Wang","year":"2021","unstructured":"Wang T, Zheng Z, Yan C, Zhang J, Sun Y, Zheng B, Yang Y (2021) Each part matters: Local patterns facilitate cross-view geo-localization. IEEE Trans Circ Syst Video Technol 32(2):867\u2013879","journal-title":"IEEE Trans Circ Syst Video Technol"},{"issue":"9","key":"575_CR7","doi-asserted-by":"publisher","first-page":"3300","DOI":"10.3390\/s21093300","volume":"21","author":"B Kang","year":"2021","unstructured":"Kang B, Lee S, Zou S (2021) Developing sidewalk inventory data using street view images. Sensors 21(9):3300","journal-title":"Sensors"},{"issue":"11","key":"575_CR8","doi-asserted-by":"publisher","first-page":"10330","DOI":"10.1109\/TVT.2018.2865836","volume":"67","author":"M Cheng","year":"2018","unstructured":"Cheng M, Zhang Y, Su Y, Alvarez JM, Kong H (2018) Curb detection for road and sidewalk detection. IEEE Trans Veh Technol 67(11):10330\u201310342","journal-title":"IEEE Trans Veh Technol"},{"issue":"1","key":"575_CR9","doi-asserted-by":"publisher","first-page":"7","DOI":"10.1177\/2399808321995817","volume":"49","author":"H Ning","year":"2022","unstructured":"Ning H, Ye X, Chen Z, Liu T, Cao T (2022) Sidewalk extraction using aerial and street view images. Environ Plan B Urban Anal City Sci 49(1):7\u201322","journal-title":"Environ Plan B Urban Anal City Sci"},{"key":"575_CR10","doi-asserted-by":"publisher","unstructured":"Z\u00fcnd D, Bettencourt LM (2021) Street view imaging for automated assessments of urban infrastructure and services. Springer, Singapore, pp 29\u201340. https:\/\/doi.org\/10.1007\/978-981-15-8983-6_4","DOI":"10.1007\/978-981-15-8983-6_4"},{"key":"575_CR11","doi-asserted-by":"crossref","unstructured":"Xia D, Huang J, Yang J, Liu X, Wang H (2022) Duarus: Automatic geo-object change detection with street-view imagery for updating road database at baidu maps. In: Proceedings of the 31st ACM international conference on information & knowledge management, pp 3565\u20133574","DOI":"10.1145\/3511808.3557118"},{"key":"575_CR12","doi-asserted-by":"publisher","unstructured":"Lyu F, Ma X, Song Y, Zhu E, Wang S (2023) Large-scale google street view images for urban change detection. https:\/\/doi.org\/10.5703\/1288284317674","DOI":"10.5703\/1288284317674"},{"issue":"2","key":"575_CR13","doi-asserted-by":"publisher","first-page":"0263775","DOI":"10.1371\/journal.pone.0263775","volume":"17","author":"G Byun","year":"2022","unstructured":"Byun G, Kim Y (2022) A street-view-based method to detect urban growth and decline: A case study of midtown in detroit, michigan, usa. PLoS ONE 17(2):0263775","journal-title":"PLoS ONE"},{"key":"575_CR14","doi-asserted-by":"crossref","unstructured":"Deng X, Zhu Y, Newsam S (2018) What is it like down there? generating dense ground-level views and image features from overhead imagery using conditional generative adversarial networks. In: Proceedings of the 26th ACM SIGSPATIAL international conference on advances in geographic information systems, pp 43\u201352","DOI":"10.1145\/3274895.3274969"},{"key":"575_CR15","doi-asserted-by":"crossref","unstructured":"Ren B, Tang H, Sebe N, et al (2021) Cascaded cross mlp-mixer gans for cross-view image translation. In: British machine vision conference (BMVC\u201921). BMVA, pp 1\u201314","DOI":"10.5244\/C.35.40"},{"key":"575_CR16","doi-asserted-by":"publisher","first-page":"108884","DOI":"10.1016\/j.patcog.2022.108884","volume":"131","author":"S Wu","year":"2022","unstructured":"Wu S, Tang H, Jing X-Y, Qian J, Sebe N, Yan Y, Zhang Q (2022) Cross-view panorama image synthesis with progressive attention gans. Pattern Recogn 131:108884","journal-title":"Pattern Recogn"},{"key":"575_CR17","doi-asserted-by":"crossref","unstructured":"Regmi K, Borji A (2018) Cross-view image synthesis using conditional gans. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3501\u20133510","DOI":"10.1109\/CVPR.2018.00369"},{"key":"575_CR18","doi-asserted-by":"crossref","unstructured":"Tang H, Xu D, Yan Y, Torr P.H, Sebe N (2020) Local class-specific and global image-level generative adversarial networks for semantic-guided scene generation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 7870\u20137879","DOI":"10.1109\/CVPR42600.2020.00789"},{"key":"575_CR19","doi-asserted-by":"crossref","unstructured":"Tang H, Xu D, Sebe N, Wang Y, Corso JJ, Yan Y (2019) Multi-channel attention selection gan with cascaded semantic guidance for cross-view image translation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 2417\u20132426","DOI":"10.1109\/CVPR.2019.00252"},{"key":"575_CR20","doi-asserted-by":"crossref","unstructured":"Regmi K, Shah M (2019) Bridging the domain gap for ground-to-aerial image matching. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 470\u2013479","DOI":"10.1109\/ICCV.2019.00056"},{"key":"575_CR21","doi-asserted-by":"crossref","unstructured":"Workman S, Souvenir R, Jacobs N (2015) Wide-area image geolocalization with aerial reference imagery. In: Proceedings of the IEEE international conference on computer vision (ICCV), pp 3961\u20133969","DOI":"10.1109\/ICCV.2015.451"},{"key":"575_CR22","doi-asserted-by":"crossref","unstructured":"Shi Y, Yu X, Campbell D, Li H (2020) Where am i looking at? joint location and orientation estimation by cross-view matching. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 4064\u20134072","DOI":"10.1109\/CVPR42600.2020.00412"},{"key":"575_CR23","doi-asserted-by":"crossref","unstructured":"Vo NN, Hays J (2016) Localizing and orienting street views using overhead imagery. In: Computer vision\u2013ECCV 2016: 14th European conference, Amsterdam, The Netherlands, October 11\u201314, 2016, Proceedings, Part I 14. Springer, pp 494\u2013509","DOI":"10.1007\/978-3-319-46448-0_30"},{"key":"575_CR24","doi-asserted-by":"crossref","unstructured":"Lin T.-Y, Cui Y, Belongie S, Hays J (2015) Learning deep representations for ground-to-aerial geolocalization. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR), pp 5007\u20135015","DOI":"10.1109\/CVPR.2015.7299135"},{"key":"575_CR25","doi-asserted-by":"crossref","unstructured":"Hu S, Feng M, Nguyen RM, Lee GH (2018) Cvm-net: Cross-view matching network for image-based ground-to-aerial geo-localization. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 7258\u20137267","DOI":"10.1109\/CVPR.2018.00758"},{"key":"575_CR26","doi-asserted-by":"crossref","unstructured":"Liu L, Li H (2019) Lending orientation to neural networks for cross-view geo-localization. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 5624\u20135633","DOI":"10.1109\/CVPR.2019.00577"},{"key":"575_CR27","doi-asserted-by":"crossref","unstructured":"Isola P, Zhu J-Y, Zhou T, Efros AA (2017) Image-to-image translation with conditional adversarial networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 1125\u20131134","DOI":"10.1109\/CVPR.2017.632"},{"key":"575_CR28","doi-asserted-by":"crossref","unstructured":"Wang T-C, Liu M-Y, Zhu J-Y, Tao A, Kautz J, Catanzaro B (2018) High-resolution image synthesis and semantic manipulation with conditional gans. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 8798\u20138807","DOI":"10.1109\/CVPR.2018.00917"},{"key":"575_CR29","doi-asserted-by":"crossref","unstructured":"Zhu J-Y, Park T, Isola P, Efros AA (2017) Unpaired image-to-image translation using cycle-consistent adversarial networks. In: Proceedings of the IEEE international conference on computer vision, pp 2223\u20132232","DOI":"10.1109\/ICCV.2017.244"},{"key":"575_CR30","unstructured":"Goodfellow I, Pouget-Abadie J, Mirza M, Xu B, Warde-Farley D, Ozair S, Courville A, Bengio Y (2014) Generative adversarial nets. Adv Neural Inf Process Syst 27"},{"key":"575_CR31","unstructured":"Mirza M, Osindero S (2014) Conditional generative adversarial nets. arXiv preprint arXiv:1411.1784"},{"key":"575_CR32","doi-asserted-by":"crossref","unstructured":"Deng X, Zhu Y, Newsam S (2018) What is it like down there? generating dense ground-level views and image features from overhead imagery using conditional generative adversarial networks. In: Proceedings of the 26th ACM SIGSPATIAL international conference on advances in geographic information systems, pp 43\u201352","DOI":"10.1145\/3274895.3274969"},{"key":"575_CR33","doi-asserted-by":"crossref","unstructured":"Zhu X, Yin Z, Shi J, Li H, Lin D (2018) Generative adversarial frontal view to bird view synthesis. In: 2018 International conference on 3D vision (3DV). IEEE, pp 454\u2013463","DOI":"10.1109\/3DV.2018.00059"},{"key":"575_CR34","doi-asserted-by":"crossref","unstructured":"Qian M, Xiong J, Xia G-S, Xue N (2023) Sat2density: Faithful density learning from satellite-ground image pairs. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 3683\u20133692","DOI":"10.1109\/ICCV51070.2023.00341"},{"key":"575_CR35","doi-asserted-by":"crossref","unstructured":"Liu L, Li H (2019) Lending orientation to neural networks for cross-view geo-localization. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 5624\u20135633","DOI":"10.1109\/CVPR.2019.00577"},{"key":"575_CR36","doi-asserted-by":"crossref","unstructured":"Zhai M, Bessinger Z, Workman S, Jacobs N (2017) Predicting ground-level scene layout from aerial imagery. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 867\u2013875","DOI":"10.1109\/CVPR.2017.440"},{"issue":"8","key":"575_CR37","doi-asserted-by":"publisher","first-page":"5625","DOI":"10.1109\/TPAMI.2024.3369699","volume":"46","author":"J Zhang","year":"2024","unstructured":"Zhang J, Huang J, Jin S, Lu S (2024) Vision-language models for vision tasks: a survey. IEEE Trans Pattern Anal Mach Intell 46(8):5625\u20135644","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"575_CR38","doi-asserted-by":"crossref","unstructured":"Rombach R, Blattmann A, Lorenz D, Esser P, Ommer B (2022) High-resolution image synthesis with latent diffusion models. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 10684\u201310695","DOI":"10.1109\/CVPR52688.2022.01042"},{"issue":"9","key":"575_CR39","doi-asserted-by":"publisher","first-page":"10850","DOI":"10.1109\/TPAMI.2023.3261988","volume":"45","author":"F-A Croitoru","year":"2023","unstructured":"Croitoru F-A, Hondru V, Ionescu RT, Shah M (2023) Diffusion models in vision: a survey. IEEE Trans Pattern Anal Mach Intell 45(9):10850\u201310869","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"575_CR40","doi-asserted-by":"crossref","unstructured":"Ronneberger O, Fischer P, Brox T (2015) U-net: convolutional networks for biomedical image segmentation. In: Medical image computing and computer-assisted intervention\u2013MICCAI 2015: 18th International Conference, Munich, Germany, October 5-9, 2015, Proceedings, Part III 18. Springer, pp 234\u2013241","DOI":"10.1007\/978-3-319-24574-4_28"},{"key":"575_CR41","unstructured":"Radford A, Kim JW, Hallacy C, Ramesh A, Goh G, Agarwal S, Sastry G, Askell A, Mishkin P, Clark J, et al (2021) Learning transferable visual models from natural language supervision. In: International conference on machine learning. PMLR, pp 8748\u20138763"},{"key":"575_CR42","unstructured":"Kingma DP, Welling M (2014) Auto-encoding variational bayes stat 1050:1"},{"key":"575_CR43","doi-asserted-by":"crossref","unstructured":"Zhang L, Rao A, Agrawala M (2023) Adding conditional control to text-to-image diffusion models. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 3836\u20133847","DOI":"10.1109\/ICCV51070.2023.00355"},{"key":"575_CR44","unstructured":"Loshchilov I (2017) Decoupled weight decay regularization. arXiv preprint arXiv:1711.05101"},{"key":"575_CR45","doi-asserted-by":"crossref","unstructured":"Regmi K, Borji A (2018) Cross-view image synthesis using conditional gans. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3501\u20133510","DOI":"10.1109\/CVPR.2018.00369"},{"key":"575_CR46","unstructured":"Li W, He J, Ye J, Zhong H, Zheng Z, Huang Z, Lin D, He C (2024) Crossviewdiff: A cross-view diffusion model for satellite-to-street view synthesis. arXiv preprint arXiv:2408.14765"},{"issue":"9","key":"575_CR47","doi-asserted-by":"publisher","first-page":"1477","DOI":"10.3390\/rs16091477","volume":"16","author":"Y Bazi","year":"2024","unstructured":"Bazi Y, Bashmal L, Al Rahhal MM, Ricci R, Melgani F (2024) Rs-llava: A large vision-language model for joint captioning and question answering in remote sensing imagery. Remote Sensing 16(9):1477","journal-title":"Remote Sensing"},{"key":"575_CR48","doi-asserted-by":"crossref","unstructured":"Rotstein N, Bensaid D, Brody S, Ganz R, Kimmel R (2024) Fusecap: Leveraging large language models for enriched fused image captions. In: Proceedings of the IEEE\/CVF winter conference on applications of computer vision, pp 5689\u20135700","DOI":"10.1109\/WACV57701.2024.00559"}],"container-title":["GeoInformatica"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10707-026-00575-1.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10707-026-00575-1","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10707-026-00575-1.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,6,22]],"date-time":"2026-06-22T02:58:51Z","timestamp":1782097131000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10707-026-00575-1"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,5,2]]},"references-count":48,"journal-issue":{"issue":"1","published-print":{"date-parts":[[2026,6]]}},"alternative-id":["575"],"URL":"https:\/\/doi.org\/10.1007\/s10707-026-00575-1","relation":{},"ISSN":["1384-6175","1573-7624"],"issn-type":[{"value":"1384-6175","type":"print"},{"value":"1573-7624","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,5,2]]},"assertion":[{"value":"6 December 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"26 March 2026","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"26 April 2026","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"2 May 2026","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"6 May 2026","order":6,"name":"change_date","label":"Change Date","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"Update","order":7,"name":"change_type","label":"Change Type","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"The second and third affiliations have been swapped to reflect the correct information.","order":8,"name":"change_details","label":"Change Details","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no competing interests.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Competing interests"}}],"article-number":"17"}}