{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,20]],"date-time":"2026-07-20T17:03:19Z","timestamp":1784566999439,"version":"3.55.0"},"reference-count":90,"publisher":"Springer Science and Business Media LLC","issue":"6","license":[{"start":{"date-parts":[[2026,6,1]],"date-time":"2026-06-01T00:00:00Z","timestamp":1780272000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,6,1]],"date-time":"2026-06-01T00:00:00Z","timestamp":1780272000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"ELSA - European Lighthouse on Secure and Safe AI","award":["101070617"],"award-info":[{"award-number":["101070617"]}]},{"name":"ELSA - European Lighthouse on Secure and Safe AI","award":["101070617"],"award-info":[{"award-number":["101070617"]}]},{"name":"French project SIGHT","award":["ANR-20-CE23-0016"],"award-info":[{"award-number":["ANR-20-CE23-0016"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int J Comput Vis"],"published-print":{"date-parts":[[2026,6]]},"DOI":"10.1007\/s11263-026-02875-3","type":"journal-article","created":{"date-parts":[[2026,6,3]],"date-time":"2026-06-03T16:45:18Z","timestamp":1780505118000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Domain Adaptation with a Single Vision-Language Embedding"],"prefix":"10.1007","volume":"134","author":[{"given":"Mohammad","family":"Fahes","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Tuan-Hung","family":"Vu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Andrei","family":"Bursuc","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Patrick","family":"P\u00e9rez","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Raoul","family":"de Charette","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,6,3]]},"reference":[{"key":"2875_CR1","unstructured":"Arjovsky, M., Bottou, L., Gulrajani, I., & Lopez-Paz, D. (2019). Invariant risk minimization. arXiv preprint arXiv:1907.02893."},{"key":"2875_CR2","unstructured":"Bai, S., Chen, K., Liu, X., Wang, J., Ge, W., Song, S., Dang,K., Wang, P., Wang, S., Tang, J., et al. (2025). Qwen2. 5-vl technical report. arXiv preprint arXiv:2502.13923"},{"issue":"1","key":"2875_CR3","doi-asserted-by":"publisher","first-page":"151","DOI":"10.1007\/s10994-009-5152-4","volume":"79","author":"S Ben-David","year":"2010","unstructured":"Ben-David, S., Blitzer, J., Crammer, K., Kulesza, A., Pereira, F., & Vaughan, J. W. (2010). A theory of learning from different domains. Machine learning, 79(1), 151\u2013175.","journal-title":"Machine learning"},{"key":"2875_CR4","unstructured":"Bommasani, R., Hudson, DA., Adeli, E., Altman, R., Arora, S., von Arx, S., Bernstein, M.S., Bohg, J., Bosselut, A., Brunskill, E., & Brynjolfsson, E. (2021). On the opportunities and risks of foundation models. arXiv preprint arXiv:2108.07258"},{"key":"2875_CR5","doi-asserted-by":"crossref","unstructured":"Bottou, L. (2010). Large-scale machine learning with stochastic gradient descent. In: Computer Statistics or Comparative Statistics.","DOI":"10.1007\/978-3-7908-2604-3_16"},{"issue":"4","key":"2875_CR6","doi-asserted-by":"publisher","first-page":"834","DOI":"10.1109\/TPAMI.2017.2699184","volume":"40","author":"LC Chen","year":"2017","unstructured":"Chen, L. C., Papandreou, G., Kokkinos, I., Murphy, K., & Yuille, A. L. (2017). Deeplab: Semantic image segmentation with deep convolutional nets, atrous convolution, and fully connected crfs. IEEE transactions on pattern analysis and machine intelligence, 40(4), 834\u2013848.","journal-title":"IEEE transactions on pattern analysis and machine intelligence"},{"key":"2875_CR7","doi-asserted-by":"crossref","unstructured":"Chen, L.C., Zhu, Y., Papandreou, G., Schroff, F., & Adam, H. (2018). Encoder-decoder with atrous separable convolution for semantic image segmentation. In Proceedings of the European conference on computer vision (ECCV).","DOI":"10.1007\/978-3-030-01234-2_49"},{"key":"2875_CR8","doi-asserted-by":"crossref","unstructured":"Chen, Q., Wang, Y., Yang, T., Zhang, X., Cheng, J., & Sun, J. (2021a). You only look one-level feature. In:Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition.","DOI":"10.1109\/CVPR46437.2021.01284"},{"issue":"7","key":"2875_CR9","doi-asserted-by":"publisher","first-page":"2223","DOI":"10.1007\/s11263-021-01447-x","volume":"129","author":"Y Chen","year":"2021","unstructured":"Chen, Y., Wang, H., Li, W., Sakaridis, C., Dai, D., & Van Gool, L. (2021). Scale-aware domain adaptive faster r-cnn. International Journal of Computer Vision, 129(7), 2223\u20132243.","journal-title":"International Journal of Computer Vision"},{"key":"2875_CR10","doi-asserted-by":"crossref","unstructured":"Cheng, B., Collins, MD., Zhu Y, Liu, T., Huang, T.S., Adam, H., & Chen, L.C. (2020). Panoptic-deeplab: A simple, strong, and fast baseline for bottom-up panoptic segmentation. In:Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 12475-12485).","DOI":"10.1109\/CVPR42600.2020.01249"},{"key":"2875_CR11","doi-asserted-by":"crossref","unstructured":"Choi, S., Jung, S., Yun, H., Kim, J.T., Kim, S., & Choo, J. (2021). Robustnet: Improving domain generalization in urban-scene segmentation via instance selective whitening. In:Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 11580-11590).","DOI":"10.1109\/CVPR46437.2021.01141"},{"key":"2875_CR12","doi-asserted-by":"crossref","unstructured":"Cordts, M., Omran, M., Ramos, S., Rehfeld, T., Enzweiler, M., Benenson, R., Franke, U., Roth, S., & Schiele, B. (2016). The cityscapes dataset for semantic urban scene understanding. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 3213-3223).","DOI":"10.1109\/CVPR.2016.350"},{"key":"2875_CR13","doi-asserted-by":"crossref","unstructured":"Fahes, M., Vu, TH., Bursuc, A., P\u00e9rez, P., & De Charette, R. (2023). Poda: Prompt-driven zero-shot domain adaptation. In: Proceedings of the IEEE\/CVF international conference on computer vision (pp. 18623-18633).","DOI":"10.1109\/ICCV51070.2023.01707"},{"key":"2875_CR14","unstructured":"Fahes, M., Vu, TH., Bursuc, A., P\u00e9rez, P., & De Charette, R. (2024a). Clip\u2019s visual embedding projector is a few-shot cornucopia. arXiv preprint arXiv:2410.05270."},{"key":"2875_CR15","doi-asserted-by":"crossref","unstructured":"Fahes, M., Vu, TH., Bursuc, A., P\u00e9rez, P., & De Charette, R. (2024b). A simple recipe for language-guided domain generalized segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition.","DOI":"10.1109\/CVPR52733.2024.02211"},{"key":"2875_CR16","unstructured":"Fan, Q., Segu, M., Tai, YW., Yu, F., Tang, C.K., Schiele, B., & Dai, D. (2023). Towards robust object detection invariant to real-world domain shifts. In: The Eleventh International Conference on Learning Representations."},{"issue":"4","key":"2875_CR17","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1145\/3528223.3530164","volume":"41","author":"R Gal","year":"2022","unstructured":"Gal, R., Patashnik, O., Maron, H., Bermano, A. H., Chechik, G., & Cohen-Or, D. (2022). StyleGAN-NADA: CLIP-guided domain adaptation of image generators. ACM Transactions on Graphics (TOG), 41(4), 1\u201313.","journal-title":"ACM Transactions on Graphics (TOG)"},{"key":"2875_CR18","unstructured":"Ganin, Y., & Lempitsky, V. (2015). Unsupervised domain adaptation by backpropagation. In: International conference on machine learning (pp. 1180-1189)."},{"issue":"59","key":"2875_CR19","first-page":"1","volume":"17","author":"Y Ganin","year":"2016","unstructured":"Ganin, Y., Ustinova, E., Ajakan, H., Germain, P., Larochelle, H., Laviolette, F., March, M., & Lempitsky, V. (2016). Domain-adversarial training of neural networks. Journal of machine learning research, 17(59), 1\u201335.","journal-title":"Journal of machine learning research"},{"issue":"2","key":"2875_CR20","doi-asserted-by":"publisher","first-page":"581","DOI":"10.1007\/s11263-023-01891-x","volume":"132","author":"P Gao","year":"2024","unstructured":"Gao, P., Geng, S., Zhang, R., Ma, T., Fang, R., Zhang, Y., Li, H., & Qiao, Y. (2024). Clip-adapter: Better vision-language models with feature adapters. International journal of computer vision, 132(2), 581\u2013595.","journal-title":"International journal of computer vision"},{"key":"2875_CR21","doi-asserted-by":"crossref","unstructured":"Gatys, LA., Ecker, AS., Bethge, M. (2015). A neural algorithm of artistic style. arXiv preprint arXiv:1508.06576.","DOI":"10.1167\/16.12.326"},{"key":"2875_CR22","doi-asserted-by":"crossref","unstructured":"Gatys, L.A., Ecker, A.S., Bethge, M. (2016). Image style transfer using convolutional neural networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 2414-2423).","DOI":"10.1109\/CVPR.2016.265"},{"key":"2875_CR23","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., & Sun, J.(2016). Deep residual learning for image recognition. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 770-778).","DOI":"10.1109\/CVPR.2016.90"},{"key":"2875_CR24","unstructured":"Hoffman, J., Tzeng, E., Park, T., Zhu, J.Y., Isola, P., Saenko, K., Efros, A., & Darrell, T. (2018). Cycada: Cycle-consistent adversarial domain adaptation. In: International conference on machine learning (pp. 1989-1998). Pmlr."},{"key":"2875_CR25","doi-asserted-by":"crossref","unstructured":"Hoyer, L., Dai, D.,& Van\u00a0Gool, L. (2022). Daformer: Improving network architectures and training strategies for domain-adaptive semantic segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 9924-9935).","DOI":"10.1109\/CVPR52688.2022.00969"},{"key":"2875_CR26","doi-asserted-by":"crossref","unstructured":"Huang, X., & Belongie, S. (2017). Arbitrary style transfer in real-time with adaptive instance normalization. In: Proceedings of the IEEE international conference on computer vision (pp. 1501-1510).","DOI":"10.1109\/ICCV.2017.167"},{"key":"2875_CR27","unstructured":"Ioffe, S., & Szegedy, C. (2015). Batch normalization: Accelerating deep network training by reducing internal covariate shift. In: International conference on machine learning (pp. 448-456)."},{"key":"2875_CR28","doi-asserted-by":"crossref","unstructured":"Jatavallabhula, KM., Kuwajerwala, A., Gu, Q., Omama, M., Chen, T., Maalouf, A., Li, S., Iyer, G., Saryazdi, S., Keetha, N., & Tewari, A. (2023). Conceptfusion: Open-set multimodal 3d mapping. arXiv preprint arXiv:2302.07241","DOI":"10.15607\/RSS.2023.XIX.066"},{"key":"2875_CR29","unstructured":"Jia, C., Yang, Y., Xia, Y., Chen, Y.T., Parekh, Z., Pham, H., Le, Q., Sung, Y.H., Li, Z., & Duerig, T. (2021). Scaling up visual and vision-language representation learning with noisy text supervision. In: International conference on machine learning."},{"key":"2875_CR30","doi-asserted-by":"publisher","first-page":"423","DOI":"10.1162\/tacl_a_00324","volume":"8","author":"Z Jiang","year":"2020","unstructured":"Jiang, Z., Xu, F. F., Araki, J., & Neubig, G. (2020). How can we know what language models know? Transactions of the Association for Computational Linguistics, 8, 423\u2013438.","journal-title":"Transactions of the Association for Computational Linguistics"},{"key":"2875_CR31","doi-asserted-by":"crossref","unstructured":"Kang, G., Jiang, L., Yang, Y., & Hauptmann, A.G. (2019). Contrastive adaptation network for unsupervised domain adaptation. In:Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 4893-4902).","DOI":"10.1109\/CVPR.2019.00503"},{"key":"2875_CR32","doi-asserted-by":"crossref","unstructured":"Karras, T., Laine, S., & Aila, T. (2019). A style-based generator architecture for generative adversarial networks. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 4401-4410).","DOI":"10.1109\/CVPR.2019.00453"},{"key":"2875_CR33","doi-asserted-by":"crossref","unstructured":"Kirillov, A., Girshick, R., He, K., & Doll\u00e1r, P. (2019). Panoptic feature pyramid networks. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 6399-6408).","DOI":"10.1109\/CVPR.2019.00656"},{"key":"2875_CR34","doi-asserted-by":"crossref","unstructured":"Kirillov, A., Mintun, E., Ravi, N., Mao, H., Rolland, C., Gustafson, L., Xiao, T., Whitehead, S., Berg, A.C., Lo, W.Y., & Doll\u00e1r, P. (2023). Segment anything. In: Proceedings of the IEEE\/CVF international conference on computer vision (pp. 4015-4026).","DOI":"10.1109\/ICCV51070.2023.00371"},{"key":"2875_CR35","unstructured":"Krizhevsky, A., Sutskever, I., & Hinton, GE. (2012). Imagenet classification with deep convolutional neural networks. Advances in neural information processing systems 25."},{"key":"2875_CR36","doi-asserted-by":"crossref","unstructured":"Kwon, G., & Ye, J.C. (2022). Clipstyler: Image style transfer with a single text condition. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 18062-18071).","DOI":"10.1109\/CVPR52688.2022.01753"},{"key":"2875_CR37","doi-asserted-by":"crossref","unstructured":"Lee, S., Seong, H., Lee, S., & Kim, E. (2022). Wildnet: Learning domain generalized semantic segmentation from the wild. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 9936-9946).","DOI":"10.1109\/CVPR52688.2022.00970"},{"key":"2875_CR38","doi-asserted-by":"crossref","unstructured":"Lengyel, A., Garg, S., Milford, M., & Van Gemert, J.C. (2021). Zero-shot day-night domain adaptation with a physics prior. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision (pp. 4399-4409).","DOI":"10.1109\/ICCV48922.2021.00436"},{"key":"2875_CR39","unstructured":"Li, B., Weinberger, KQ., Belongie, S., Koltun, V., & Ranftl, R. (2022). Language-driven semantic segmentation. In: International Conference on Learning Representations (ICLR)."},{"key":"2875_CR40","doi-asserted-by":"crossref","unstructured":"Li, G., Kang, G., Liu, W., Wei, Y., & Yang, Y. (2020). Content-consistent matching for domain adaptive semantic segmentation. In: European conference on computer vision (pp. 440-456).","DOI":"10.1007\/978-3-030-58568-6_26"},{"key":"2875_CR41","first-page":"9694","volume":"34","author":"J Li","year":"2021","unstructured":"Li, J., Selvaraju, R., Gotmare, A., Joty, S., Xiong, C., & Hoi, S. C. H. (2021). Align before fuse: Vision and language representation learning with momentum distillation. Advances in neural information processing systems, 34, 9694\u20139705.","journal-title":"Advances in neural information processing systems"},{"key":"2875_CR42","doi-asserted-by":"crossref","unstructured":"Li, P., Li, D., Li, W., Gong, S., Fu, Y. & Hospedales, T.M. (2021b). A simple feature augmentation for domain generalization. In: Proceedings of the IEEE\/CVF international conference on computer vision (pp. 8886-8895).","DOI":"10.1109\/ICCV48922.2021.00876"},{"key":"2875_CR43","doi-asserted-by":"crossref","unstructured":"Li, Y., Yuan, L., & Vasconcelos, N. (2019). Bidirectional learning for domain adaptation of semantic segmentation. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 2117-2125).","DOI":"10.1109\/CVPR.2019.00710"},{"key":"2875_CR44","doi-asserted-by":"crossref","unstructured":"Lin, T.Y., Doll\u00e1r, P., Girshick, R., He, K., Hariharan, B., & Belongie, S. (2017). Feature pyramid networks for object detection. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 2117-2125).","DOI":"10.1109\/CVPR.2017.106"},{"key":"2875_CR45","doi-asserted-by":"crossref","unstructured":"Long, J., Shelhamer, E., & Darrell, T. (2015). Fully convolutional networks for semantic segmentation. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 3431-3440).","DOI":"10.1109\/CVPR.2015.7298965"},{"key":"2875_CR46","unstructured":"Long, M., Cao, Z., Wang, J.& Jordan, M.I. (2018). Conditional adversarial domain adaptation. Advances in neural information processing systems 31."},{"key":"2875_CR47","doi-asserted-by":"crossref","unstructured":"Lu, Y., Liu, J., Zhang, Y., Liu, Y., & Tian, X. (2022). Prompt distribution learning. In:Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 5206-5215).","DOI":"10.1109\/CVPR52688.2022.00514"},{"key":"2875_CR48","first-page":"20612","volume":"33","author":"Y Luo","year":"2020","unstructured":"Luo, Y., Liu, P., Guan, T., Yu, J., & Yang, Y. (2020). Adversarial style mining for one-shot unsupervised domain adaptation. Advances in neural information processing systems, 33, 20612\u201320623.","journal-title":"Advances in neural information processing systems"},{"key":"2875_CR49","doi-asserted-by":"crossref","unstructured":"Minderer, M., Gritsenko, A., Stone, A., Neumann, M., Weissenborn, D., Dosovitskiy, A., Mahendran, A., Arnab, A., Dehghani, M., Shen, Z., & Wang, X. (2022). Simple open-vocabulary object detection with vision transformers. In: European conference on computer vision (pp. 728-755).","DOI":"10.1007\/978-3-031-20080-9_42"},{"key":"2875_CR50","unstructured":"Ovadia, Y., Fertig, E., Ren, J., Nado, Z., Sculley, D., Nowozin, S., Dillon, J., Lakshminarayanan, B., & Snoek, J. (2019). Can you trust your model\u2019s uncertainty? evaluating predictive uncertainty under dataset shift. Advances in neural information processing systems 32."},{"key":"2875_CR51","doi-asserted-by":"crossref","unstructured":"Pan, F., Shin, I., Rameau, F., Lee, S., & Kweon, I.S. (2020). Unsupervised intra-domain adaptation for semantic segmentation through self-supervision. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 3764-3773).","DOI":"10.1109\/CVPR42600.2020.00382"},{"key":"2875_CR52","doi-asserted-by":"crossref","unstructured":"Pan, X., Luo, P., Shi, J.,& Tang, X.(2018). Two at once: Enhancing learning and generalization capacities via ibn-net. In: Proceedings of the european conference on computer vision (ECCV) (pp. 464-479).","DOI":"10.1007\/978-3-030-01225-0_29"},{"key":"2875_CR53","doi-asserted-by":"crossref","unstructured":"Patashnik, O., Wu, Z., Shechtman, E., Cohen-Or, D., & Lischinski, D. (2021). Styleclip: Text-driven manipulation of stylegan imagery. In: Proceedings of the IEEE\/CVF international conference on computer vision (pp. 2085-2094).","DOI":"10.1109\/ICCV48922.2021.00209"},{"key":"2875_CR54","doi-asserted-by":"crossref","unstructured":"Pizzati, F., Cerri, P., & de\u00a0Charette, R. (2021). Comogan: continuous model-guided image-to-image translation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (pp. 14288-14298).","DOI":"10.1109\/CVPR46437.2021.01406"},{"key":"2875_CR55","doi-asserted-by":"crossref","unstructured":"Pizzati, F., Lalonde, JF.,& de\u00a0Charette, R. (2022). Manifest: Manifold deformation for few-shot image translation. In: European Conference on Computer Vision (pp. 440-456).","DOI":"10.1007\/978-3-031-19790-1_27"},{"key":"2875_CR56","unstructured":"Radford, A., Kim, JW., Hallacy, C., Ramesh, A., Goh, G., Agarwal, S., Sastry, G., Askell, A., Mishkin, P., Clark, J., & Krueger, G. (2021). Learning transferable visual models from natural language supervision. In: International conference on machine learning (pp. 8748-8763)."},{"key":"2875_CR57","unstructured":"Ren, S., He, K., Girshick, R., & Sun, J. (2015). Faster r-cnn: Towards real-time object detection with region proposal networks. Advances in neural information processing systems 28."},{"key":"2875_CR58","doi-asserted-by":"crossref","unstructured":"Rezaeianaran, F., Shetty, R., Aljundi, R., Reino, D.O., Zhang, S., & Schiele, B. (2021). Seeking similarities over differences: Similarity-based domain alignment for adaptive object detection. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision (pp. 9204-9213).","DOI":"10.1109\/ICCV48922.2021.00907"},{"key":"2875_CR59","doi-asserted-by":"crossref","unstructured":"Richter SR, Vineet V, Roth S, & Koltun, V. (2016). Playing for data: Ground truth from computer games. In: European conference on computer vision (pp. 102-118).","DOI":"10.1007\/978-3-319-46475-6_7"},{"key":"2875_CR60","doi-asserted-by":"crossref","unstructured":"Ronneberger, O., Fischer, P., & Brox, T. (2015). U-net: Convolutional networks for biomedical image segmentation. In: International Conference on Medical image computing and computer-assisted intervention (pp. 234-241).","DOI":"10.1007\/978-3-319-24574-4_28"},{"issue":"9","key":"2875_CR61","doi-asserted-by":"publisher","first-page":"973","DOI":"10.1007\/s11263-018-1072-8","volume":"126","author":"C Sakaridis","year":"2018","unstructured":"Sakaridis, C., Dai, D., & Van Gool, L. (2018). Semantic foggy scene understanding with synthetic data. International Journal of Computer Vision, 126(9), 973\u2013992.","journal-title":"International Journal of Computer Vision"},{"key":"2875_CR62","doi-asserted-by":"crossref","unstructured":"Sakaridis, C., Dai, D., Van\u00a0Gool, L. (2021). Acdc: The adverse conditions dataset with correspondences for semantic driving scene understanding. In: Proceedings of the IEEE\/CVF international conference on computer vision (pp. 10765-10775).","DOI":"10.1109\/ICCV48922.2021.01059"},{"key":"2875_CR63","unstructured":"Santurkar, S., Tsipras, D., Ilyas, A., & Madry, A. (2018). How does batch normalization help optimization? Advances in neural information processing systems 31."},{"key":"2875_CR64","doi-asserted-by":"crossref","unstructured":"Shin, T., Razeghi, Y., Logan\u00a0IV, RL., Wallace, E., & Singh, S. (2020). Autoprompt: Eliciting knowledge from language models with automatically generated prompts. In: Proceedings of the 2020 conference on empirical methods in natural language processing (EMNLP) (pp. 4222-4235).","DOI":"10.18653\/v1\/2020.emnlp-main.346"},{"key":"2875_CR65","doi-asserted-by":"crossref","unstructured":"Sun, B., & Saenko, K. (2016). Deep coral: Correlation alignment for deep domain adaptation. In: European conference on computer vision (pp. 443-450).","DOI":"10.1007\/978-3-319-49409-8_35"},{"key":"2875_CR66","doi-asserted-by":"crossref","unstructured":"Sun, P., Kretzschmar, H., Dotiwalla, X., Chouard, A., Patnaik, V., Tsui, P., Guo, J., Zhou, Y., Chai, Y., Caine, B., & Vasudevan, V. (2020). Scalability in perception for autonomous driving: Waymo open dataset. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 2446-2454).","DOI":"10.1109\/CVPR42600.2020.00252"},{"key":"2875_CR67","doi-asserted-by":"crossref","unstructured":"Tsai, Y.H., Hung, W.C., Schulter, S., Sohn, K., Yang, M.H., & Chandraker, M. (2018). Learning to adapt structured output space for semantic segmentation. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 7472-7481).","DOI":"10.1109\/CVPR.2018.00780"},{"key":"2875_CR68","unstructured":"Ulyanov, D., Vedaldi, A., & Lempitsky, V. (2016). Instance normalization: The missing ingredient for fast stylization. arXiv preprint arXiv:1607.08022"},{"key":"2875_CR69","doi-asserted-by":"crossref","unstructured":"Ulyanov, D., Vedaldi, A., & Lempitsky, V. (2017). Improved texture networks: Maximizing quality and diversity in feed-forward stylization and texture synthesis. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 6924-6932).","DOI":"10.1109\/CVPR.2017.437"},{"key":"2875_CR70","doi-asserted-by":"crossref","unstructured":"Vidit, V., Engilberge, M., & Salzmann, M. (2023). Clip the gap: A single domain generalization approach for object detection. arXiv preprint arXiv:2301.05499","DOI":"10.1109\/CVPR52729.2023.00314"},{"key":"2875_CR71","doi-asserted-by":"crossref","unstructured":"Vu, T.H., Jain, H., Bucher, M., Cord, M., & P\u00e9rez, P. (2019). Advent: Adversarial entropy minimization for domain adaptation in semantic segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 2517-2526).","DOI":"10.1109\/CVPR.2019.00262"},{"key":"2875_CR72","unstructured":"Wah, C., Branson, S., Welinder, P., Perona, P., & Belongie, S. (2011). The caltech-ucsd birds-200-2011 dataset. California Institute of Technology."},{"key":"2875_CR73","unstructured":"Wang, D., Shelhamer, E., Liu, S., Olshausen, B., & Darrell, T. (2021a). Tent: Fully test-time adaptation by entropy minimization. In: International Conference on Learning Representations (ICLR)."},{"issue":"10","key":"2875_CR74","doi-asserted-by":"publisher","first-page":"3349","DOI":"10.1109\/TPAMI.2020.2983686","volume":"43","author":"J Wang","year":"2020","unstructured":"Wang, J., Sun, K., Cheng, T., Jiang, B., Deng, C., Zhao, Y., Liu, D., Mu, Y., Tan, M., Wang, X., & Liu, W. (2020). Deep high-resolution representation learning for visual recognition. IEEE transactions on pattern analysis and machine intelligence, 43(10), 3349\u20133364.","journal-title":"IEEE transactions on pattern analysis and machine intelligence"},{"key":"2875_CR75","doi-asserted-by":"crossref","unstructured":"Wang, S., Chen, X., Wang, Y., Long, M., & Wang, J. (2020b). Progressive adversarial networks for fine-grained domain adaptation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 9213-9222).","DOI":"10.1109\/CVPR42600.2020.00923"},{"key":"2875_CR76","doi-asserted-by":"crossref","unstructured":"Wang, Y., Li, W., Dai, D., Van & Gool, L. (2017). Deep domain adaptation by geodesic distance minimization. In: Proceedings of the IEEE international conference on computer vision workshops (pp. 2651-2657).","DOI":"10.1109\/ICCVW.2017.315"},{"key":"2875_CR77","doi-asserted-by":"crossref","unstructured":"Wang, Z., Luo, Y., Qiu, R., Huang, Z., & Baktashmotlagh, M. (2021b). Learning to diversify for single domain generalization. In: Proceedings of the IEEE\/CVF international conference on computer vision (pp. 834-843).","DOI":"10.1109\/ICCV48922.2021.00087"},{"key":"2875_CR78","doi-asserted-by":"crossref","unstructured":"Wu, A.,& Deng, C. (2022). Single-domain generalized object detection in urban scene via cyclic-disentangled self-distillation. In: Proceedings of the IEEE\/CVF Conference on computer vision and pattern recognition (pp. 847-856).","DOI":"10.1109\/CVPR52688.2022.00092"},{"key":"2875_CR79","doi-asserted-by":"crossref","unstructured":"Wu, X., Wu, Z., Lu, Y., Ju, L., & Wang, S. (2022). Style mixing and patchwise prototypical matching for one-shot unsupervised domain adaptive semantic segmentation. In: proceedings of the AAAI conference on artificial intelligenc 36(3), pp. 2740-2749.","DOI":"10.1609\/aaai.v36i3.20177"},{"key":"2875_CR80","doi-asserted-by":"crossref","unstructured":"Xiao, N., & Zhang, L .(2021). Dynamic weighted learning for unsupervised domain adaptation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 15242-15251).","DOI":"10.1109\/CVPR46437.2021.01499"},{"key":"2875_CR81","doi-asserted-by":"crossref","unstructured":"Yang, Y., Soatto, S. (2020). Fda: Fourier domain adaptation for semantic segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 4085-4095).","DOI":"10.1109\/CVPR42600.2020.00414"},{"key":"2875_CR82","doi-asserted-by":"crossref","unstructured":"Zhao, H., Shi, J., Qi, X., Wang, X., & Jia, J. (2017). Pyramid scene parsing network. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 2881-2890).","DOI":"10.1109\/CVPR.2017.660"},{"key":"2875_CR83","doi-asserted-by":"crossref","unstructured":"Zhao, H., Qi, X., Shen, X., Shi, J., & Jia, J. (2018). Icnet for real-time semantic segmentation on high-resolution images. In: Proceedings of the European conference on computer vision (ECCV) (pp. 405-420).","DOI":"10.1007\/978-3-030-01219-9_25"},{"key":"2875_CR84","doi-asserted-by":"crossref","unstructured":"Zhong, Z., Friedman, D., & Chen, D. (2021). Factual probing is [mask]: Learning vs. learning to recall. In: Proceedings of the 2021 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies (pp. 5017-5033).","DOI":"10.18653\/v1\/2021.naacl-main.398"},{"key":"2875_CR85","doi-asserted-by":"crossref","unstructured":"Zhou, C., Loy, C. C., & Dai, B. (2022). Extract free dense labels from clip. European conference on computer vision (pp. 696\u2013712). Springer Nature Switzerland: Cham.","DOI":"10.1007\/978-3-031-19815-1_40"},{"key":"2875_CR86","unstructured":"Zhou, K., Yang, Y., Qiao, Y., & Xiang, T. (2021). Domain generalization with mixstyle. In: International Conference on Learning Representations."},{"key":"2875_CR87","doi-asserted-by":"crossref","unstructured":"Zhou, K., Yang, J., Loy, C.C., & Liu, Z. (2022b). Conditional prompt learning for vision-language models. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 16816-16825).","DOI":"10.1109\/CVPR52688.2022.01631"},{"issue":"9","key":"2875_CR88","doi-asserted-by":"publisher","first-page":"2337","DOI":"10.1007\/s11263-022-01653-1","volume":"130","author":"K Zhou","year":"2022","unstructured":"Zhou, K., Yang, J., Loy, C. C., & Liu, Z. (2022). Learning to prompt for vision-language models. International journal of computer vision, 130(9), 2337\u20132348.","journal-title":"International journal of computer vision"},{"key":"2875_CR89","doi-asserted-by":"crossref","unstructured":"Zou, Y., Yu, Z., Kumar, B.V.K., & Wang, J. (2018). Unsupervised domain adaptation for semantic segmentation via class-balanced self-training. In: Proceedings of the European conference on computer vision (ECCV) (pp. 289-305).","DOI":"10.1007\/978-3-030-01219-9_18"},{"key":"2875_CR90","doi-asserted-by":"crossref","unstructured":"Zou, Y., Yu, Z., Liu, X., Kumar, B.V.K., & Wang, J. (2019). Confidence regularized self-training. In: Proceedings of the IEEE\/CVF international conference on computer vision (pp. 5982-5991).","DOI":"10.1109\/ICCV.2019.00608"}],"container-title":["International Journal of Computer Vision"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-026-02875-3.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11263-026-02875-3","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-026-02875-3.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,7,20]],"date-time":"2026-07-20T16:15:31Z","timestamp":1784564131000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11263-026-02875-3"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,6]]},"references-count":90,"journal-issue":{"issue":"6","published-print":{"date-parts":[[2026,6]]}},"alternative-id":["2875"],"URL":"https:\/\/doi.org\/10.1007\/s11263-026-02875-3","relation":{},"ISSN":["0920-5691","1573-1405"],"issn-type":[{"value":"0920-5691","type":"print"},{"value":"1573-1405","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,6]]},"assertion":[{"value":"26 October 2024","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"27 April 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"3 June 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}}],"article-number":"305"}}