{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,4,22]],"date-time":"2026-04-22T20:19:28Z","timestamp":1776889168699,"version":"3.51.2"},"reference-count":73,"publisher":"Springer Science and Business Media LLC","issue":"12","license":[{"start":{"date-parts":[[2024,6,24]],"date-time":"2024-06-24T00:00:00Z","timestamp":1719187200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,6,24]],"date-time":"2024-06-24T00:00:00Z","timestamp":1719187200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100010909","name":"Excellent Young Scientists Fund","doi-asserted-by":"publisher","award":["62022078"],"award-info":[{"award-number":["62022078"]}],"id":[{"id":"10.13039\/501100010909","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int J Comput Vis"],"published-print":{"date-parts":[[2024,12]]},"DOI":"10.1007\/s11263-024-02086-8","type":"journal-article","created":{"date-parts":[[2024,6,24]],"date-time":"2024-06-24T20:24:40Z","timestamp":1719260680000},"page":"5681-5697","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":4,"title":["HybridPrompt: Domain-Aware Prompting for Cross-Domain Few-Shot Learning"],"prefix":"10.1007","volume":"132","author":[{"given":"Jiamin","family":"Wu","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-0764-6106","authenticated-orcid":false,"given":"Tianzhu","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yongdong","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2024,6,24]]},"reference":[{"key":"2086_CR1","unstructured":"Antoniou, A., Edwards, H., & Storkey, A. (2018). How to train your MAML. In International conference on learning representations"},{"key":"2086_CR2","doi-asserted-by":"crossref","unstructured":"Bateni, P., Goyal, R., Masrani, V., Wood, F., & Sigal, L. (2020) Improved few-shot visual classification. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 14493\u201314502).","DOI":"10.1109\/CVPR42600.2020.01450"},{"key":"2086_CR3","unstructured":"Brown, T., Mann, B., Ryder, N., Subbiah, M., Kaplan, J. D., Dhariwal, P., Neelakantan, A., Shyam, P., Sastry, G., Askell, A., Agarwal, S., Herbert-Voss, A., Krueger, G., Henighan, T., Child, R., Ramesh, A., Ziegler, D., Wu, J., Winter, C., Hesse, C., Chen, M., Sigler, E., Litwin, M., Gray, S., Chess, B., Clark, J., Berner, C., McCandlish, S., Radford, A., Sutskever, I., & Amodei, D. (2020). Language models are few-shot learners. Advances in Neural Information Processing Systems, 33, 1877\u20131901."},{"key":"2086_CR4","doi-asserted-by":"crossref","unstructured":"Bulat, A., Guerrero, R., Martinez, B., & Tzimiropoulos, G. (2023) FS-DETR: Few-shot detection transformer with prompting and without re-training. In Proceedings of the IEEE\/CVF international conference on computer vision. (pp. 11793\u201311802).","DOI":"10.1109\/ICCV51070.2023.01083"},{"key":"2086_CR5","doi-asserted-by":"crossref","unstructured":"Carion, N., Massa, F., Synnaeve, G., Usunier, N., Kirillov, A., & Zagoruyko, S. (2020) End-to-end object detection with transformers. In European conference on computer vision (pp. 213\u2013229).","DOI":"10.1007\/978-3-030-58452-8_13"},{"key":"2086_CR6","first-page":"9912","volume":"33","author":"M Caron","year":"2020","unstructured":"Caron, M., Misra, I., Mairal, J., Goyal, P., Bojanowski, P., & Joulin, A. (2020). Unsupervised learning of visual features by contrasting cluster assignments. Advances in Neural Information Processing Systems, 33, 9912\u20139924.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"2086_CR7","doi-asserted-by":"crossref","unstructured":"Caron, M., Touvron, H., Misra, I., J\u00e9gou, H., Mairal, J., Bojanowski, P., & Joulin, A. (2021). Emerging properties in self-supervised vision transformers. In Proceedings of the IEEE international conference on computer vision (pp. 9650\u20139660).","DOI":"10.1109\/ICCV48922.2021.00951"},{"issue":"4","key":"2086_CR8","first-page":"4650","volume":"45","author":"G Cheng","year":"2022","unstructured":"Cheng, G., Lang, C., & Han, J. (2022). Holistic prototype activation for few-shot segmentation. IEEE Transactions on Pattern Analysis and Machine Intelligence, 45(4), 4650\u20134666.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"2086_CR9","doi-asserted-by":"crossref","unstructured":"Cimpoi, M., Maji, S., Kokkinos, I., Mohamed, S., & Vedaldi, A. (2014). Describing textures in the wild. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 3606\u20133613).","DOI":"10.1109\/CVPR.2014.461"},{"key":"2086_CR10","doi-asserted-by":"crossref","unstructured":"Cui, Y., Song, Y., Sun, C., Howard, A., & Belongie, S. (2018). Large scale fine-grained categorization and domain-specific transfer learning. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 4109\u20134118).","DOI":"10.1109\/CVPR.2018.00432"},{"key":"2086_CR11","first-page":"2292","volume":"26","author":"M Cuturi","year":"2013","unstructured":"Cuturi, M. (2013). Sinkhorn distances: Lightspeed computation of optimal transport. Advances in Neural Information Processing Systems, 26, 2292\u20132300.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"2086_CR12","doi-asserted-by":"crossref","unstructured":"Deng, J., Dong, W., Socher, R., Li, L. J., Li, K., & Fei-Fei, L. (2009). ImageNet: A large-scale hierarchical image database. In Proceedings of the IEEE international conference on computer vision (pp. 248\u2013255). IEEE.","DOI":"10.1109\/CVPR.2009.5206848"},{"key":"2086_CR13","unstructured":"Dosovitskiy, A., Beyer, L., Kolesnikov, A., Weissenborn, D., Zhai, X., Unterthiner, T., Dehghani, M., Minderer, M., Heigold, G., Gelly, S., Uszkoreit, J., & Houlsby, N. (2021). An image is worth 16 x 16 words: Transformers for image recognition at scale. In International conference on learning representations."},{"key":"2086_CR14","doi-asserted-by":"crossref","unstructured":"Dvornik, N., Schmid, C., & Mairal, J. (2020). Selecting relevant features from a multi-domain representation for few-shot classification. In European conference on computer vision (pp. 769\u2013786).","DOI":"10.1007\/978-3-030-58607-2_45"},{"issue":"4","key":"2086_CR15","doi-asserted-by":"publisher","first-page":"594","DOI":"10.1109\/TPAMI.2006.79","volume":"28","author":"L Fei-Fei","year":"2006","unstructured":"Fei-Fei, L., Fergus, R., & Perona, P. (2006). One-shot learning of object categories. IEEE Transactions on Pattern Analysis and Machine Intelligence, 28(4), 594\u2013611.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"2086_CR16","unstructured":"Finn, C., Abbeel, P., & Levine, S. (2017). Model-agnostic meta-learning for fast adaptation of deep networks. In International conference on machine learning (pp. 1126\u20131135)."},{"key":"2086_CR17","doi-asserted-by":"crossref","unstructured":"Guo, Y., Codella, N. C., Karlinsky, L., Codella, J. V., Smith, J. R., Saenko, K., Rosing, T., & Feris, R. (2020). A broader study of cross-domain few-shot learning. In European conference on computer vision (pp. 124\u2013141).","DOI":"10.1007\/978-3-030-58583-9_8"},{"key":"2086_CR18","unstructured":"Hou, R., Chang, H., Ma, B., Shan, S., & Chen, X. (2019). Cross attention network for few-shot classification. In Advances in neural information processing systems (pp. 4003\u20134014)."},{"key":"2086_CR19","doi-asserted-by":"crossref","unstructured":"Houben, S., Stallkamp, J., Salmen, J., Schlipsing, M., & Igel, C. (2013). Detection of traffic signs in real-world images: The German traffic sign detection benchmark. In International joint conference on neural networks (pp. 1\u20138). IEEE","DOI":"10.1109\/IJCNN.2013.6706807"},{"key":"2086_CR20","doi-asserted-by":"crossref","unstructured":"Hu, S. X., Li, D., St\u00fchmer, J., Kim, M., & Hospedales, T. M. (2022). Pushing the limits of simple pipelines for few-shot learning: External data and fine-tuning make a difference. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 9068\u20139077).","DOI":"10.1109\/CVPR52688.2022.00886"},{"key":"2086_CR21","doi-asserted-by":"crossref","unstructured":"Jia, M., Tang, L., Chen, B. C., Cardie, C., Belongie, S., Hariharan, B., & Lim, S. N. (2022). Visual prompt tuning. In European conference on computer vision (pp. 709\u2013727).","DOI":"10.1007\/978-3-031-19827-4_41"},{"key":"2086_CR22","unstructured":"Jongejan, J., Rowley, H., Kawashima, T., Kim, J., & Fox-Gieg, N. (2016). The quick, draw!-ai experiment. http:\/\/quickdraw.withgoogle.com"},{"key":"2086_CR23","unstructured":"Krizhevsky, A., & Hinton, G. (2009). Learning multiple layers of features from tiny images. Technical report, Citeseer."},{"key":"2086_CR24","doi-asserted-by":"crossref","unstructured":"Kumar Dwivedi, S., Gupta, V., Mitra, R., Ahmed, S., & Jain, A. (2019). Protogan: Towards few shot learning for action recognition. In Proceedings of the IEEE\/CVF international conference on computer vision workshops (pp. 0\u20130).","DOI":"10.1109\/ICCVW.2019.00166"},{"issue":"6266","key":"2086_CR25","doi-asserted-by":"publisher","first-page":"1332","DOI":"10.1126\/science.aab3050","volume":"350","author":"BM Lake","year":"2015","unstructured":"Lake, B. M., Salakhutdinov, R., & Tenenbaum, J. B. (2015). Human-level concept learning through probabilistic program induction. Science, 350(6266), 1332\u20131338.","journal-title":"Science"},{"key":"2086_CR26","doi-asserted-by":"crossref","unstructured":"Lang, C., Cheng, G., Tu, B., & Han, J. (2023a). Few-shot segmentation via divide-and-conquer proxies. International Journal of Computer Vision, 132, 1\u201323.","DOI":"10.1007\/s11263-023-01886-8"},{"key":"2086_CR27","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2023.3265865","author":"C Lang","year":"2023","unstructured":"Lang, C., Cheng, G., Tu, B., Li, C., & Han, J. (2023b). Base and meta: A new perspective on few-shot segmentation. IEEE Transactions on Pattern Analysis and Machine Intelligence. https:\/\/doi.org\/10.1109\/TPAMI.2023.3265865","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"2086_CR28","doi-asserted-by":"publisher","DOI":"10.1109\/TIP.2023.3315555","author":"C Lang","year":"2023","unstructured":"Lang, C., Cheng, G., Tu, B., Li, C., & Han, J. (2023c). Retain and recover: Delving into information loss for few-shot segmentation. IEEE Transactions on Image Processing. https:\/\/doi.org\/10.1109\/TIP.2023.3315555","journal-title":"IEEE Transactions on Image Processing"},{"key":"2086_CR29","unstructured":"LeCun, Y., & Cortes, C. (2010). MNIST handwritten digit database. http:\/\/yann.lecun.com\/exdb\/mnist"},{"key":"2086_CR30","doi-asserted-by":"crossref","unstructured":"Lee, K., Maji, S., Ravichandran, A., & Soatto, S. (2019). Meta-learning with differentiable convex optimization. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 10657\u201310665).","DOI":"10.1109\/CVPR.2019.01091"},{"key":"2086_CR31","doi-asserted-by":"crossref","unstructured":"Lester, B., Al-Rfou, R., & Constant, N. (2021). The power of scale for parameter-efficient prompt tuning. In Proceedings of the conference on empirical methods in natural language processing (pp. 3045\u20133059).","DOI":"10.18653\/v1\/2021.emnlp-main.243"},{"key":"2086_CR32","doi-asserted-by":"crossref","unstructured":"Li, W., Liu, X., & Bilen, H. (2022). Cross-domain few-shot learning with task-specific adapters. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 7161\u20137170).","DOI":"10.1109\/CVPR52688.2022.00702"},{"key":"2086_CR33","doi-asserted-by":"crossref","unstructured":"Li, W. H., Liu, X., & Bilen, H. (2021). Universal representation learning from multiple domains for few-shot classification. In Proceedings of the IEEE international conference on computer vision (pp. 9526\u20139535).","DOI":"10.1109\/ICCV48922.2021.00939"},{"key":"2086_CR34","doi-asserted-by":"crossref","unstructured":"Li, X. L., & Liang, P. (2021). Prefix-tuning: Optimizing continuous prompts for generation. In Proceedings of the 59th annual meeting of the association for computational linguistics and the 11th international joint conference on natural language processing (pp. 4582\u20134597).","DOI":"10.18653\/v1\/2021.acl-long.353"},{"key":"2086_CR35","doi-asserted-by":"crossref","unstructured":"Lin, T. Y., Maire, M., Belongie, S., Hays, J., Perona, P., Ramanan, D., Doll\u00e1r, P., & Zitnick, C. L. (2014). Microsoft coco: Common objects in context. In European conference on computer vision (pp. 740\u2013755). Springer.","DOI":"10.1007\/978-3-319-10602-1_48"},{"key":"2086_CR36","doi-asserted-by":"crossref","unstructured":"Liu, B., Cao, Y., Lin, Y., Li, Q., Zhang, Z., Long, M., & Hu, H. (2020). Negative margin matters: Understanding margin in few-shot classification. In European conference on computer vision (pp. 438\u2013455).","DOI":"10.1007\/978-3-030-58548-8_26"},{"key":"2086_CR37","doi-asserted-by":"crossref","unstructured":"Liu, L., Hamilton, W., Long, G., Jiang, J., & Larochelle, H. (2021a). A universal representation transformer layer for few-shot image classification. In International conference on learning representations.","DOI":"10.1109\/ICCV48922.2021.00939"},{"issue":"9","key":"2086_CR38","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1145\/3560815","volume":"55","author":"P Liu","year":"2023","unstructured":"Liu, P., Yuan, W., Fu, J., Jiang, Z., Hayashi, H., & Neubig, G. (2023). Pre-train, prompt, and predict: A systematic survey of prompting methods in natural language processing. ACM Computing Surveys, 55(9), 1\u201335.","journal-title":"ACM Computing Surveys"},{"key":"2086_CR39","doi-asserted-by":"crossref","unstructured":"Liu, Y., Lee, J., Zhu, L., Chen, L., Shi, H., & Yang, Y. (2021b). A multi-mode modulator for multi-domain few-shot classification. In Proceedings of the IEEE international conference on computer vision (pp. 8453\u20138462).","DOI":"10.1109\/ICCV48922.2021.00834"},{"key":"2086_CR40","doi-asserted-by":"crossref","unstructured":"Ma, T., Sun, Y., Yang, Z., & Yang, Y. (2023). Prod: Prompting-to-disentangle domain knowledge for cross-domain few-shot image classification. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 19754\u201319763).","DOI":"10.1109\/CVPR52729.2023.01892"},{"key":"2086_CR41","unstructured":"Maji, S., Rahtu, E., Kannala, J., Blaschko, M., & Vedaldi, A. (2013). Fine-grained visual classification of aircraft. arXiv preprint arXiv:1306.5151"},{"key":"2086_CR42","doi-asserted-by":"crossref","unstructured":"Nilsback, M. E., & Zisserman, A. (2008). Automated flower classification over a large number of classes. In Indian Conference on Computer Vision (pp. 722\u2013729). IEEE: Graphics & Image Processing.","DOI":"10.1109\/ICVGIP.2008.47"},{"key":"2086_CR43","unstructured":"Oreshkin, B., Rodr\u00edguez\u00a0L\u00f3pez, P., & Lacoste, A. (2018). TADAM: Task dependent adaptive metric for improved few-shot learning. In Advances in neural information processing systems (pp. 721\u2013731)."},{"key":"2086_CR44","doi-asserted-by":"crossref","unstructured":"Perrett, T., Masullo, A., Burghardt, T., Mirmehdi, M., & Damen, D. (2021). Temporal-relational crosstransformers for few-shot action recognition. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 475\u2013484).","DOI":"10.1109\/CVPR46437.2021.00054"},{"key":"2086_CR45","unstructured":"Radford, A., Kim, J. W., Hallacy, C., Ramesh, A., Goh, G., Agarwal, S., Sastry, G., Askell, A., Mishkin, P., Clark, J., Krueger, G., & Sutskever, I. (2021). Learning transferable visual models from natural language supervision. In International conference on machine learning (pp. 8748\u20138763). PMLR."},{"key":"2086_CR46","unstructured":"Raghu, M., Unterthiner, T., Kornblith, S., Zhang, C., & Dosovitskiy, A. (2021). Do vision transformers see like convolutional neural networks?. In Advances in neural information processing systems (pp. 12116\u201312128)."},{"key":"2086_CR47","unstructured":"Ravi, S., & Larochelle, H. (2017). Optimization as a model for few-shot learning. In International conference on learning representations."},{"key":"2086_CR48","unstructured":"Requeima, J., Gordon, J., Bronskill, J., Nowozin, S., & Turner, R. E. (2019). Fast and flexible multi-task classification using conditional neural adaptive processes. In Advances in neural information processing systems (pp. 7959\u20137970)."},{"key":"2086_CR49","doi-asserted-by":"crossref","unstructured":"Rubner, Y., Tomasi, C., & Guibas, L. J. (1998). A metric for distributions with applications to image databases. In Sixth international conference on computer vision (IEEE Cat. No. 98CH36271) (pp. 59\u201366). IEEE.","DOI":"10.1109\/ICCV.1998.710701"},{"key":"2086_CR50","unstructured":"Schroeder, B., & Cui, Y. (2018). FGVCx fungi classification challenge 2018. https:\/\/github.com\/visipedia\/fgvcx_fungi_comp"},{"key":"2086_CR51","doi-asserted-by":"crossref","unstructured":"Shin, T., Razeghi, Y., Logan IV, R. L., Wallace, E., & Singh, S. (2020). Autoprompt: Eliciting knowledge from language models with automatically generated prompts. In EMNLP (pp. 4222\u20134235).","DOI":"10.18653\/v1\/2020.emnlp-main.346"},{"key":"2086_CR52","doi-asserted-by":"crossref","unstructured":"Simon, C., Koniusz, P., Nock, R., & Harandi, M. (2020). Adaptive subspaces for few-shot learning. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 4136\u20134145).","DOI":"10.1109\/CVPR42600.2020.00419"},{"key":"2086_CR53","unstructured":"Snell, J., Swersky, K., & Zemel, R. (2017). Prototypical networks for few-shot learning. In Advances in neural information processing systems (pp. 4077\u20134087)."},{"key":"2086_CR54","doi-asserted-by":"crossref","unstructured":"Sun, B., Li, B., Cai, S., Yuan, Y., & Zhang, C. (2021). FSCE: Few-shot object detection via contrastive proposal encoding. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 7352\u20137362).","DOI":"10.1109\/CVPR46437.2021.00727"},{"key":"2086_CR55","doi-asserted-by":"crossref","unstructured":"Sun, Q., Liu, Y., Chua, T. S., & Schiele, B. (2019). Meta-transfer learning for few-shot learning. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 403\u2013412).","DOI":"10.1109\/CVPR.2019.00049"},{"key":"2086_CR56","doi-asserted-by":"crossref","unstructured":"Sung, F., Yang, Y., Zhang, L., Xiang, T., Torr, P. H., & Hospedales, T. M. (2018). Learning to compare: Relation network for few-shot learning. In Proceedings of the IEEE conference on computer vision and pattern recognition.","DOI":"10.1109\/CVPR.2018.00131"},{"key":"2086_CR57","doi-asserted-by":"crossref","unstructured":"Tian, Y., Wang, Y., Krishnan, D., Tenenbaum, J. B., & Isola, P. (2020). Rethinking few-shot image classification: A good embedding is all you need? In European conference on computer vision (pp. 266\u2013282).","DOI":"10.1007\/978-3-030-58568-6_16"},{"key":"2086_CR58","unstructured":"Triantafillou, E., Zhu, T., Dumoulin, V., Lamblin, P., Evci, U., Xu, K., Goroshin, R., Gelada, C., Swersky, K., Manzagol, P. A., & Larochelle, H. (2019). Meta-dataset: A dataset of datasets for learning to learn from few examples. arXiv preprint arXiv:1903.03096"},{"key":"2086_CR59","unstructured":"Triantafillou, E., Larochelle, H., Zemel, R., & Dumoulin, V. (2021). Learning a universal template for few-shot dataset generalization. In International conference on machine learning (pp. 10424\u201310433)."},{"key":"2086_CR60","unstructured":"Vinyals, O., Blundell, C., Lillicrap, T., & Wierstra, D. (2016). Matching networks for one shot learning. In Advances in neural information processing systems (pp. 3630\u20133638)."},{"key":"2086_CR61","unstructured":"Wah, C., Branson, S., Welinder, P., Perona, P., & Belongie, S. (2011). The Caltech-UCSD birds-200-2011 dataset. Technical report."},{"key":"2086_CR62","doi-asserted-by":"crossref","unstructured":"Wang, Z., Zhang, Z., Ebrahimi, S., Sun, R., Zhang, H., Lee, C. Y., Ren, X., Su, G., Perot, V., Dy, J. & Pfister, T. (2022). Dualprompt: Complementary prompting for rehearsal-free continual learning. In European conference on computer vision (pp. 631\u2013648).","DOI":"10.1007\/978-3-031-19809-0_36"},{"key":"2086_CR63","doi-asserted-by":"crossref","unstructured":"Wu, J., Zhang, T., Zhang, Y., & Wu, F. (2021). Task-aware part mining network for few-shot learning. In Proceedings of the IEEE international conference on computer vision (pp. 8433\u20138442).","DOI":"10.1109\/ICCV48922.2021.00832"},{"key":"2086_CR64","doi-asserted-by":"crossref","unstructured":"Wu, J., Zhang, T., Zhang, Z., Wu, F., & Zhang, Y. (2022). Motion-modulated temporal fragment alignment network for few-shot action recognition. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 9151\u20139160).","DOI":"10.1109\/CVPR52688.2022.00894"},{"key":"2086_CR65","first-page":"12077","volume":"34","author":"E Xie","year":"2021","unstructured":"Xie, E., Wang, W., Yu, Z., Anandkumar, A., Alvarez, J. M., & Luo, P. (2021). SegFormer: Simple and efficient design for semantic segmentation with transformers. Advances in Neural Information Processing Systems, 34, 12077\u201312090.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"2086_CR66","doi-asserted-by":"crossref","unstructured":"Ye, H. J., Hu, H., Zhan, D. C., & Sha, F. (2020). Few-shot learning via embedding adaptation with set-to-set functions. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 8808\u20138817).","DOI":"10.1109\/CVPR42600.2020.00883"},{"key":"2086_CR67","unstructured":"Zeiler, M. D. (2012). Adadelta: An adaptive learning rate method. arXiv preprint arXiv:1212.5701"},{"key":"2086_CR68","doi-asserted-by":"crossref","unstructured":"Zeiler, M. D., & Fergus, R. (2014). Visualizing and understanding convolutional networks. In European conference on computer vision (pp. 818\u2013833). Springer.","DOI":"10.1007\/978-3-319-10590-1_53"},{"key":"2086_CR69","doi-asserted-by":"crossref","unstructured":"Zhang, C., Cai, Y., Lin, G., & Shen, C. (2020). DeepEMD: Few-shot image classification with differentiable earth mover\u2019s distance and structured classifiers. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 12203\u201312213).","DOI":"10.1109\/CVPR42600.2020.01222"},{"key":"2086_CR70","doi-asserted-by":"crossref","unstructured":"Zhang, R., Hu, X., Li, B., Huang, S., Deng, H., Qiao, Y., Gao, P., & Li, H. (2023) Prompt, generate, then cache: Cascade of foundation models makes strong few-shot learners. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 15211\u201315222).","DOI":"10.1109\/CVPR52729.2023.01460"},{"issue":"9","key":"2086_CR71","doi-asserted-by":"publisher","first-page":"2337","DOI":"10.1007\/s11263-022-01653-1","volume":"130","author":"K Zhou","year":"2022","unstructured":"Zhou, K., Yang, J., Loy, C. C., & Liu, Z. (2022a). Learning to prompt for vision-language models. International Journal of Computer Vision, 130(9), 2337\u20132348.","journal-title":"International Journal of Computer Vision"},{"key":"2086_CR72","doi-asserted-by":"crossref","unstructured":"Zhou, K., Yang, J., Loy, C. C., & Liu, Z. (2022b) Conditional prompt learning for vision-language models. In Proceedings of the IEEE conference on computer vision and pattern recognition (pp. 16816\u201316825).","DOI":"10.1109\/CVPR52688.2022.01631"},{"key":"2086_CR73","doi-asserted-by":"crossref","unstructured":"Zhu, C., Chen, F., Ahmed, U., Shen, Z., & Savvides, M. (2021). Semantic relation reasoning for shot-stable few-shot object detection. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (pp. 8782\u20138791).","DOI":"10.1109\/CVPR46437.2021.00867"}],"container-title":["International Journal of Computer Vision"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-024-02086-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11263-024-02086-8\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-024-02086-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,11,15]],"date-time":"2024-11-15T10:14:42Z","timestamp":1731665682000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11263-024-02086-8"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,6,24]]},"references-count":73,"journal-issue":{"issue":"12","published-print":{"date-parts":[[2024,12]]}},"alternative-id":["2086"],"URL":"https:\/\/doi.org\/10.1007\/s11263-024-02086-8","relation":{},"ISSN":["0920-5691","1573-1405"],"issn-type":[{"value":"0920-5691","type":"print"},{"value":"1573-1405","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,6,24]]},"assertion":[{"value":"5 December 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"17 April 2024","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"24 June 2024","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors have no relevant financial or non-financial interests to disclose.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}]}}