{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,24]],"date-time":"2026-06-24T16:18:15Z","timestamp":1782317895748,"version":"3.54.5"},"reference-count":50,"publisher":"Springer Science and Business Media LLC","issue":"2","license":[{"start":{"date-parts":[[2024,4,1]],"date-time":"2024-04-01T00:00:00Z","timestamp":1711929600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,4,1]],"date-time":"2024-04-01T00:00:00Z","timestamp":1711929600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"The National Natural Science Foundation of China","doi-asserted-by":"crossref","award":["62172155"],"award-info":[{"award-number":["62172155"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"crossref"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimedia Systems"],"published-print":{"date-parts":[[2024,4]]},"DOI":"10.1007\/s00530-024-01316-w","type":"journal-article","created":{"date-parts":[[2024,4,3]],"date-time":"2024-04-03T16:02:02Z","timestamp":1712160122000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":8,"title":["Intelligent-paint: a Chinese painting process generation method based on vision transformer"],"prefix":"10.1007","volume":"30","author":[{"given":"Zunfu","family":"Wang","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Fang","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Zhixiong","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Changjuan","family":"Ran","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Mohan","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2024,4,3]]},"reference":[{"key":"1316_CR1","doi-asserted-by":"crossref","unstructured":"Hertzmann, A.: A Survey of Stroke-Based Rendering. Institute of Electrical and Electronics Engineers (2003)","DOI":"10.1109\/MCG.2003.1210867"},{"issue":"6","key":"1316_CR2","doi-asserted-by":"publisher","first-page":"2806","DOI":"10.1109\/TPAMI.2020.3045007","volume":"44","author":"S Oprea","year":"2020","unstructured":"Oprea, S., Martinez-Gonzalez, P., Garcia-Garcia, A., Castro-Vargas, J.A., Orts-Escolano, S., Garcia-Rodriguez, J., Argyros, A.: A review on deep learning techniques for video prediction. IEEE Trans. Pattern Anal. Mach. Intell. 44(6), 2806\u20132826 (2020)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"1316_CR3","doi-asserted-by":"crossref","unstructured":"Zhao, A., Balakrishnan, G., Lewis, K.M., Durand, F., Guttag, J.V., Dalca, A.V.: Painting many pasts: Synthesizing time lapse videos of paintings. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 8435\u20138445 (2020)","DOI":"10.1109\/CVPR42600.2020.00846"},{"key":"1316_CR4","doi-asserted-by":"crossref","unstructured":"Yan, X., Yang, J., Sohn, K., Lee, H.: Attribute2image: Conditional image generation from visual attributes. In: Computer Vision\u2014ECCV 2016: 14th European Conference, Amsterdam, The Netherlands, October 11\u201314, 2016, Proceedings, Part IV 14, pp. 776\u2013791. Springer (2016)","DOI":"10.1007\/978-3-319-46493-0_47"},{"key":"1316_CR5","first-page":"2672","volume":"7","author":"I Goodfellow","year":"2014","unstructured":"Goodfellow, I., Pouget-Abadie, J., Mirza, M., Xu, B., Warde-Farley, D., Ozair, S., Courville, A., Bengio, Y.: Generative adversarial nets. Adv. Neural Inf. Process. Syst. 7, 2672\u20132680 (2014)","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"1316_CR6","doi-asserted-by":"crossref","unstructured":"Li, C., Wand, M.: Precomputed real-time texture synthesis with Markovian generative adversarial networks. In: Computer Vision\u2014ECCV 2016: 14th European Conference, Amsterdam, The Netherlands, October 11\u201314, 2016, Proceedings, Part III 14, pp. 702\u2013716. Springer (2016)","DOI":"10.1007\/978-3-319-46487-9_43"},{"key":"1316_CR7","doi-asserted-by":"crossref","unstructured":"Zhu, J.-Y., Park, T., Isola, P., Efros, A.A.: Unpaired image-to-image translation using cycle-consistent adversarial networks. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 2223\u20132232 (2017)","DOI":"10.1109\/ICCV.2017.244"},{"issue":"3","key":"1316_CR8","doi-asserted-by":"publisher","first-page":"853","DOI":"10.1145\/1073204.1073273","volume":"24","author":"Y Chuang","year":"2005","unstructured":"Chuang, Y., Goldman, D., Zheng, K., Curless, B., Salesin, D., Szeliski, R.: Animating pictures with stochastic motion textures. ACM Trans. Graph. 24(3), 853\u2013860 (2005)","journal-title":"ACM Trans. Graph."},{"key":"1316_CR9","doi-asserted-by":"crossref","unstructured":"Torbunov, D., Huang, Y., Yu, H., Huang, J., Yoo, S., Lin, M., Viren, B., Ren, Y.: UVCGAN: UNet vision transformer cycle-consistent gan for unpaired image-to-image translation. In: Proceedings of the IEEE\/CVF Winter Conference on Applications of Computer Vision, pp. 702\u2013712 (2023)","DOI":"10.1109\/WACV56688.2023.00077"},{"key":"1316_CR10","doi-asserted-by":"crossref","unstructured":"He, B., Gao, F., Ma, D., Shi, B., Duan, L.-Y.: ChipGAN: a generative adversarial network for chinese ink wash painting style transfer. In: Proceedings of the 26th ACM International Conference on Multimedia, pp. 1172\u20131180 (2018)","DOI":"10.1145\/3240508.3240655"},{"key":"1316_CR11","doi-asserted-by":"crossref","unstructured":"Xie, S., Tu, Z.: Holistically-nested edge detection. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 1395\u20131403 (2015)","DOI":"10.1109\/ICCV.2015.164"},{"key":"1316_CR12","doi-asserted-by":"crossref","unstructured":"Haeberli, P.: Paint by numbers: abstract image representations. In: Proceedings of the 17th Annual Conference on Computer Graphics and Interactive Techniques, pp. 207\u2013214 (1990)","DOI":"10.1145\/97879.97902"},{"key":"1316_CR13","doi-asserted-by":"crossref","unstructured":"Lewis, J.-P.: Texture synthesis for digital painting. In: Proceedings of the 11th Annual Conference on Computer Graphics and Interactive Techniques, pp. 245\u2013252 (1984)","DOI":"10.1145\/800031.808605"},{"issue":"3","key":"1316_CR14","doi-asserted-by":"publisher","first-page":"313","DOI":"10.1145\/325165.325250","volume":"19","author":"WT Reeves","year":"1985","unstructured":"Reeves, W.T., Blau, R.: Approximate and probabilistic algorithms for shading and rendering structured particle systems. ACM SIGGRAPH Comput. Graph. 19(3), 313\u2013322 (1985)","journal-title":"ACM SIGGRAPH Comput. Graph."},{"key":"1316_CR15","doi-asserted-by":"crossref","unstructured":"Hertzmann, A.: Painterly rendering with curved brush strokes of multiple sizes. In: Proceedings of the 25th Annual Conference on Computer Graphics and Interactive Techniques, pp. 453\u2013460 (1998)","DOI":"10.1145\/280814.280951"},{"key":"1316_CR16","doi-asserted-by":"crossref","unstructured":"Fu, H., Zhou, S., Liu, L., Mitra, N.J.: Animated construction of line drawings. In: Proceedings of the 2011 SIGGRAPH Asia Conference, pp. 1\u201310 (2011)","DOI":"10.1145\/2024156.2024167"},{"issue":"12","key":"1316_CR17","doi-asserted-by":"publisher","first-page":"3019","DOI":"10.1109\/TVCG.2017.2774292","volume":"24","author":"F Tang","year":"2017","unstructured":"Tang, F., Dong, W., Meng, Y., Mei, X., Huang, F., Zhang, X., Deussen, O.: Animated construction of Chinese brush paintings. IEEE Trans. Vis. Comput. Graph. 24(12), 3019\u20133031 (2017)","journal-title":"IEEE Trans. Vis. Comput. Graph."},{"key":"1316_CR18","unstructured":"Frans, K., Cheng, C.-Y.: Unsupervised image to sequence translation with canvas-drawer networks. arXiv preprint arXiv:1809.08340 (2018)"},{"key":"1316_CR19","doi-asserted-by":"crossref","unstructured":"Huang, Z., Heng, W., Zhou, S.: Learning to paint with model-based deep reinforcement learning. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 8709\u20138718 (2019)","DOI":"10.1109\/ICCV.2019.00880"},{"key":"1316_CR20","doi-asserted-by":"crossref","unstructured":"Singh, J., Zheng, L.: Combining semantic guidance and deep reinforcement learning for generating human level paintings. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 16387\u201316396 (2021)","DOI":"10.1109\/CVPR46437.2021.01612"},{"key":"1316_CR21","doi-asserted-by":"crossref","unstructured":"Liu, S., Lin, T., He, D., Li, F., Deng, R., Li, X., Ding, E., Wang, H.: Paint transformer: feed forward neural painting with stroke prediction. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 6598\u20136607 (2021)","DOI":"10.1109\/ICCV48922.2021.00653"},{"key":"1316_CR22","doi-asserted-by":"crossref","unstructured":"Zou, Z., Shi, T., Qiu, S., Yuan, Y., Shi, Z.: Stylized neural painting. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 15689\u201315698 (2021)","DOI":"10.1109\/CVPR46437.2021.01543"},{"key":"1316_CR23","doi-asserted-by":"crossref","unstructured":"Litwinowicz, P.: Processing images and video for an impressionist effect. In: Proceedings of the 24th Annual Conference on Computer Graphics and Interactive Techniques, pp. 407\u2013414 (1997)","DOI":"10.1145\/258734.258893"},{"key":"1316_CR24","unstructured":"Ha, D., Eck, D.: A neural representation of sketch drawings. arXiv preprint arXiv:1704.03477 (2017)"},{"key":"1316_CR25","unstructured":"Ganin, Y., Kulkarni, T., Babuschkin, I., Eslami, S.A., Vinyals, O.: Synthesizing programs for images using reinforced adversarial learning. In: International Conference on Machine Learning, pp. 1666\u20131675. PMLR (2018)"},{"issue":"5","key":"1316_CR26","doi-asserted-by":"publisher","first-page":"1134","DOI":"10.1587\/transinf.E96.D.1134","volume":"96","author":"N Xie","year":"2013","unstructured":"Xie, N., Hachiya, H., Sugiyama, M.: Artist Agent: a reinforcement learning approach to automatic stroke generation in oriental ink painting. IEICE Trans. Inf. Syst. 96(5), 1134\u20131144 (2013)","journal-title":"IEICE Trans. Inf. Syst."},{"key":"1316_CR27","unstructured":"Zhou, T., Fang, C., Wang, Z., Yang, J., Kim, B., Chen, Z., Brandt, J., Terzopoulos, D.: Learning to sketch with deep q networks and demonstrated strokes. arXiv preprint arXiv:1810.05977 (2018)"},{"key":"1316_CR28","unstructured":"Zheng, N., Jiang, Y., Huang, D.: StrokeNet: a neural painting environment. In: International Conference on Learning Representations (2018)"},{"key":"1316_CR29","doi-asserted-by":"crossref","unstructured":"Zhao, A., Balakrishnan, G., Lewis, K.M., Durand, F., Guttag, J.V., Dalca, A.V.: Painting many pasts: Synthesizing time lapse videos of paintings. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 8435\u20138445 (2020)","DOI":"10.1109\/CVPR42600.2020.00846"},{"key":"1316_CR30","unstructured":"Van Den\u00a0Oord, A., Kalchbrenner, N., Kavukcuoglu, K.: Pixel recurrent neural networks. In: International Conference on Machine Learning, pp. 1747\u20131756. PMLR (2016)"},{"key":"1316_CR31","unstructured":"Oord, A., Kalchbrenner, N., Espeholt, L., Vinyals, O., Graves, A., et al.: Conditional image generation with pixelcnn decoders. Adv. Neural Inf. Process. Syst. 29 (2016)"},{"key":"1316_CR32","unstructured":"Dosovitskiy, A., Beyer, L., Kolesnikov, A., Weissenborn, D., Zhai, X., Unterthiner, T., Dehghani, M., Minderer, M., Heigold, G., Gelly, S., et al.: An image is worth $$16 \\times 16$$ words: transformers for image recognition at scale. arXiv preprint arXiv:2010.11929 (2020)"},{"key":"1316_CR33","first-page":"30392","volume":"34","author":"T Xiao","year":"2021","unstructured":"Xiao, T., Singh, M., Mintun, E., Darrell, T., Doll\u00e1r, P., Girshick, R.: Early convolutions help transformers see better. Adv. Neural Inf. Process. Syst. 34, 30392\u201330400 (2021)","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"1316_CR34","doi-asserted-by":"crossref","unstructured":"Guo, J., Han, K., Wu, H., Tang, Y., Chen, X., Wang, Y., Xu, C.: CMT: convolutional neural networks meet vision transformers. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 12175\u201312185 (2022)","DOI":"10.1109\/CVPR52688.2022.01186"},{"key":"1316_CR35","first-page":"8780","volume":"34","author":"P Dhariwal","year":"2021","unstructured":"Dhariwal, P., Nichol, A.: Diffusion models beat gans on image synthesis. Adv. Neural Inf. Process. Syst. 34, 8780\u20138794 (2021)","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"1316_CR36","unstructured":"Ming, R.: Chinese brush painting: an academic approach for painting flowers and fish [paperback]"},{"key":"1316_CR37","unstructured":"Dwight, J.: The Chinese brush painting bible: over 200 motifs with step-by-step illustrated instructions (2011)"},{"key":"1316_CR38","doi-asserted-by":"crossref","unstructured":"Ronneberger, O., Fischer, P., Brox, T.: U-net: Convolutional networks for biomedical image segmentation. In: Medical Image Computing and Computer-Assisted Intervention\u2014MICCAI 2015: 18th International Conference, Munich, Germany, October 5\u20139, 2015, Proceedings, Part III 18, pp. 234\u2013241. Springer (2015)","DOI":"10.1007\/978-3-319-24574-4_28"},{"key":"1316_CR39","unstructured":"Devlin, J., Chang, M.-W., Lee, K., Toutanova, K.: Bert: pre-training of deep bidirectional transformers for language understanding. arXiv preprint arXiv:1810.04805 (2018)"},{"key":"1316_CR40","doi-asserted-by":"crossref","unstructured":"Anokhin, I., Demochkin, K., Khakhulin, T., Sterkin, G., Lempitsky, V., Korzhenkov, D.: Image generators with conditionally-independent pixel synthesis. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 14278\u201314287 (2021)","DOI":"10.1109\/CVPR46437.2021.01405"},{"key":"1316_CR41","unstructured":"Bachlechner, T., Majumder, B.P., Mao, H., Cottrell, G., McAuley, J.: Rezero is all you need: Fast convergence at large depth. In: Uncertainty in Artificial Intelligence, pp. 1352\u20131361. PMLR (2021)"},{"key":"1316_CR42","doi-asserted-by":"crossref","unstructured":"Wang, Y., Wu, C., Herranz, L., Weijer, J., Gonzalez-Garcia, A., Raducanu, B.: Transferring gans: generating images from limited data. In: Proceedings of the European Conference on Computer Vision (ECCV), pp. 218\u2013234 (2018)","DOI":"10.1007\/978-3-030-01231-1_14"},{"key":"1316_CR43","doi-asserted-by":"crossref","unstructured":"Noguchi, A., Harada, T.: Image generation from small datasets via batch statistics adaptation. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 2750\u20132758 (2019)","DOI":"10.1109\/ICCV.2019.00284"},{"key":"1316_CR44","unstructured":"Zhao, M., Cong, Y., Carin, L.: On leveraging pretrained gans for limited-data generation. In: Proc. ICML, pp. 11340\u201311351 (2020)"},{"key":"1316_CR45","doi-asserted-by":"crossref","unstructured":"Wang, Y., Gonzalez-Garcia, A., Berga, D., Herranz, L., Khan, F.S., Weijer, J.v.d.: Minegan: effective knowledge transfer from gans to target domains with few images. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 9332\u20139341 (2020)","DOI":"10.1109\/CVPR42600.2020.00935"},{"key":"1316_CR46","unstructured":"Grigoryev, T., Voynov, A., Babenko, A.: When, why, and which pretrained gans are useful? arXiv preprint arXiv:2202.08937 (2022)"},{"key":"1316_CR47","unstructured":"Heusel, M., Ramsauer, H., Unterthiner, T., Nessler, B., Hochreiter, S.: Gans trained by a two time-scale update rule converge to a local Nash equilibrium. Adv. Neural Inf. Process. Syst. 30 (2017)"},{"issue":"4","key":"1316_CR48","doi-asserted-by":"publisher","first-page":"600","DOI":"10.1109\/TIP.2003.819861","volume":"13","author":"Z Wang","year":"2004","unstructured":"Wang, Z., Bovik, A.C., Sheikh, H.R., Simoncelli, E.P.: Image quality assessment: from error visibility to structural similarity. IEEE Trans. Image Process. 13(4), 600\u2013612 (2004)","journal-title":"IEEE Trans. Image Process."},{"key":"1316_CR49","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1007\/s11432-019-2757-1","volume":"63","author":"G Zhai","year":"2020","unstructured":"Zhai, G., Min, X.: Perceptual image quality assessment: a survey. Sci. China Inf. Sci. 63, 1\u201352 (2020)","journal-title":"Sci. China Inf. Sci."},{"key":"1316_CR50","unstructured":"Ranzato, M., Szlam, A., Bruna, J., Mathieu, M., Collobert, R., Chopra, S.: Video (language) modeling: a baseline for generative models of natural videos. arXiv preprint arXiv:1412.6604 (2014)"}],"container-title":["Multimedia Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-024-01316-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00530-024-01316-w\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-024-01316-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,4,12]],"date-time":"2024-04-12T13:18:30Z","timestamp":1712927910000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00530-024-01316-w"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,4]]},"references-count":50,"journal-issue":{"issue":"2","published-print":{"date-parts":[[2024,4]]}},"alternative-id":["1316"],"URL":"https:\/\/doi.org\/10.1007\/s00530-024-01316-w","relation":{},"ISSN":["0942-4962","1432-1882"],"issn-type":[{"value":"0942-4962","type":"print"},{"value":"1432-1882","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,4]]},"assertion":[{"value":"28 October 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"6 March 2024","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"3 April 2024","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"112"}}