{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,23]],"date-time":"2026-07-23T09:00:47Z","timestamp":1784797247330,"version":"3.55.0"},"reference-count":76,"publisher":"Springer Science and Business Media LLC","issue":"4","license":[{"start":{"date-parts":[[2026,6,11]],"date-time":"2026-06-11T00:00:00Z","timestamp":1781136000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,6,11]],"date-time":"2026-06-11T00:00:00Z","timestamp":1781136000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Mach. Intell. Res."],"published-print":{"date-parts":[[2026,8]]},"DOI":"10.1007\/s11633-025-1606-9","type":"journal-article","created":{"date-parts":[[2026,6,11]],"date-time":"2026-06-11T03:05:59Z","timestamp":1781147159000},"page":"804-822","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["T2IW: Joint Text to Image &amp; Watermark Generation"],"prefix":"10.1007","volume":"23","author":[{"ORCID":"https:\/\/orcid.org\/0009-0003-7012-6485","authenticated-orcid":false,"given":"Guokai","family":"Zhang","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-5165-204X","authenticated-orcid":false,"given":"Yuting","family":"Su","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-7696-5330","authenticated-orcid":false,"given":"Lanjun","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0008-2004-0508","authenticated-orcid":false,"given":"Guochang","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-5467-0910","authenticated-orcid":false,"given":"Dan","family":"Song","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-5755-9145","authenticated-orcid":false,"given":"An-An","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,6,11]]},"reference":[{"issue":"8","key":"1606_CR1","doi-asserted-by":"publisher","first-page":"6809","DOI":"10.1109\/TCSVT.2024.3427488","volume":"34","author":"S Li","year":"2024","unstructured":"S. Li, X. Li, L. Chiariglione, J. Luo, W. Wang, Z. Yang, D. Mandic, H. Fujita. Introduction to the special issue on AI-generated content for multimedia. IEEE Transactions on Circuits and Systems for Video Technology, vol. 34, no. 8, pp. 6809\u20136813, 2024. DOI: https:\/\/doi.org\/10.1109\/TCSVT.2024.3427488.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"issue":"8","key":"1606_CR2","doi-asserted-by":"publisher","first-page":"6814","DOI":"10.1109\/TCSVT.2024.3351601","volume":"34","author":"F Nazarieh","year":"2024","unstructured":"F. Nazarieh, Z. Feng, M. Awais, W. Wang, J. Kittler. A survey of cross-modal visual content generation. IEEE Transactions on Circuits and Systems for Video Technology, vol. 34, no. 8, pp. 6814\u20136832, 2024. DOI: https:\/\/doi.org\/10.1109\/tcsvt.2024.3351601.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"key":"1606_CR3","doi-asserted-by":"publisher","first-page":"16494","DOI":"10.1109\/CVPR52688.2022.01602","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"M Tao","year":"2022","unstructured":"M. Tao, H. Tang, F. Wu, X. Jing, B. K. Bao, C. Xu. DF-GAN: A simple and effective baseline for text-to-image synthesis. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, New Orleans, USA, pp. 16494\u201316504, 2022. DOI: https:\/\/doi.org\/10.1109\/CVPR52688.2022.01602."},{"key":"1606_CR4","doi-asserted-by":"publisher","first-page":"462","DOI":"10.1109\/TMM.2023.3266607","volume":"26","author":"S Ye","year":"2024","unstructured":"S. Ye, H. Wang, M. Tan, F. Liu. Recurrent affine transformation for text-to-image synthesis. IEEE Transactions on Multimedia, vol. 26, pp. 462\u2013473, 2024. DOI: https:\/\/doi.org\/10.1109\/TMM.2023.3266607.","journal-title":"IEEE Transactions on Multimedia"},{"key":"1606_CR5","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1109\/IJCNN52387.2021.9533527","volume-title":"Proceedings of International Joint Conference on Neural Networks","author":"Z Zhang","year":"2021","unstructured":"Z. Zhang, L. Schomaker. DTGAN: Dual attention generative adversarial networks for text-to-image generation. In Proceedings of International Joint Conference on Neural Networks, IEEE, Shenzhen, China, pp. 1\u20138, 2021. DOI: https:\/\/doi.org\/10.1109\/IJCNN52387.2021.9533527."},{"key":"1606_CR6","doi-asserted-by":"publisher","first-page":"5908","DOI":"10.1109\/ICCV.2017.629","volume-title":"Proceedings of IEEE International Conference on Computer Vision","author":"H Zhang","year":"2017","unstructured":"H. Zhang, T. Xu, H. Li, S. Zhang, X. Wang, X. Huang, D. Metaxas. StackGAN: Text to photo-realistic image synthesis with stacked generative adversarial networks. In Proceedings of IEEE International Conference on Computer Vision, Venice, Italy, pp. 5908\u20135916, 2017. DOI: https:\/\/doi.org\/10.1109\/ICCV.2017.629."},{"issue":"8","key":"1606_CR7","doi-asserted-by":"publisher","first-page":"6901","DOI":"10.1109\/TCSVT.2024.3349567","volume":"34","author":"C Jin","year":"2024","unstructured":"C. Jin, R. Zhu, Z. Zhu, L. Yang, M. Yang, J. Luo. MtArtGPT: A multi-task art generation system with pre-trained transformer. IEEE Transactions on Circuits and Systems for Video Technology, vol. 34, no. 8, pp. 6901\u20136912, 2024. DOI: https:\/\/doi.org\/10.1109\/tcsvt.2024.3349567.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"issue":"8","key":"1606_CR8","doi-asserted-by":"publisher","first-page":"6913","DOI":"10.1109\/TCSVT.2023.3298811","volume":"34","author":"Y Zhao","year":"2024","unstructured":"Y. Zhao, H. Li, Z. Zhang, Y. Chen, Q. Liu, X. Zhang. Regional traditional painting generation based on controllable disentanglement model. IEEE Transactions on Circuits and Systems for Video Technology, vol. 34, no. 8, pp. 6913\u20136925, 2024. DOI: https:\/\/doi.org\/10.1109\/tcsvt.2023.3298811.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"key":"1606_CR9","volume-title":"Proceedings of the 34th International Conference on Neural Information Processing Systems","author":"J Ho","year":"2020","unstructured":"J. Ho, A. Jain, P. Abbeel. Denoising diffusion probabilistic models. In Proceedings of the 34th International Conference on Neural Information Processing Systems, Vancouver, Canada, Article number 574, 2020."},{"key":"1606_CR10","doi-asserted-by":"publisher","first-page":"10674","DOI":"10.1109\/CVPR52688.2022.01042","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"R Rombach","year":"2022","unstructured":"R. Rombach, A. Blattmann, D. Lorenz, P. Esser, B. Ommer. High-resolution image synthesis with latent diffusion models. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, New Orleans, USA, pp. 10674\u201310685, 2022. DOI: https:\/\/doi.org\/10.1109\/CVPR52688.2022.01042."},{"key":"1606_CR11","doi-asserted-by":"publisher","first-page":"14214","DOI":"10.1109\/CVPR52729.2023.01366","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"M Tao","year":"2023","unstructured":"M. Tao, B. K. Bao, H. Tang, C. Xu. GALIP: Generative adversarial clips for text-to-image synthesis. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, Vancouver, Canada, pp. 14214\u201314223, 2023. DOI: https:\/\/doi.org\/10.1109\/CVPR52729.2023.01366."},{"issue":"4","key":"1606_CR12","doi-asserted-by":"publisher","first-page":"491","DOI":"10.1007\/s40319-023-01321-y","volume":"54","author":"A Strowel","year":"2023","unstructured":"A. Strowel. ChatGPT and generative AI tools: Theft of intellectual labor? IIC-International Review of Intellectual Property and Competition Law, vol. 54, no. 4, pp. 491\u2013494, 2023. DOI: https:\/\/doi.org\/10.1007\/s40319-023-01321-y.","journal-title":"IIC-International Review of Intellectual Property and Competition Law"},{"key":"1606_CR13","unstructured":"W. Wu, S. Liu. A comprehensive review and systematic analysis of artificial intelligence regulation policies, [Online], Available: https:\/\/arxiv.org\/abs\/2307.12218, 2023."},{"issue":"8","key":"1606_CR14","doi-asserted-by":"publisher","first-page":"3221","DOI":"10.1007\/s12652-019-01500-1","volume":"11","author":"A Mohanarathinam","year":"2020","unstructured":"A. Mohanarathinam, S. Kamalraj, G. K. D. Prasanna Venkatesan, R. V. Ravi, C. S. Manikandababu. RETRACTED ARTICLE: Digital watermarking techniques for image security: A review. Journal of Ambient Intelligence and Humanized Computing, vol. 11, no. 8, pp. 3221\u20133229, 2020. DOI: https:\/\/doi.org\/10.1007\/s12652-019-01500-1.","journal-title":"Journal of Ambient Intelligence and Humanized Computing"},{"key":"1606_CR15","doi-asserted-by":"publisher","first-page":"1668","DOI":"10.1145\/3581783.3612448","volume-title":"Proceedings of the 31st ACM International Conference on Multimedia","author":"C Xiong","year":"2023","unstructured":"C. Xiong, C. Qin, G. Feng, X. Zhang. Flexible and secure watermarking for latent diffusion model. In Proceedings of the 31st ACM International Conference on Multimedia, Ottawa, Canada, pp. 1668\u20131676, 2023. DOI: https:\/\/doi.org\/10.1145\/3581783.3612448."},{"issue":"9","key":"1606_CR16","doi-asserted-by":"publisher","first-page":"4660","DOI":"10.1109\/TCSVT.2023.3245650","volume":"33","author":"Y Huang","year":"2023","unstructured":"Y. Huang, H. Guan, J. Liu, S. Zhang, B. Niu, G. Zhang. Robust texture-aware local adaptive image watermarking with perceptual guarantee. IEEE Transactions on Circuits and Systems for Video Technology, vol. 33, no. 9, pp. 4660\u20134674, 2023. DOI: https:\/\/doi.org\/10.1109\/tcsvt.2023.3245650.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"key":"1606_CR17","doi-asserted-by":"publisher","first-page":"682","DOI":"10.1007\/978-3-030-01267-0_40","volume-title":"Proceedings of the 15th European Conference on Computer Vision","author":"J Zhu","year":"2018","unstructured":"J. Zhu, R. Kaplan, J. Johnson, F. F. Li. HiDDeN: Hiding data with deep networks. In Proceedings of the 15th European Conference on Computer Vision, Munich, Germany, pp. 682\u2013697, 2018. DOI: https:\/\/doi.org\/10.1007\/978-3-030-01267-0_40."},{"issue":"3","key":"1606_CR18","doi-asserted-by":"publisher","first-page":"613","DOI":"10.1109\/TETCI.2021.3055520","volume":"6","author":"W Ding","year":"2022","unstructured":"W. Ding, Y. Ming, Z. Cao, C. T. Lin. A generalized deep neural network approach for digital watermarking analysis. IEEE Transactions on Emerging Topics in Computational Intelligence, vol. 6, no. 3, pp. 613\u2013627, 2022. DOI: https:\/\/doi.org\/10.1109\/TETCI.2021.3055520.","journal-title":"IEEE Transactions on Emerging Topics in Computational Intelligence"},{"key":"1606_CR19","doi-asserted-by":"publisher","first-page":"271","DOI":"10.1109\/COMSWA.2008.4554423","volume-title":"Proceedings of the 3rd International Conference on Communication Systems software and Middleware and Workshops","author":"K A Navas","year":"2008","unstructured":"K. A. Navas, M. C. Ajay, M. Lekshmi, T. S. Archana, M. Sasikumar. DWT-DCT-SVD based watermarking. In Proceedings of the 3rd International Conference on Communication Systems software and Middleware and Workshops, IEEE, Bangalore, India, pp. 271\u2013274, 2008. DOI: https:\/\/doi.org\/10.1109\/COMSWA.2008.4554423."},{"key":"1606_CR20","doi-asserted-by":"publisher","first-page":"3054","DOI":"10.1109\/ICASSP43922.2022.9746058","volume-title":"Proceedings of IEEE International Conference on Acoustics, Speech and Signal Processing","author":"P Fernandez","year":"2022","unstructured":"P. Fernandez, A. Sablayrolles, T. Furon, H. J\u00e9gou, M. Douze. Watermarking images in self-supervised latent spaces. In Proceedings of IEEE International Conference on Acoustics, Speech and Signal Processing, Singapore, pp. 3054\u20133058, 2022. DOI: https:\/\/doi.org\/10.1109\/ICASSP43922.2022.9746058."},{"key":"1606_CR21","doi-asserted-by":"publisher","first-page":"12162","DOI":"10.1109\/CVPR52733.2024.01156","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Z Yang","year":"2024","unstructured":"Z. Yang, K. Zeng, K. Chen, H. Fang, W. Zhang, N. Yu. Gaussian shading: Provable performance-lossless image watermarking for diffusion models. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, Seattle, USA, pp. 12162\u201312171, 2024. DOI: https:\/\/doi.org\/10.1109\/CVPR52733.2024.01156."},{"key":"1606_CR22","volume-title":"Proceedings of the 37th International Conference on Neural Information Processing Systems","author":"Y Wen","year":"2023","unstructured":"Y. Wen, J. Kirchenbauer, J. Geiping, T. Goldstein. Tree-rings watermarks: Invisible fingerprints for diffusion images. In Proceedings of the 37th International Conference on Neural Information Processing Systems, New Orleans, USA, Article number 2529, 2023."},{"key":"1606_CR23","volume-title":"Proceedings of the 38th International Conference on Neural Information Processing Systems","author":"L Zhang","year":"2024","unstructured":"L. Zhang, X. Liu, A. V. Martin, C. X. Bearfield, Y. Brun, H. Guan. Attack-resilient image watermarking using stable diffusion. In Proceedings of the 38th International Conference on Neural Information Processing Systems, Vancouver, Canada, Article number 1215, 2024."},{"issue":"5","key":"1606_CR24","doi-asserted-by":"publisher","first-page":"991","DOI":"10.1007\/s11760-013-0534-2","volume":"9","author":"Q Su","year":"2015","unstructured":"Q. Su, G. Wang, S. Jia, X. Zhang, Q. Liu, X. Liu. Embedding color image watermark in color image based on two-level DCT. Signal, Image and Video Processing, vol. 9, no. 5, pp. 991\u20131007, 2015. DOI: https:\/\/doi.org\/10.1007\/s11760-013-0534-2.","journal-title":"Signal, Image and Video Processing"},{"key":"1606_CR25","doi-asserted-by":"publisher","unstructured":"J. Wang, Z. Du. A method of processing color image watermarking based on the Haar wavelet. Journal of Visual Communication and Image Representation, vol. 64, Article number 102627, 2019. DOI: https:\/\/doi.org\/10.1016\/j.jvcir.2019.102627.","DOI":"10.1016\/j.jvcir.2019.102627"},{"issue":"4","key":"1606_CR26","doi-asserted-by":"publisher","first-page":"20133","DOI":"10.1007\/s11042-019-7326-9","volume":"78","author":"F Zhang","year":"2019","unstructured":"F. Zhang, T. Luo, G. Jiang, M. Yu, H. Xu, W. Zhou. A novel robust color image watermarking method using RGB correlations. Multimedia Tools and Applications, vol. 78, no. 4, pp. 20133\u201320155, 2019. DOI: https:\/\/doi.org\/10.1007\/s11042-019-7326-9.","journal-title":"Multimedia Tools and Applications"},{"issue":"4","key":"1606_CR27","doi-asserted-by":"publisher","first-page":"321","DOI":"10.1504\/IJICS.2018.095296","volume":"10","author":"A K Pal","year":"2018","unstructured":"A. K. Pal, S. Roy. A robust and blind image watermarking scheme in DCT domain. International Journal of Information and Computer Security, vol. 10, no. 4, pp. 321\u2013340, 2018. DOI: https:\/\/doi.org\/10.1504\/IJICS.2018.095296.","journal-title":"International Journal of Information and Computer Security"},{"key":"1606_CR28","doi-asserted-by":"publisher","first-page":"448","DOI":"10.1109\/ICCSP.2016.7754176","volume-title":"Proceedings of International Conference on Communication and Signal Processing","author":"N S Kumaran","year":"2016","unstructured":"N. S. Kumaran, S. Abinaya. Comparison analysis of digital image watermarking using DWT and LSB technique. In Proceedings of International Conference on Communication and Signal Processing, IEEE, Melmaruvathur, India, pp. 448\u2013451, 2016. DOI: https:\/\/doi.org\/10.1109\/ICCSP.2016.7754176."},{"issue":"12","key":"1606_CR29","doi-asserted-by":"publisher","first-page":"17027","DOI":"10.1007\/s11042-018-7085-z","volume":"78","author":"A K Abdulrahman","year":"2019","unstructured":"A. K. Abdulrahman, S. Ozturk. A novel hybrid DCT and DWT based robust watermarking algorithm for color images. Multimedia Tools and Applications, vol. 78, no. 12, pp. 17027\u201317049, 2019. DOI: https:\/\/doi.org\/10.1007\/s11042-018-7085-z.","journal-title":"Multimedia Tools and Applications"},{"issue":"1","key":"1606_CR30","doi-asserted-by":"publisher","first-page":"89","DOI":"10.1007\/s11633-023-1455-3","volume":"21","author":"C Wu","year":"2024","unstructured":"Chaojie Wu, Mingyang Li, Ying Gao, Xinyan Xie, Wing W. Y. Ng, Ahmad Musyafa. Weakly Supervised Object Localization with Background Suppression Erasing for Art Authentication and Copyright Protection. Machine Intelligence Research, vol. 21, no. 1, pp. 89\u2013103, 2024. DOI: https:\/\/doi.org\/10.1007\/s11633-023-1455-3.","journal-title":"Machine Intelligence Research"},{"key":"1606_CR31","doi-asserted-by":"publisher","first-page":"234","DOI":"10.1007\/978-3-319-24574-4_28","volume-title":"Proceedings of the 18th International Conference on Medical Image Computing and Computer-assisted intervention","author":"O Ronneberger","year":"2015","unstructured":"O. Ronneberger, P. Fischer, T. Brox. U-Net: Convolutional networks for biomedical image segmentation. In Proceedings of the 18th International Conference on Medical Image Computing and Computer-assisted intervention, Munich, Germany, pp. 234\u2013241, 2015. DOI: https:\/\/doi.org\/10.1007\/978-3-319-24574-4_28."},{"key":"1606_CR32","doi-asserted-by":"publisher","DOI":"10.1093\/oso\/9780199247851.001.0001","volume-title":"Foundations of Non-cooperative Game Theory","author":"K Ritzberger","year":"2002","unstructured":"K. Ritzberger. Foundations of Non-cooperative Game Theory, Oxford, UK: Oxford University Press, 2002."},{"issue":"8","key":"1606_CR33","doi-asserted-by":"publisher","first-page":"6847","DOI":"10.1109\/TCSVT.2023.3294261","volume":"34","author":"Y Wang","year":"2024","unstructured":"Y. Wang, W. Zhou, J. Bao, W. Wang, L. Li, H. Li. CLIP2GAN: Toward bridging text with the latent space of GANs. IEEE Transactions on Circuits and Systems for Video Technology, vol. 34, no. 8, pp. 6847\u20136859, 2024. DOI: https:\/\/doi.org\/10.1109\/tcsvt.2023.3294261.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"issue":"11","key":"1606_CR34","doi-asserted-by":"publisher","first-page":"4258","DOI":"10.1109\/TCSVT.2019.2953753","volume":"30","author":"M Yuan","year":"2020","unstructured":"M. Yuan, Y. Peng. Bridge-GAN: Interpretable representation learning for text-to-image synthesis. IEEE Transactions on Circuits and Systems for Video Technology, vol. 30, no. 11, pp. 4258\u20134268, 2020. DOI: https:\/\/doi.org\/10.1109\/TCSVT.2019.2953753.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"issue":"7","key":"1606_CR35","doi-asserted-by":"publisher","first-page":"5400","DOI":"10.1109\/TCSVT.2023.3347971","volume":"34","author":"H Tan","year":"2024","unstructured":"H. Tan, B. Yin, K. Xu, H. Wang, X. Liu, X. Li. Attention-bridged modal interaction for text-to-image generation. IEEE Transactions on Circuits and Systems for Video Technology, vol. 34, no. 7, pp. 5400\u20135413, 2024. DOI: https:\/\/doi.org\/10.1109\/TCSVT.2023.3347971.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"issue":"8","key":"1606_CR36","doi-asserted-by":"publisher","first-page":"6860","DOI":"10.1109\/TCSVT.2024.3369757","volume":"34","author":"H Chen","year":"2024","unstructured":"H. Chen, Y. Zhang, X. Wang, X. Duan, Y. Zhou, W. Zhu. DisenDreamer: Subject-driven text-to-image generation with sample-aware disentangled tuning. IEEE Transactions on Circuits and Systems for Video Technology, vol. 34, no. 8, pp. 6860\u20136873, 2024. DOI: https:\/\/doi.org\/10.1109\/TCSVT.2024.3369757.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"issue":"8","key":"1606_CR37","doi-asserted-by":"publisher","first-page":"1947","DOI":"10.1109\/TPAMI.2018.2856256","volume":"41","author":"H Zhang","year":"2019","unstructured":"H. Zhang, T. Xu, H. Li, S. Zhang, X. Wang, X. Huang, D. N. Metaxas. StackGAN++: Realistic image synthesis with stacked generative adversarial networks. IEEE Transactions on Pattern Analysis and Machine Intelligence, vol. 41, no. 8, pp. 1947\u20131962, 2019. DOI: https:\/\/doi.org\/10.1109\/TPAMI.2018.2856256.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"1606_CR38","doi-asserted-by":"publisher","first-page":"1316","DOI":"10.1109\/CVPR.2018.00143","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"T Xu","year":"2018","unstructured":"T. Xu, P. Zhang, Q. Huang, H. Zhang, Z. Gan, X. Huang, X. He. AttnGAN: Fine-grained text to image generation with attentional generative adversarial networks. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, Salt Lake City, USA, pp. 1316\u20131324, 2018. DOI: https:\/\/doi.org\/10.1109\/CVPR.2018.00143."},{"key":"1606_CR39","first-page":"1060","volume-title":"Proceedings of the 33rd International Conference on International Conference on Machine Learning","author":"S Reed","year":"2016","unstructured":"S. Reed, Z. Akata, X. Yan, L. Logeswaran, B. Schiele, H. Lee. Generative adversarial text to image synthesis. In Proceedings of the 33rd International Conference on International Conference on Machine Learning, New York, USA, pp. 1060\u20131069, 2016."},{"key":"1606_CR40","doi-asserted-by":"publisher","first-page":"5707","DOI":"10.1109\/ICCV.2017.608","volume-title":"Proceedings of IEEE International Conference on Computer Vision","author":"H Dong","year":"2017","unstructured":"H. Dong, S. Yu, C. Wu, Y. Guo. Semantic image synthesis via adversarial learning. In Proceedings of IEEE International Conference on Computer Vision, Venice, Italy, pp. 5707\u20135715, 2017. DOI: https:\/\/doi.org\/10.1109\/ICCV.2017.608."},{"key":"1606_CR41","doi-asserted-by":"publisher","DOI":"10.1609\/aaai.v33i01.33013272","volume-title":"Proceedings of the 33rd AAAI Conference on Artificial Intelligence","author":"M Cha","year":"2019","unstructured":"M. Cha, Y. L. Gwon, H. T. Kung. Adversarial learning of semantic relevance in text to image synthesis. In Proceedings of the 33rd AAAI Conference on Artificial Intelligence, Honolulu, USA, Article number 402, 2019. DOI: https:\/\/doi.org\/10.1609\/aaai.v33i01.33013272."},{"key":"1606_CR42","doi-asserted-by":"publisher","first-page":"1275","DOI":"10.1109\/TIP.2020.3026728","volume":"30","author":"H Tan","year":"2021","unstructured":"H. Tan, X. Liu, M. Liu, B. Yin, X. Li. KT-GAN: Knowledge-transfer generative adversarial network for text-to-image synthesis. IEEE Transactions on Image Processing, vol. 30, pp. 1275\u20131290, 2021. DOI: https:\/\/doi.org\/10.1109\/TIP.2020.3026728.","journal-title":"IEEE Transactions on Image Processing"},{"key":"1606_CR43","first-page":"6000","volume-title":"Proceedings of the 31st International Conference on Neural Information Processing Systems","author":"A Vaswani","year":"2017","unstructured":"A. Vaswani, N. Shazeer, N. Parmar, J. Uszkoreit, L. Jones, A. N. Gomez, L. Kaiser, I. Polosukhin. Attention is all you need. In Proceedings of the 31st International Conference on Neural Information Processing Systems, Long Beach, USA, pp. 6000\u20136010, 2017."},{"key":"1606_CR44","doi-asserted-by":"publisher","first-page":"2322","DOI":"10.1109\/CVPR.2019.00243","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"G Yin","year":"2019","unstructured":"G. Yin, B. Liu, L. Sheng, N. Yu, X. Wang, J. Shao. Semantics disentangling for text-to-image generation. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, Long Beach, USA, pp. 2322\u20132331, 2019. DOI: https:\/\/doi.org\/10.1109\/CVPR.2019.00243."},{"key":"1606_CR45","first-page":"2234","volume-title":"Proceedings of the 30th International Conference on Neural Information Processing Systems","author":"T Salimans","year":"2016","unstructured":"T. Salimans, I. Goodfellow, W. Zaremba, V. Cheung, A. Radford, X. Chen. Improved techniques for training GANs. In Proceedings of the 30th International Conference on Neural Information Processing Systems, Barcelona, Spain, pp. 2234\u20132242, 2016."},{"key":"1606_CR46","first-page":"6629","volume-title":"Proceedings of the 31st International Conference on Neural Information Processing Systems","author":"M Heusel","year":"2017","unstructured":"M. Heusel, H. Ramsauer, T. Unterthiner, B. Nessler, S. Hochreiter. GANs trained by a two time-scale update rule converge to a local nash equilibrium. In Proceedings of the 31st International Conference on Neural Information Processing Systems, Long Beach, USA, pp. 6629\u20136640, 2017."},{"key":"1606_CR47","doi-asserted-by":"publisher","first-page":"4356","DOI":"10.1109\/TMM.2021.3116416","volume":"24","author":"J Peng","year":"2022","unstructured":"J. Peng, Y. Zhou, X. Sun, L. Cao, Y. Wu, F. Huang, R. Ji. Knowledge-driven generative adversarial network for text-to-image synthesis. IEEE Transactions on Multimedia, vol. 24, pp. 4356\u20134366, 2022. DOI: https:\/\/doi.org\/10.1109\/TMM.2021.3116416.","journal-title":"IEEE Transactions on Multimedia"},{"key":"1606_CR48","first-page":"217","volume-title":"Proceedings of the 30th International Conference on Neural Information Processing Systems","author":"S Reed","year":"2016","unstructured":"S. Reed, Z. Akata, S. Mohan, S. Tenka, B. Schiele, H. Lee. Learning what and where to draw. In Proceedings of the 30th International Conference on Neural Information Processing Systems, Barcelona, Spain, pp. 217\u2013225, 2016."},{"key":"1606_CR49","doi-asserted-by":"publisher","first-page":"18092","DOI":"10.1109\/CVPR52688.2022.01758","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"F Wu","year":"2022","unstructured":"F. Wu, L. Liu, F. Hao, F. He, J. Cheng. Text-to-image synthesis based on object-guided joint-decoding transformer. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, New Orleans, USA, pp. 18092\u201318101, 2022. DOI: https:\/\/doi.org\/10.1109\/CVPR52688.2022.01758."},{"key":"1606_CR50","volume-title":"Proceedings of International Conference on Learning Representations","author":"T Hinz","year":"2019","unstructured":"T. Hinz, S. Heinrich, S. Wermter. Generating multiple objects at spatially distinct locations. In Proceedings of International Conference on Learning Representations, New Orleans, USA, 2019."},{"key":"1606_CR51","doi-asserted-by":"publisher","first-page":"12166","DOI":"10.1109\/CVPR.2019.01245","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"W Li","year":"2019","unstructured":"W. Li, P. Zhang, L. Zhang, Q. Huang, X. He, S. Lyu, J. Gao. Object-driven text-to-image synthesis via adversarial training. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, Long Beach, USA, pp. 12166\u201312174, 2019. DOI: https:\/\/doi.org\/10.1109\/CVPR.2019.01245."},{"key":"1606_CR52","doi-asserted-by":"publisher","first-page":"8576","DOI":"10.1109\/CVPR.2019.00878","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"B Zhao","year":"2019","unstructured":"B. Zhao, L. Meng, W. Yin, L. Sigal. Image generation from layout. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, Long Beach, USA, pp. 8576\u20138585, 2019. DOI: https:\/\/doi.org\/10.1109\/CVPR.2019.00878."},{"issue":"3","key":"1606_CR53","doi-asserted-by":"publisher","first-page":"1552","DOI":"10.1109\/TPAMI.2020.3021209","volume":"44","author":"T Hinz","year":"2022","unstructured":"T. Hinz, S. Heinrich, S. Wermter. Semantic object accuracy for generative text-to-image synthesis. IEEE Transactions on Pattern Analysis and Machine Intelligence, vol. 44, no. 3, pp. 1552\u20131565, 2022. DOI: https:\/\/doi.org\/10.1109\/TPAMI.2020.3021209.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"1606_CR54","volume-title":"Proceedings of the 33rd International Conference on Neural Information Processing Systems","author":"T Qiao","year":"2019","unstructured":"T. Qiao, J. Zhang, D. Xu, D. Tao. Learn, imagine and create: Text-to-image generation from prior knowledge. In Proceedings of the 33rd International Conference on Neural Information Processing Systems, Vancouver, Canada, Article number 80, 2019."},{"key":"1606_CR55","doi-asserted-by":"publisher","first-page":"5795","DOI":"10.1109\/CVPR.2019.00595","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"M Zhu","year":"2019","unstructured":"M. Zhu, P. Pan, W. Chen, Y. Yang. DM-GAN: Dynamic memory generative adversarial networks for text-to-image synthesis. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, Long Beach, USA, pp. 5795\u20135803, 2019. DOI: https:\/\/doi.org\/10.1109\/CVPR.2019.00595."},{"issue":"2","key":"1606_CR56","doi-asserted-by":"publisher","first-page":"1602","DOI":"10.1109\/TCSVT.2024.3474029","volume":"35","author":"K Wang","year":"2025","unstructured":"K. Wang, S. Wu, X. Yin, W. Lu, X. Luo, R. Yang. Robust image watermarking with synchronization using template enhanced-extracted network. IEEE Transactions on Circuits and Systems for Video Technology, vol. 35, no. 2, pp. 1602\u20131614, 2025. DOI: https:\/\/doi.org\/10.1109\/TCSVT.2024.3474029.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"issue":"7","key":"1606_CR57","doi-asserted-by":"publisher","first-page":"2591","DOI":"10.1109\/TCSVT.2020.3030671","volume":"31","author":"H Wu","year":"2021","unstructured":"H. Wu, G. Liu, Y. Yao, X. Zhang. Watermarking neural networks with watermarked images. IEEE Transactions on Circuits and Systems for Video Technology, vol. 31, no. 7, pp. 2591\u20132601, 2021. DOI: https:\/\/doi.org\/10.1109\/TCSVT.2020.3030671.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"key":"1606_CR58","doi-asserted-by":"publisher","first-page":"3629","DOI":"10.1109\/CVPR46437.2021.00363","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"D S Ong","year":"2021","unstructured":"D. S. Ong, C. S. Chan, K. W. Ng, L. Fan, Q. Yang. Protecting intellectual property of generative adversarial networks from ambiguity attacks. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, Nashville, USA, pp. 3629\u20133638, 2021. DOI: https:\/\/doi.org\/10.1109\/CVPR46437.2021.00363."},{"key":"1606_CR59","doi-asserted-by":"publisher","first-page":"14428","DOI":"10.1109\/ICCV48922.2021.01418","volume-title":"Proceedings of IEEE\/CVF International Conference on Computer Vision","author":"N Yu","year":"2021","unstructured":"N. Yu, V. Skripniuk, S. Abdelnabi, M. Fritz. Artificial fingerprinting for generative models: Rooting deepfake attribution in training data. In Proceedings of IEEE\/CVF International Conference on Computer Vision, IEEE, Montreal, Canada, pp. 14428\u201314437, 2021. DOI: https:\/\/doi.org\/10.1109\/ICCV48922.2021.01418."},{"key":"1606_CR60","doi-asserted-by":"publisher","first-page":"247","DOI":"10.1016\/j.cose.2016.11.016","volume":"65","author":"H Kandi","year":"2017","unstructured":"H. Kandi, D. Mishra, S. R. K. S. Gorthi. Exploring the learning capabilities of convolutional neural networks for robust image watermarking. Computers & Security, vol. 65, pp. 247\u2013268, 2017. DOI: https:\/\/doi.org\/10.1016\/j.cose.2016.11.016.","journal-title":"Computers & Security"},{"key":"1606_CR61","unstructured":"K. Zhang, A. Cuesta-Infante, L. Xu, K. Veeramachaneni, SteganoGAN: High Capacity Image Steganography with GANs, [Online], Available: https:\/\/arxiv.org\/abs\/1901.03892."},{"key":"1606_CR62","doi-asserted-by":"publisher","first-page":"2114","DOI":"10.1109\/CVPR42600.2020.00219","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"M Tancik","year":"2020","unstructured":"M. Tancik, B. Mildenhall, R. Ng. StegaStamp: Invisible hyperlinks in physical photographs. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, Seattle, USA, pp. 2114\u20132123, 2020. DOI: https:\/\/doi.org\/10.1109\/CVPR42600.2020.00219."},{"issue":"1","key":"1606_CR63","doi-asserted-by":"publisher","first-page":"122","DOI":"10.1016\/S1665-6423(14)71612-8","volume":"12","author":"H Tao","year":"2014","unstructured":"H. Tao, C. M. Li, J. M. Zain, A. N. Abdalla. Robust image watermarking theories and techniques: A review. Journal of Applied Research and Technology, vol. 12, no. 1, pp. 122\u2013138, 2014. DOI: https:\/\/doi.org\/10.1016\/S1665-6423(14)71612-8.","journal-title":"Journal of Applied Research and Technology"},{"key":"1606_CR64","doi-asserted-by":"publisher","first-page":"596","DOI":"10.1007\/978-3-540-24672-5_47","volume-title":"Proceedings of the 8th European Conference on Computer Vision","author":"D B Russakoff","year":"2004","unstructured":"D. B. Russakoff, C. Tomasi, T. Rohlfing, C. R.Jr. Maurer. Image similarity using mutual information of regions. In Proceedings of the 8th European Conference on Computer Vision, Prague, Czech Republic, pp. 596\u2013607, 2004. DOI: https:\/\/doi.org\/10.1007\/978-3-540-24672-5_47."},{"key":"1606_CR65","doi-asserted-by":"publisher","first-page":"722","DOI":"10.1109\/ICVGIP.2008.47","volume-title":"Proceedings of the 6th Indian Conference on Computer Vision, Graphics & Image Processing","author":"M E Nilsback","year":"2008","unstructured":"M. E. Nilsback, A. Zisserman. Automated flower classification over a large number of classes. In Proceedings of the 6th Indian Conference on Computer Vision, Graphics & Image Processing, IEEE, Bhubaneswar, India, pp. 722\u2013729, 2008. DOI: https:\/\/doi.org\/10.1109\/ICVGIP.2008.47."},{"key":"1606_CR66","unstructured":"C. Wah, S. Branson, P. Welinder, P. Perona, S. Belongie. Caltech-UCSD Birds-200-2011 Dataset, [Online], Available: https:\/\/ieee-dataport.org\/documents\/caltech-ucsd-birds-200-2011, 2011."},{"key":"1606_CR67","doi-asserted-by":"publisher","first-page":"740","DOI":"10.1007\/978-3-319-10602-1_48","volume-title":"Proceedings of the 13th European Conference on Computer Vision","author":"T Y Lin","year":"2014","unstructured":"T. Y. Lin, M. Maire, S. Belongie, J. Hays, P. Perona, D. Ramanan, P. Doll\u00e1r, C. L. Zitnick. Microsoft COCO: Common objects in context. In Proceedings of the 13th European Conference on Computer Vision, Zurich, Switzerland, pp. 740\u2013755, 2014. DOI: https:\/\/doi.org\/10.1007\/978-3-319-10602-1_48."},{"key":"1606_CR68","doi-asserted-by":"publisher","first-page":"586","DOI":"10.1109\/CVPR.2018.00068","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"R Zhang","year":"2018","unstructured":"R. Zhang, P. Isola, A. A. Efros, E. Shechtman, O. Wang. The unreasonable effectiveness of deep features as a perceptual metric. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, IEEE, Salt Lake City, USA pp. 586\u2013595, 2018. DOI: https:\/\/doi.org\/10.1109\/CVPR.2018.00068."},{"key":"1606_CR69","doi-asserted-by":"publisher","DOI":"10.1109\/ICPECTS56089.2022.10047136","volume-title":"Proceedings of International Conference on Power, Energy, Control and Transmission Systems","author":"C Jeeva","year":"2022","unstructured":"C. Jeeva, T. Porselvi, B. Krithika, R. Shreya, G. S. Priyaa, K. Sivasankari. Intelligent image text reader using easy OCR, NRCLex & NLTK. In Proceedings of International Conference on Power, Energy, Control and Transmission Systems, IEEE, Chennai, India, 2022. DOI: https:\/\/doi.org\/10.1109\/ICPECTS56089.2022.10047136."},{"issue":"9","key":"1606_CR70","doi-asserted-by":"publisher","first-page":"926","DOI":"10.1109\/34.232078","volume":"15","author":"A Marzal","year":"1993","unstructured":"A. Marzal, E. Vidal. Computation of normalized edit distance and applications. IEEE Transactions on Pattern Analysis and Machine Intelligence, vol. 15, no. 9, pp. 926\u2013932, 1993. DOI: https:\/\/doi.org\/10.1109\/34.232078.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"1606_CR71","volume-title":"Digital Watermarking and Steganography","author":"I J Cox","year":"2008","unstructured":"I. J. Cox, M. L. Miller, J. A. Bloom, J. Fridrich, T. Kalker. Digital Watermarking and Steganography, 2nd ed., San Francisco, USA: Morgan Kaufmann Publishers, 2008.","edition":"2nd ed."},{"key":"1606_CR72","doi-asserted-by":"publisher","unstructured":"J. Huang, T. Luo, L. Li, G. Yang, H. Xu, C. C. Chang. ARWGAN: Attention-guided robust image watermarking model based on GAN. IEEE Transactions on Instrumentation and Measurement, vol. 72, Article number 5018417, 2023. DOI: https:\/\/doi.org\/10.1109\/TIM.2023.3285981.","DOI":"10.1109\/TIM.2023.3285981"},{"issue":"9","key":"1606_CR73","doi-asserted-by":"publisher","first-page":"14371","DOI":"10.1109\/TITS.2025.3550120","volume":"26","author":"F Kou","year":"2025","unstructured":"F. Kou, Y. Yao, J. Han, J. Wang, H. Li, X. Li, J. Zhang. DualFocus GAN for robust watermarking in transportation cyber-physical systems. IEEE Transactions on Intelligent Transportation Systems, vol. 26, no. 9, pp. 14371\u201314382, 2025. DOI: https:\/\/doi.org\/10.1109\/TITS.2025.3550120.","journal-title":"IEEE Transactions on Intelligent Transportation Systems"},{"key":"1606_CR74","unstructured":"G. Gao, T. Xu, F. Hua. Robust image watermarking based on generative adversarial networks for copyright protection. Research Square, [Online], Available: https:\/\/assets-eu.researchsquare.com\/files\/rs-4039149\/v1_covered_1699b3eb-c849-459f-b4b3-ee874887721f.pdf."},{"issue":"3","key":"1606_CR75","doi-asserted-by":"publisher","first-page":"218","DOI":"10.1038\/s41566-022-01141-5","volume":"17","author":"M Yako","year":"2023","unstructured":"M. Yako, Y. Yamaoka, T. Kiyohara, C. Hosokawa, A. Noda, K. Tack, N. Spooren, T. Hirasawa, A. Ishikawa. Video-rate hyperspectral camera based on a CMOS-compatible random array of Fabry\u2013P\u00e9rot filters. Nature Photonics, vol. 17, no. 3, pp. 218\u2013223, 2023. DOI: https:\/\/doi.org\/10.1038\/s41566-022-01141-5.","journal-title":"Nature Photonics"},{"key":"1606_CR76","doi-asserted-by":"publisher","DOI":"10.1109\/ICME52920.2022.9859725","volume-title":"Proceedings of IEEE International Conference on Multimedia and Expo","author":"J Lu","year":"2022","unstructured":"J. Lu, J. Ni, W. Su, H. Xie. Wavelet-based CNN for robust and high-capacity image watermarking. In Proceedings of IEEE International Conference on Multimedia and Expo, Taipei, China, 2022. DOI: https:\/\/doi.org\/10.1109\/ICME52920.2022.9859725."}],"container-title":["Machine Intelligence Research"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11633-025-1606-9.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11633-025-1606-9","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11633-025-1606-9.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,7,23]],"date-time":"2026-07-23T08:02:29Z","timestamp":1784793749000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11633-025-1606-9"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,6,11]]},"references-count":76,"journal-issue":{"issue":"4","published-print":{"date-parts":[[2026,8]]}},"alternative-id":["1606"],"URL":"https:\/\/doi.org\/10.1007\/s11633-025-1606-9","relation":{},"ISSN":["2731-538X","2731-5398"],"issn-type":[{"value":"2731-538X","type":"print"},{"value":"2731-5398","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,6,11]]},"assertion":[{"value":"3 June 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"9 October 2025","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"11 June 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}}]}}