{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,1,30]],"date-time":"2025-01-30T16:40:05Z","timestamp":1738255205122,"version":"3.35.0"},"reference-count":73,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2024,12,20]],"date-time":"2024-12-20T00:00:00Z","timestamp":1734652800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,12,20]],"date-time":"2024-12-20T00:00:00Z","timestamp":1734652800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"Special Project on Regional Collaborative Innovation in Xinjiang Uygur Autonomous Region","award":["2022E02035"],"award-info":[{"award-number":["2022E02035"]}]},{"name":"Hubei Province Key Research and Development Special Project of Science and Technology Innovation Plan","award":["2023BAB087"],"award-info":[{"award-number":["2023BAB087"]}]},{"name":"Hubei Provincial Administration of Traditional Chinese Medicine Research Project on Traditional Chinese Medicine","award":["ZY2023M064"],"award-info":[{"award-number":["ZY2023M064"]}]},{"name":"Wuhan knowledge innovation special Dawn project","award":["2023010201020465"],"award-info":[{"award-number":["2023010201020465"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Appl Intell"],"published-print":{"date-parts":[[2025,2]]},"DOI":"10.1007\/s10489-024-06043-3","type":"journal-article","created":{"date-parts":[[2024,12,20]],"date-time":"2024-12-20T07:20:25Z","timestamp":1734679225000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["MSDM: multi-space diffusion with dynamic loss weight"],"prefix":"10.1007","volume":"55","author":[{"given":"Zhou","family":"Liu","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Zheng","family":"Ye","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Jing","family":"Liu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Jun","family":"Qin","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Ben","family":"He","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Cathal","family":"Gurrin","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2024,12,20]]},"reference":[{"key":"6043_CR1","doi-asserted-by":"publisher","unstructured":"Capel EH, Dumas J (2023) Denoising diffusion probabilistic models for probabilistic energy forecasting. In: 2023 IEEE Belgrade PowerTech, pp 1\u20136. https:\/\/doi.org\/10.1109\/powertech55446.2023.10202713. IEEE","DOI":"10.1109\/powertech55446.2023.10202713"},{"key":"6043_CR2","doi-asserted-by":"publisher","unstructured":"Goodfellow IJ, Pouget-Abadie J, Mirza M, Xu B, Warde-Farley D, Ozair S, Courville A, Bengio Y (2014) Generative adversarial nets. In: Proceedings of the 27th International Conference on Neural Information Processing Systems - vol 2, NIPS\u201914, pp 2672\u20132680. MIT Press, Cambridge, MA, USA. https:\/\/doi.org\/10.5555\/2969033.2969125","DOI":"10.5555\/2969033.2969125"},{"key":"6043_CR3","doi-asserted-by":"publisher","unstructured":"Ramesh A, Dhariwal P, Nichol A, Chu C, Chen M (2022) Hierarchical text-conditional image generation with clip latents. arXiv preprint arXiv:2204.06125. 1(2):3. https:\/\/doi.org\/10.48550\/arXiv.2204.06125","DOI":"10.48550\/arXiv.2204.06125"},{"key":"6043_CR4","doi-asserted-by":"publisher","first-page":"36479","DOI":"10.5555\/3600270.3602913","volume":"35","author":"C Saharia","year":"2022","unstructured":"Saharia C, Chan W, Saxena S, Li L, Whang J, Denton EL, Ghasemipour K, Gontijo Lopes R, Karagol Ayan B, Salimans T et al (2022) Photorealistic text-to-image diffusion models with deep language understanding. Adv Neural Inf Process Syst 35:36479\u201336494. https:\/\/doi.org\/10.5555\/3600270.3602913","journal-title":"Adv Neural Inf Process Syst"},{"key":"6043_CR5","doi-asserted-by":"publisher","unstructured":"Tewel Y, Shalev Y, Schwartz I, Wolf L (2022) Zerocap: Zero-shot image-to-text generation for visual-semantic arithmetic. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp 17918\u201317928. https:\/\/doi.org\/10.1109\/cvpr52688.2022.01739","DOI":"10.1109\/cvpr52688.2022.01739"},{"key":"6043_CR6","unstructured":"Xu M, Yu L, Song Y, Shi C, Ermon S, Tang J (2021) Geodiff: A geometric diffusion model for molecular conformation generation. In: International Conference on Learning Representations"},{"key":"6043_CR7","first-page":"21696","volume":"34","author":"D Kingma","year":"2021","unstructured":"Kingma D, Salimans T, Poole B, Ho J (2021) Variational diffusion models. Adv Neural Inf Process Syst 34:21696\u201321707","journal-title":"Adv Neural Inf Process Syst"},{"key":"6043_CR8","doi-asserted-by":"publisher","DOI":"10.1101\/2022.11.18.517004","author":"Y Takagi","year":"2022","unstructured":"Takagi Y, Nishimoto S (2022). High-resolution image reconstruction with latent diffusion models from human brain activity. https:\/\/doi.org\/10.1101\/2022.11.18.517004","journal-title":"High-resolution image reconstruction with latent diffusion models from human brain activity."},{"key":"6043_CR9","doi-asserted-by":"publisher","unstructured":"Meng C, Rombach R, Gao R, Kingma D, Ermon S, Ho J, Salimans T (2023) On distillation of guided diffusion models. In: 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR), pp 14297\u201314306. https:\/\/doi.org\/10.1109\/CVPR52729.2023.01374","DOI":"10.1109\/CVPR52729.2023.01374"},{"issue":"1","key":"6043_CR10","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1080\/1206212X.2022.2143026","volume":"2","author":"J An","year":"2015","unstructured":"An J, Cho S (2015) Variational autoencoder based anomaly detection using reconstruction probability. Special lecture on IE 2(1):1\u201318. https:\/\/doi.org\/10.1080\/1206212X.2022.2143026","journal-title":"Special lecture on IE"},{"key":"6043_CR11","doi-asserted-by":"publisher","unstructured":"Yang R, Mandt S (2022) Lossy image compression with conditional diffusion models. https:\/\/doi.org\/10.48550\/arXiv.2209.06950","DOI":"10.48550\/arXiv.2209.06950"},{"key":"6043_CR12","unstructured":"Salimans T, Ho J (2021) Progressive distillation for fast sampling of diffusion models. In: International Conference on Learning Representations"},{"key":"6043_CR13","doi-asserted-by":"publisher","unstructured":"Hang T, Gu S, Li C, Bao J, Chen D, Hu H, Geng X, Guo B (2023) Efficient diffusion training via min-snr weighting strategy. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp 7441\u20137451. https:\/\/doi.org\/10.1109\/iccv51070.2023.00684","DOI":"10.1109\/iccv51070.2023.00684"},{"key":"6043_CR14","doi-asserted-by":"publisher","unstructured":"You A, Zhou C, Zhang Q, Xu L (2021) Towards controllable and photorealistic region-wise image manipulation. In: Proceedings of the 29th ACM International Conference on Multimedia, pp 535\u2013543.https:\/\/doi.org\/10.1145\/3474085.3475206","DOI":"10.1145\/3474085.3475206"},{"key":"6043_CR15","unstructured":"Song J, Meng C, Ermon S (2020) Denoising diffusion implicit models. In: International Conference on Learning Representations"},{"key":"6043_CR16","doi-asserted-by":"publisher","unstructured":"Patel Y, Appalaraju S, Manmatha R (2019) Deep perceptual compression. https:\/\/doi.org\/10.48550\/arXiv.1907.08310","DOI":"10.48550\/arXiv.1907.08310"},{"key":"6043_CR17","unstructured":"Nichol AQ, Dhariwal P (2021) Improved denoising diffusion probabilistic models. In: International Conference on Machine Learning, pp 8162\u20138171. PMLR"},{"issue":"1","key":"6043_CR18","doi-asserted-by":"publisher","first-page":"47","DOI":"10.1109\/tci.2016.2644865","volume":"3","author":"H Zhao","year":"2017","unstructured":"Zhao H, Gallo O, Frosio I, Kautz J (2017) Loss functions for image restoration with neural networks. IEEE Trans Comput Imaging 3(1):47\u201357. https:\/\/doi.org\/10.1109\/tci.2016.2644865","journal-title":"IEEE Trans Comput Imaging"},{"key":"6043_CR19","doi-asserted-by":"publisher","unstructured":"Papagiannis G, Li Y (2022) Imitation learning with sinkhorn distances. In: Joint European Conference on Machine Learning and Knowledge Discovery in Databases, pp 116\u2013131. https:\/\/doi.org\/10.1007\/978-3-031-26412-2sps8 . Springer","DOI":"10.1007\/978-3-031-26412-2sps8"},{"key":"6043_CR20","doi-asserted-by":"publisher","unstructured":"Kendall A, Gal Y, Cipolla R (2018) Multi-task learning using uncertainty to weigh losses for scene geometry and semantics. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp 7482\u20137491. https:\/\/doi.org\/10.1109\/cvpr.2018.00781","DOI":"10.1109\/cvpr.2018.00781"},{"key":"6043_CR21","doi-asserted-by":"publisher","unstructured":"Sener O, Koltun V (2018) Multi-task learning as multi-objective optimization. In: Proceedings of the 32nd International Conference on Neural Information Processing Systems. NIPS\u201918, pp. 525\u2013536. Curran Associates Inc., Red Hook, NY, USA. https:\/\/doi.org\/10.5555\/3326943.3326992","DOI":"10.5555\/3326943.3326992"},{"issue":"5\u20136","key":"6043_CR22","doi-asserted-by":"publisher","first-page":"313","DOI":"10.1016\/J.CRMA.2012.03.014","volume":"350","author":"JA D\u00e9sid\u00e9ri","year":"2012","unstructured":"D\u00e9sid\u00e9ri JA (2012) Multiple-gradient descent algorithm (mgda) for multiobjective optimization. C R Math 350(5\u20136):313\u2013318. https:\/\/doi.org\/10.1016\/J.CRMA.2012.03.014","journal-title":"C R Math"},{"issue":"3","key":"6043_CR23","doi-asserted-by":"publisher","first-page":"516","DOI":"10.1080\/0305215x.2017.1327579","volume":"50","author":"A Mart\u00edn","year":"2017","unstructured":"Mart\u00edn A, Sch\u00fctze O (2017) Pareto tracer: a predictor\u2013corrector method for multi-objective optimization problems. Eng Optim 50(3):516\u2013536. https:\/\/doi.org\/10.1080\/0305215x.2017.1327579","journal-title":"Eng Optim"},{"key":"6043_CR24","doi-asserted-by":"publisher","first-page":"411","DOI":"10.1016\/b978-0-08-017639-0.50008-x","volume":"2","author":"B Riemann","year":"1854","unstructured":"Riemann B (1854) On the hypotheses which lie at the foundations of geometry. A source book in mathematics 2:411\u2013425. https:\/\/doi.org\/10.1016\/b978-0-08-017639-0.50008-x","journal-title":"A source book in mathematics"},{"key":"6043_CR25","doi-asserted-by":"publisher","unstructured":"Zhu J, Shen Y, Zhao D, Zhou B (2020) In-domain gan inversion for real image editing. In: European Conference on Computer Vision, pp 592\u2013608. https:\/\/doi.org\/10.1007\/978-3-030-58520-4sps35 . Springer","DOI":"10.1007\/978-3-030-58520-4sps35"},{"key":"6043_CR26","doi-asserted-by":"publisher","unstructured":"Zhang Y, Huang N, Tang F, Huang H, Ma C, Dong W, Xu C (2023) Inversion-based style transfer with diffusion models. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp 10146\u201310156. https:\/\/doi.org\/10.1109\/cvpr52729.2023.00978","DOI":"10.1109\/cvpr52729.2023.00978"},{"key":"6043_CR27","doi-asserted-by":"publisher","unstructured":"Lin H, Cheng X, Wu X, Shen D (2022) Cat: Cross attention in vision transformer. In: 2022 IEEE International Conference on Multimedia and Expo (ICME), pp 1\u20136. https:\/\/doi.org\/10.1109\/ICME52920.2022.9859720 . IEEE","DOI":"10.1109\/ICME52920.2022.9859720"},{"key":"6043_CR28","doi-asserted-by":"publisher","unstructured":"Chen Z, Badrinarayanan V, Lee CY, Rabinovich A (2018) Gradnorm: Gradient normalization for adaptive loss balancing in deep multitask networks. In: International Conference on Machine Learning, pp 794\u2013803. https:\/\/doi.org\/10.48550\/arXiv.1711.02257 . PMLR","DOI":"10.48550\/arXiv.1711.02257"},{"issue":"4","key":"6043_CR29","doi-asserted-by":"publisher","first-page":"600","DOI":"10.1109\/TIP.2003.819861","volume":"13","author":"Z Wang","year":"2004","unstructured":"Wang Z, Bovik AC, Sheikh HR, Simoncelli EP (2004) Image quality assessment: from error visibility to structural similarity. IEEE Trans Image Process 13(4):600\u2013612","journal-title":"IEEE Trans Image Process"},{"key":"6043_CR30","doi-asserted-by":"publisher","unstructured":"Meyer GP (2021) An alternative probabilistic interpretation of the huber loss. In: Proceedings of the Ieee\/cvf Conference on Computer Vision and Pattern Recognition, pp 5261\u20135269. https:\/\/doi.org\/10.1109\/cvpr46437.2021.00522","DOI":"10.1109\/cvpr46437.2021.00522"},{"key":"6043_CR31","doi-asserted-by":"publisher","unstructured":"Barron JT (2019) A general and adaptive robust loss function. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp 4331\u20134339. https:\/\/doi.org\/10.1109\/cvpr.2019.00446","DOI":"10.1109\/cvpr.2019.00446"},{"key":"6043_CR32","unstructured":"Song Y, Sohl-Dickstein J, Kingma DP, Kumar A, Ermon S, Poole B (2020) Score-based generative modeling through stochastic differential equations. In: International Conference on Learning Representations"},{"key":"6043_CR33","doi-asserted-by":"publisher","unstructured":"Bao F, Nie, S, Xue K, Cao Y, Li C, Su H, Zhu J (2023) All are worth words: A vit backbone for diffusion models. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp 22669\u201322679. https:\/\/doi.org\/10.1109\/CVPR52729.2023.02171","DOI":"10.1109\/CVPR52729.2023.02171"},{"key":"6043_CR34","doi-asserted-by":"publisher","unstructured":"Kingma DP, Welling M (2013) Auto-encoding variational bayes. https:\/\/doi.org\/10.48550\/arXiv.1312.6114","DOI":"10.48550\/arXiv.1312.6114"},{"key":"6043_CR35","doi-asserted-by":"publisher","unstructured":"Van Den Oord A, Vinyals O et al (2017) Neural discrete representation learning. Adv Neural Inf Process Syst 30. https:\/\/doi.org\/10.5555\/3295222.3295378","DOI":"10.5555\/3295222.3295378"},{"key":"6043_CR36","doi-asserted-by":"publisher","unstructured":"Song Y, Ermon S (2019) Generative modeling by estimating gradients of the data distribution. Adv Neural Inf Process Syst 32. https:\/\/doi.org\/10.5555\/3454287.3455354","DOI":"10.5555\/3454287.3455354"},{"key":"6043_CR37","first-page":"12533","volume":"34","author":"A Sinha","year":"2021","unstructured":"Sinha A, Song J, Meng C, Ermon S (2021) D2c: Diffusion-decoding models for few-shot conditional generation. Adv Neural Inf Process Syst 34:12533\u201312548","journal-title":"Adv Neural Inf Process Syst"},{"key":"6043_CR38","doi-asserted-by":"publisher","unstructured":"Ngatchou P, Zarei A, El-Sharkawi A (2005) Pareto multi objective optimization. In: Proceedings of the 13th International Conference On, Intelligent Systems Application to Power Systems, pp 84\u201391. https:\/\/doi.org\/10.1007\/springerreferencesps72504 . IEEE","DOI":"10.1007\/springerreferencesps72504"},{"key":"6043_CR39","doi-asserted-by":"publisher","unstructured":"Jin Y, Olhofer M, Sendhoff B (2001) Dynamic weighted aggregation for evolutionary multi-objective optimization: why does it work and how? In: Proceedings of the 3rd Annual Conference on Genetic and Evolutionary Computation. GECCO\u201901, pp 1042\u20131049. Morgan Kaufmann Publishers Inc., San Francisco, CA, USA. https:\/\/doi.org\/10.5555\/2955239.2955427","DOI":"10.5555\/2955239.2955427"},{"issue":"725\/36","key":"6043_CR40","doi-asserted-by":"publisher","first-page":"725","DOI":"10.1007\/springerreferencesps5696","volume":"10","author":"G Gordon","year":"2012","unstructured":"Gordon G, Tibshirani R (2012) Karush-kuhn-tucker conditions. Optim 10(725\/36):725. https:\/\/doi.org\/10.1007\/springerreferencesps5696","journal-title":"Karush-kuhn-tucker conditions. Optim"},{"issue":"13","key":"6043_CR41","doi-asserted-by":"publisher","first-page":"4602","DOI":"10.3390\/en15134602","volume":"15","author":"S Mirfallah Lialestani","year":"2022","unstructured":"Mirfallah Lialestani S, Parcerisa D, Himi M, Abbaszadeh Shahri A (2022) Generating 3d geothermal maps in catalonia, spain using a hybrid adaptive multitask deep learning procedure. Energ 15(13):4602. https:\/\/doi.org\/10.3390\/en15134602","journal-title":"Energ"},{"issue":"1","key":"6043_CR42","doi-asserted-by":"publisher","first-page":"267","DOI":"10.1007\/s10064-020-01922-8","volume":"80","author":"A Abbaszadeh Shahri","year":"2020","unstructured":"Abbaszadeh Shahri A, Maghsoudi Moud F (2020) Landslide susceptibility mapping using hybridized block modular intelligence model. Bull Eng Geol Environ 80(1):267\u2013284. https:\/\/doi.org\/10.1007\/s10064-020-01922-8","journal-title":"Bull Eng Geol Environ"},{"issue":"4","key":"6043_CR43","doi-asserted-by":"publisher","first-page":"838","DOI":"10.1007\/s11390-018-1859-7","volume":"33","author":"BJ Zou","year":"2018","unstructured":"Zou BJ, Guo YD, He Q, Ouyang PB, Liu K, Chen ZL (2018) 3d filtering by block matching and convolutional neural network for image denoising. J Comput Sci Technol 33(4):838\u2013848. https:\/\/doi.org\/10.1007\/s11390-018-1859-7","journal-title":"J Comput Sci Technol"},{"key":"6043_CR44","doi-asserted-by":"publisher","unstructured":"Zhou J, Ni J, Rao Y (2017) Block-based convolutional neural network for image forgery detection. In: Digital Forensics and Watermarking: 16th International Workshop, IWDW 2017, Magdeburg, Germany, August 23-25, 2017, Proceedings 16, pp 65\u201376. https:\/\/doi.org\/10.1007\/978-3-319-64185-0sps6 . Springer","DOI":"10.1007\/978-3-319-64185-0sps6"},{"issue":"1","key":"6043_CR45","doi-asserted-by":"publisher","first-page":"2383","DOI":"10.1007\/s11042-023-15649-7","volume":"83","author":"M Sabeena","year":"2024","unstructured":"Sabeena M, Abraham L (2024) Convolutional block attention based network for copy-move image forgery detection. Multimed Tools Appl 83(1):2383\u20132405. https:\/\/doi.org\/10.1007\/s11042-023-15649-7","journal-title":"Multimed Tools Appl"},{"key":"6043_CR46","doi-asserted-by":"publisher","first-page":"211","DOI":"10.1007\/s11263-015-0816-y","volume":"115","author":"O Russakovsky","year":"2015","unstructured":"Russakovsky O, Deng J, Su H, Krause J, Satheesh S, Ma S, Huang Z, Karpathy A, Khosla A, Bernstein M et al (2015) Imagenet large scale visual recognition challenge. Int J Comput Vis 115:211\u2013252. https:\/\/doi.org\/10.1007\/s11263-015-0816-y","journal-title":"Int J Comput Vis"},{"key":"6043_CR47","unstructured":"Hertz A, Mokady R, Tenenbaum J, Aberman K, Pritch Y, Cohen-or D (2022) Prompt-to-prompt image editing with cross-attention control. In: The Eleventh International Conference on Learning Representations"},{"key":"6043_CR48","doi-asserted-by":"publisher","unstructured":"Brooks T, Holynski A, Efros AA (2023) Instructpix2pix: Learning to follow image editing instructions. In: 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR), pp 18392\u201318402. https:\/\/doi.org\/10.1109\/CVPR52729.2023.01764","DOI":"10.1109\/CVPR52729.2023.01764"},{"key":"6043_CR49","doi-asserted-by":"publisher","unstructured":"Ronneberger O, Fischer P, Brox T (2015) U-net: Convolutional networks for biomedical image segmentation. In: Medical Image Computing and Computer-Assisted Intervention\u2013MICCAI 2015: 18th International Conference, Munich, Germany, October 5-9, 2015, Proceedings, Part III 18, pp 234\u2013241. https:\/\/doi.org\/10.1007\/978-3-319-24574-4sps28 . Springer","DOI":"10.1007\/978-3-319-24574-4sps28"},{"key":"6043_CR50","unstructured":"Razavi A, Oord A, Vinyals O (2019) Generating diverse high-fidelity images with vq-vae-2. Adv Neural Inf Process Syst 32"},{"issue":"4","key":"6043_CR51","doi-asserted-by":"publisher","first-page":"193","DOI":"10.1007\/BF00344251","volume":"36","author":"K Fukushima","year":"1980","unstructured":"Fukushima K (1980) Neocognitron: A self-organizing neural network model for a mechanism of pattern recognition unaffected by shift in position. Biol Cybern 36(4):193\u2013202. https:\/\/doi.org\/10.1007\/BF00344251","journal-title":"Biol Cybern"},{"key":"6043_CR52","unstructured":"Krizhevsky A, Hinton G et al (2009) Learning multiple layers of features from tiny images"},{"key":"6043_CR53","doi-asserted-by":"publisher","unstructured":"Nguyen A, Yosinski J, Clune J (2015) Deep neural networks are easily fooled: High confidence predictions for unrecognizable images. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp 427\u2013436. https:\/\/doi.org\/10.1109\/CVPR.2015.7298640","DOI":"10.1109\/CVPR.2015.7298640"},{"key":"6043_CR54","doi-asserted-by":"publisher","unstructured":"Liu Z, Luo P, Wang X, Tang X (2015) Deep learning face attributes in the wild. In: 2015 IEEE International Conference on Computer Vision (ICCV), pp 3730\u20133738. https:\/\/doi.org\/10.1109\/ICCV.2015.425","DOI":"10.1109\/ICCV.2015.425"},{"key":"6043_CR55","unstructured":"Dosovitskiy A, Beyer L, Kolesnikov A, Weissenborn D, Zhai X, Unterthiner T, Dehghani M, Minderer M, Heigold G, Gelly S et al (2020) An image is worth 16x16 words: Transformers for image recognition at scale. In: International Conference on Learning Representations"},{"key":"6043_CR56","doi-asserted-by":"publisher","first-page":"8780","DOI":"10.5555\/3540261.3540933","volume":"34","author":"P Dhariwal","year":"2021","unstructured":"Dhariwal P, Nichol A (2021) Diffusion models beat gans on image synthesis. Adv Neural Inf Process Syst 34:8780\u20138794. https:\/\/doi.org\/10.5555\/3540261.3540933","journal-title":"Adv Neural Inf Process Syst"},{"key":"6043_CR57","doi-asserted-by":"publisher","unstructured":"Wang X, Yu K, Wu S, Gu J, Liu Y, Dong C, Qiao Y, Change\u00a0Loy C (2018) Esrgan: Enhanced super-resolution generative adversarial networks. In: Proceedings of the European Conference on Computer Vision (ECCV) Workshops, pp 0\u20130. https:\/\/doi.org\/10.1007\/978-3-030-11021-5sps5","DOI":"10.1007\/978-3-030-11021-5sps5"},{"key":"6043_CR58","doi-asserted-by":"publisher","unstructured":"Ho J, Saharia C, Chan W, Fleet DJ, Norouzi M, Salimans T (2022) Cascaded diffusion models for high fidelity image generation. J Mach Learn Res 23(1). https:\/\/doi.org\/10.5555\/3586589.3586636","DOI":"10.5555\/3586589.3586636"},{"key":"6043_CR59","doi-asserted-by":"publisher","unstructured":"Heusel M, Ramsauer H, Unterthiner T, Nessler B, Hochreiter S (2017) Gans trained by a two time-scale update rule converge to a local nash equilibrium. Adv Neural Inf Process Syst 30. https:\/\/doi.org\/10.18034\/ajase.v8i1.9","DOI":"10.18034\/ajase.v8i1.9"},{"key":"6043_CR60","doi-asserted-by":"publisher","unstructured":"Salimans T, Goodfellow I, Zaremba W, Cheung V, Radford A., Chen X (2016) Improved techniques for training gans. In: Proceedings of the 30th International Conference on Neural Information Processing Systems. NIPS\u201916, pp 2234\u20132242. Curran Associates Inc., Red Hook, NY, USA.https:\/\/doi.org\/10.5555\/3157096.3157346","DOI":"10.5555\/3157096.3157346"},{"key":"6043_CR61","unstructured":"Loshchilov I, Hutter F (2018) Decoupled weight decay regularization. In: International Conference on Learning Representations"},{"key":"6043_CR62","doi-asserted-by":"publisher","first-page":"32270","DOI":"10.5555\/3600270.3602608","volume":"35","author":"D Kim","year":"2022","unstructured":"Kim D, Na B, Kwon SJ, Lee D, Kang W, Ic Moon (2022) Maximum likelihood training of implicit nonlinear diffusion model. Adv Neural Inf Process Syst 35:32270\u201332284. https:\/\/doi.org\/10.5555\/3600270.3602608","journal-title":"Adv Neural Inf Process Syst"},{"key":"6043_CR63","doi-asserted-by":"publisher","unstructured":"Li S, Chen W, Zeng D (2023) Scire-solver: Efficient sampling of diffusion probabilistic models by score-integrand solver with recursive derivative estimation. https:\/\/doi.org\/10.48550\/arXiv.2308.07896","DOI":"10.48550\/arXiv.2308.07896"},{"key":"6043_CR64","unstructured":"Liu L, Ren Y, Lin Z, Zhao Z (2021) Pseudo numerical methods for diffusion models on manifolds. In: International Conference on Learning Representations"},{"key":"6043_CR65","doi-asserted-by":"publisher","first-page":"26565","DOI":"10.5555\/3600270.3602196","volume":"35","author":"T Karras","year":"2022","unstructured":"Karras T, Aittala M, Aila T, Laine S (2022) Elucidating the design space of diffusion-based generative models. Adv Neural Inf Process Syst 35:26565\u201326577. https:\/\/doi.org\/10.5555\/3600270.3602196","journal-title":"Adv Neural Inf Process Syst"},{"key":"6043_CR66","unstructured":"Zheng H, He P, Chen W, Zhou M (2023) Truncated diffusion probabilistic models and diffusion-based adversarial auto-encoders"},{"key":"6043_CR67","doi-asserted-by":"publisher","unstructured":"Pandey K, Mukherjee A, Rai P, Kumar A (2022) Diffusevae: Efficient, controllable and high-fidelity generation from low-dimensional latents. https:\/\/doi.org\/10.48550\/arXiv.2201.00308","DOI":"10.48550\/arXiv.2201.00308"},{"key":"6043_CR68","unstructured":"Lezama J, Salimans T, Jiang L, Chang H, Ho J, Essa I (2022) Discrete predictor-corrector diffusion models for image synthesis. In: The Eleventh International Conference on Learning Representations"},{"key":"6043_CR69","unstructured":"Radford A, Kim JW, Hallacy C, Ramesh A, Goh G, Agarwal S, Sastry G, Askell A, Mishkin P, Clark J et al (2021) Learning transferable visual models from natural language supervision. In: International Conference on Machine Learning, pp 8748\u20138763. PMLR"},{"key":"6043_CR70","doi-asserted-by":"publisher","unstructured":"Cicirello VA (2024) Evolutionary computation: Theories, techniques, and applications. Appl Sci (2076-3417) 14(6). https:\/\/doi.org\/10.3390\/app14062542","DOI":"10.3390\/app14062542"},{"key":"6043_CR71","doi-asserted-by":"publisher","unstructured":"Joshi S, Pant M, Deep K (2024) Evolutionary techniques in making efficient deep-learning framework: A review. Advanced Machine Learning with Evolutionary and Metaheuristic Techniques 87. https:\/\/doi.org\/10.1007\/978-981-99-9718-3sps4","DOI":"10.1007\/978-981-99-9718-3sps4"},{"key":"6043_CR72","doi-asserted-by":"publisher","unstructured":"Bendel O (2023) Image synthesis from an ethical perspective. AI Soc 1\u201310. https:\/\/doi.org\/10.1007\/s00146-023-01780-4","DOI":"10.1007\/s00146-023-01780-4"},{"key":"6043_CR73","doi-asserted-by":"publisher","first-page":"126","DOI":"10.1016\/j.inffus.2021.02.014","volume":"72","author":"P Shamsolmoali","year":"2021","unstructured":"Shamsolmoali P, Zareapoor M, Granger E, Zhou H, Wang R, Celebi ME, Yang J (2021) Image synthesis with adversarial networks: A comprehensive survey and case studies. Inf Fusion 72:126\u2013146. https:\/\/doi.org\/10.1016\/j.inffus.2021.02.014","journal-title":"Inf Fusion"}],"container-title":["Applied Intelligence"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-024-06043-3.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10489-024-06043-3\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-024-06043-3.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,1,30]],"date-time":"2025-01-30T16:03:12Z","timestamp":1738252992000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10489-024-06043-3"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,12,20]]},"references-count":73,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2025,2]]}},"alternative-id":["6043"],"URL":"https:\/\/doi.org\/10.1007\/s10489-024-06043-3","relation":{},"ISSN":["0924-669X","1573-7497"],"issn-type":[{"type":"print","value":"0924-669X"},{"type":"electronic","value":"1573-7497"}],"subject":[],"published":{"date-parts":[[2024,12,20]]},"assertion":[{"value":"10 September 2024","order":1,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"20 December 2024","order":2,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors of this paper declare that there are no competing interests related to the content of this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}},{"value":"This paper does not involve any ethical issues related to the use of data.","order":3,"name":"Ethics","group":{"name":"EthicsHeading","label":"Ethical and informed consent for data used"}}],"article-number":"193"}}