{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,10]],"date-time":"2026-07-10T03:04:34Z","timestamp":1783652674209,"version":"3.55.0"},"publisher-location":"Cham","reference-count":49,"publisher":"Springer Nature Switzerland","isbn-type":[{"value":"9783031736605","type":"print"},{"value":"9783031736612","type":"electronic"}],"license":[{"start":{"date-parts":[[2024,11,10]],"date-time":"2024-11-10T00:00:00Z","timestamp":1731196800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,11,10]],"date-time":"2024-11-10T00:00:00Z","timestamp":1731196800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2025]]},"DOI":"10.1007\/978-3-031-73661-2_10","type":"book-chapter","created":{"date-parts":[[2024,11,9]],"date-time":"2024-11-09T11:08:35Z","timestamp":1731150515000},"page":"172-188","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":45,"title":["Concept Sliders: LoRA Adaptors for\u00a0Precise Control in\u00a0Diffusion Models"],"prefix":"10.1007","author":[{"given":"Rohit","family":"Gandikota","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Joanna","family":"Materzy\u0144ska","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Tingrui","family":"Zhou","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Antonio","family":"Torralba","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"David","family":"Bau","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2024,11,10]]},"reference":[{"key":"10_CR1","doi-asserted-by":"crossref","unstructured":"Brooks, T., Holynski, A., Efros, A.A.: InstructPix2Pix: learning to follow image editing instructions. arXiv preprint arXiv:2211.09800 (2022)","DOI":"10.1109\/CVPR52729.2023.01764"},{"key":"10_CR2","unstructured":"Burkett, J.: Ostris\/ai-toolkit: Various AI scripts. mostly stable diffusion stuff. (2023). https:\/\/github.com\/ostris\/ai-toolkit"},{"key":"10_CR3","doi-asserted-by":"crossref","unstructured":"Ceylan, D., Huang, C.H., Mitra, N.J.: Pix2video: video editing using image diffusion. In: International Conference on Computer Vision (ICCV) (2023)","DOI":"10.1109\/ICCV51070.2023.02121"},{"key":"10_CR4","unstructured":"Dhariwal, P., Nichol, A.: Diffusion models beat GANs on image synthesis. In: Advances in Neural Information Processing Systems, vol. 34, pp. 8780\u20138794 (2021)"},{"issue":"496","key":"10_CR5","doi-asserted-by":"publisher","first-page":"1602","DOI":"10.1198\/jasa.2011.tm11181","volume":"106","author":"B Efron","year":"2011","unstructured":"Efron, B.: Tweedie\u2019s formula and selection bias. J. Am. Stat. Assoc. 106(496), 1602\u20131614 (2011)","journal-title":"J. Am. Stat. Assoc."},{"key":"10_CR6","unstructured":"Gal, R., et al.: An image is worth one word: personalizing text-to-image generation using textual inversion. arXiv preprint arXiv:2208.01618 (2022)"},{"key":"10_CR7","doi-asserted-by":"crossref","unstructured":"Gandikota, R., Materzy\u0144ska, J., Fiotto-Kaufman, J., Bau, D.: Erasing concepts from diffusion models. In: Proceedings of the 2023 IEEE International Conference on Computer Vision (2023)","DOI":"10.1109\/ICCV51070.2023.00230"},{"key":"10_CR8","doi-asserted-by":"crossref","unstructured":"Gandikota, R., Orgad, H., Belinkov, Y., Materzy\u0144ska, J., Bau, D.: Unified concept editing in diffusion models. In: IEEE\/CVF Winter Conference on Applications of Computer Vision (2024)","DOI":"10.1109\/WACV57701.2024.00503"},{"key":"10_CR9","unstructured":"Google: Imagen, unprecedented photorealism x deep level of language understanding (2022). https:\/\/imagen.research.google\/"},{"key":"10_CR10","doi-asserted-by":"crossref","unstructured":"Haas, R., Huberman-Spiegelglas, I., Mulayoff, R., Michaeli, T.: Discovering interpretable directions in the semantic latent space of diffusion models. arXiv preprint arXiv:2303.11073 (2023)","DOI":"10.1109\/FG59268.2024.10581912"},{"key":"10_CR11","unstructured":"H\u00e4rk\u00f6nen, E., Hertzmann, A., Lehtinen, J., Paris, S.: GANSpace: discovering interpretable GAN controls. In: Advances in Neural Information Processing Systems, vol. 33, pp. 9841\u20139850 (2020)"},{"key":"10_CR12","unstructured":"Heng, A., Soh, H.: Selective amnesia: a continual learning approach to forgetting in deep generative models. arXiv preprint arXiv:2305.10120 (2023)"},{"key":"10_CR13","unstructured":"Hertz, A., Mokady, R., Tenenbaum, J., Aberman, K., Pritch, Y., Cohen-Or, D.: Prompt-to-prompt image editing with cross attention control. arXiv preprint arXiv:2208.01626 (2022)"},{"key":"10_CR14","unstructured":"Ho, J., Jain, A., Abbeel, P.: Denoising diffusion probabilistic models. In: Advances in Neural Information Processing Systems, vol. 33, pp. 6840\u20136851 (2020)"},{"key":"10_CR15","unstructured":"Ho, J., Salimans, T.: Classifier-free diffusion guidance. arXiv preprint arXiv:2207.12598 (2022)"},{"key":"10_CR16","unstructured":"Hu, E.J., et al.: Lora: low-rank adaptation of large language models. arXiv preprint arXiv:2106.09685 (2021)"},{"key":"10_CR17","unstructured":"Inui, N.: SD\/SDXL tricks beneath the papers and codes (2023). https:\/\/normxu.github.io\/sd-tricks\/#slider-lora"},{"key":"10_CR18","unstructured":"Jahanian, A., Chai, L., Isola, P.: On the \u201csteerability\u201d of generative adversarial networks. arXiv preprint arXiv:1907.07171 (2019)"},{"key":"10_CR19","unstructured":"Betker, J., et al.: Improving image generation with better captions. OpenAI Reports (2023). https:\/\/cdn.openai.com\/papers\/dall-e-3.pdf"},{"key":"10_CR20","unstructured":"Karras, T., et al.: Alias-free generative adversarial networks. In: Advances in Neural Information Processing Systems, vol. 34, pp. 852\u2013863 (2021)"},{"key":"10_CR21","doi-asserted-by":"crossref","unstructured":"Karras, T., Laine, S., Aila, T.: A style-based generator architecture for generative adversarial networks. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 4401\u20134410 (2019)","DOI":"10.1109\/CVPR.2019.00453"},{"key":"10_CR22","doi-asserted-by":"crossref","unstructured":"Kawar, B., et al.: Imagic: text-based real image editing with diffusion models. arXiv preprint arXiv:2210.09276 (2022)","DOI":"10.1109\/CVPR52729.2023.00582"},{"key":"10_CR23","unstructured":"Kim, S., Jung, S., Kim, B., Choi, M., Shin, J., Lee, J.: Towards safe self-distillation of internet-scale text-to-image diffusion models. arXiv preprint arXiv:2307.05977 (2023)"},{"key":"10_CR24","doi-asserted-by":"crossref","unstructured":"Kumari, N., Zhang, B., Wang, S.Y., Shechtman, E., Zhang, R., Zhu, J.Y.: Ablating concepts in text-to-image diffusion models. In: International Conference on Computer Vision (ICCV) (2023)","DOI":"10.1109\/ICCV51070.2023.02074"},{"key":"10_CR25","doi-asserted-by":"crossref","unstructured":"Kumari, N., Zhang, B., Zhang, R., Shechtman, E., Zhu, J.Y.: Multi-concept customization of text-to-image diffusion. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2023)","DOI":"10.1109\/CVPR52729.2023.00192"},{"key":"10_CR26","unstructured":"Kwon, M., Jeong, J., Uh, Y.: Diffusion models already have a semantic latent space. arXiv preprint arXiv:2210.10960 (2022)"},{"key":"10_CR27","unstructured":"Kynk\u00e4\u00e4nniemi, T., Karras, T., Aittala, M., Aila, T., Lehtinen, J.: The role of imagenet classes in Fr\u00e9chet inception distance. arXiv preprint arXiv:2203.06026 (2022)"},{"key":"10_CR28","doi-asserted-by":"crossref","unstructured":"Liu, N., Li, S., Du, Y., Torralba, A., Tenenbaum, J.B.: Compositional visual generation with composable diffusion models. arXiv preprint arXiv:2206.01714 (2022)","DOI":"10.1007\/978-3-031-19790-1_26"},{"key":"10_CR29","unstructured":"Luo, C.: Understanding diffusion models: a unified perspective. arXiv preprint arXiv:2208.11970 (2022)"},{"key":"10_CR30","unstructured":"Meng, C., Song, Y., Song, J., Wu, J., Zhu, J.Y., Ermon, S.: SDEdit: image synthesis and editing with stochastic differential equations. arXiv preprint arXiv:2108.01073 (2021)"},{"key":"10_CR31","unstructured":"Mothrider: Can an AI draw hands? (2022). https:\/\/www.reddit.com\/r\/StableDiffusion\/comments\/ym37xi\/can_an_ai_draw_hands\/"},{"key":"10_CR32","doi-asserted-by":"crossref","unstructured":"Orgad, H., Kawar, B., Belinkov, Y.: Editing implicit assumptions in text-to-image diffusion models. In: Proceedings of the 2023 IEEE International Conference on Computer Vision (2023)","DOI":"10.1109\/ICCV51070.2023.00649"},{"key":"10_CR33","unstructured":"Park, Y.H., Kwon, M., Choi, J., Jo, J., Uh, Y.: Understanding the latent space of diffusion models through the lens of Riemannian geometry. arXiv preprint arXiv:2307.12868 (2023)"},{"key":"10_CR34","unstructured":"Park, Y.H., Kwon, M., Jo, J., Uh, Y.: Unsupervised discovery of semantic latent directions in diffusion models. arXiv preprint arXiv:2302.12469 (2023)"},{"key":"10_CR35","doi-asserted-by":"crossref","unstructured":"Parmar, G., Kumar\u00a0Singh, K., Zhang, R., Li, Y., Lu, J., Zhu, J.Y.: Zero-shot image-to-image translation. In: ACM SIGGRAPH 2023 Conference Proceedings, pp. 1\u201311 (2023)","DOI":"10.1145\/3588432.3591513"},{"key":"10_CR36","unstructured":"Podell, D., et al.: SDXL: improving latent diffusion models for high-resolution image synthesis. arXiv preprint arXiv:2307.01952 (2023)"},{"key":"10_CR37","doi-asserted-by":"crossref","unstructured":"Rombach, R., Blattmann, A., Lorenz, D., Esser, P., Ommer, B.: High-resolution image synthesis with latent diffusion models. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR) (2022). https:\/\/github.com\/CompVis\/latent-diffusion","DOI":"10.1109\/CVPR52688.2022.01042"},{"key":"10_CR38","doi-asserted-by":"crossref","unstructured":"Ruiz, N., Li, Y., Jampani, V., Pritch, Y., Rubinstein, M., Aberman, K.: DreamBooth: fine tuning text-to-image diffusion models for subject-driven generation. arXiv preprint arXiv:2208.12242 (2022)","DOI":"10.1109\/CVPR52729.2023.02155"},{"key":"10_CR39","unstructured":"Ryu, S.: Cloneofsimo\/lora: using low-rank adaptation to quickly fine-tune diffusion models (2023). https:\/\/github.com\/cloneofsimo\/lora"},{"key":"10_CR40","doi-asserted-by":"crossref","unstructured":"Schramowski, P., Brack, M., Deiseroth, B., Kersting, K.: Safe latent diffusion: mitigating inappropriate degeneration in diffusion models. arXiv preprint arXiv:2211.05105 (2022)","DOI":"10.1109\/CVPR52729.2023.02157"},{"key":"10_CR41","doi-asserted-by":"crossref","unstructured":"Schroff, F., Kalenichenko, D., Philbin, J.: FaceNet: a unified embedding for face recognition and clustering. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 815\u2013823 (2015)","DOI":"10.1109\/CVPR.2015.7298682"},{"key":"10_CR42","doi-asserted-by":"crossref","unstructured":"Shen, Y., Gu, J., Tang, X., Zhou, B.: Interpreting the latent space of GANs for semantic face editing. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2020)","DOI":"10.1109\/CVPR42600.2020.00926"},{"key":"10_CR43","unstructured":"Staffell: The sheer number of options and sliders using stable diffusion is overwhelming (2023). https:\/\/www.reddit.com\/r\/StableDiffusion\/comments\/11am5zi\/the_sheer_number_of_options_and_sliders_using\/"},{"key":"10_CR44","doi-asserted-by":"crossref","unstructured":"Wu, Q., et al.: Uncovering the disentanglement capability in text-to-image diffusion models. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 1900\u20131910 (2023)","DOI":"10.1109\/CVPR52729.2023.00189"},{"key":"10_CR45","doi-asserted-by":"crossref","unstructured":"Wu, Z., Lischinski, D., Shechtman, E.: Stylespace analysis: disentangled controls for StyleGAN image generation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 12863\u201312872 (2021)","DOI":"10.1109\/CVPR46437.2021.01267"},{"key":"10_CR46","doi-asserted-by":"crossref","unstructured":"Zhang, E., Wang, K., Xu, X., Wang, Z., Shi, H.: Forget-me-not: learning to forget in text-to-image diffusion models. arXiv preprint arXiv:2303.17591 (2023)","DOI":"10.1109\/CVPRW63382.2024.00182"},{"key":"10_CR47","unstructured":"Zhou, T.: Github - p1atdev\/leco: low-rank adaptation for erasing concepts from diffusion models (2023). https:\/\/github.com\/p1atdev\/LECO"},{"key":"10_CR48","unstructured":"Zllrunning: Using modified bisenet for face parsing in PyTorch (2019). https:\/\/github.com\/zllrunning\/face-parsing.PyTorch"},{"key":"10_CR49","unstructured":"Zou, A., et\u00a0al.: Representation engineering: a top-down approach to AI transparency. arXiv preprint arXiv:2310.01405 (2023)"}],"container-title":["Lecture Notes in Computer Science","Computer Vision \u2013 ECCV 2024"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/978-3-031-73661-2_10","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,11,9]],"date-time":"2024-11-09T12:04:22Z","timestamp":1731153862000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/978-3-031-73661-2_10"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,11,10]]},"ISBN":["9783031736605","9783031736612"],"references-count":49,"URL":"https:\/\/doi.org\/10.1007\/978-3-031-73661-2_10","relation":{},"ISSN":["0302-9743","1611-3349"],"issn-type":[{"value":"0302-9743","type":"print"},{"value":"1611-3349","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,11,10]]},"assertion":[{"value":"10 November 2024","order":1,"name":"first_online","label":"First Online","group":{"name":"ChapterHistory","label":"Chapter History"}},{"value":"ECCV","order":1,"name":"conference_acronym","label":"Conference Acronym","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"European Conference on Computer Vision","order":2,"name":"conference_name","label":"Conference Name","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Milan","order":3,"name":"conference_city","label":"Conference City","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Italy","order":4,"name":"conference_country","label":"Conference Country","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"2024","order":5,"name":"conference_year","label":"Conference Year","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"29 September 2024","order":7,"name":"conference_start_date","label":"Conference Start Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"4 October 2024","order":8,"name":"conference_end_date","label":"Conference End Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"18","order":9,"name":"conference_number","label":"Conference Number","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"eccv2024","order":10,"name":"conference_id","label":"Conference ID","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"https:\/\/eccv2024.ecva.net\/","order":11,"name":"conference_url","label":"Conference URL","group":{"name":"ConferenceInfo","label":"Conference Information"}}]}}