{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,21]],"date-time":"2026-07-21T23:19:35Z","timestamp":1784675975178,"version":"3.55.0"},"publisher-location":"New York, NY, USA","reference-count":81,"publisher":"ACM","license":[{"start":{"date-parts":[[2026,6,25]],"date-time":"2026-06-25T00:00:00Z","timestamp":1782345600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by\/4.0\/legalcode"}],"funder":[{"name":"Carnegie Mellon University Block Center for Technology and Society","award":["62020.1.5007718"],"award-info":[{"award-number":["62020.1.5007718"]}]},{"name":"National Institute of Standards and Technology","award":["60NANB24D231"],"award-info":[{"award-number":["60NANB24D231"]}]}],"content-domain":{"domain":["dl.acm.org"],"crossmark-restriction":true},"short-container-title":[],"published-print":{"date-parts":[[2026,6,25]]},"DOI":"10.1145\/3805689.3806723","type":"proceedings-article","created":{"date-parts":[[2026,6,23]],"date-time":"2026-06-23T16:20:39Z","timestamp":1782231639000},"page":"816-839","update-policy":"https:\/\/doi.org\/10.1145\/crossmark-policy","source":"Crossref","is-referenced-by-count":1,"title":["Beyond the Single Turn: Reframing Refusals as Dynamic Experiences Embedded in the Context of Mental Health Support Interactions with LLMs"],"prefix":"10.1145","author":[{"ORCID":"https:\/\/orcid.org\/0009-0003-2605-9000","authenticated-orcid":false,"given":"Ningjing","family":"Tang","sequence":"first","affiliation":[{"name":"Carnegie Mellon University, Pittsburgh, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0005-6407-6981","authenticated-orcid":false,"given":"Alice","family":"Qian","sequence":"additional","affiliation":[{"name":"Carnegie Mellon University, Pittsburgh, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-5296-5440","authenticated-orcid":false,"given":"Qiaosi","family":"Wang","sequence":"additional","affiliation":[{"name":"Carnegie Mellon University, Pittsburgh, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-8429-4883","authenticated-orcid":false,"given":"Esther","family":"Howe","sequence":"additional","affiliation":[{"name":"Department of Psychiatry and Behavioral Sciences, University of Washington School of Medicine, Seattle, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0008-6623-8181","authenticated-orcid":false,"given":"Blake","family":"Bullwinkel","sequence":"additional","affiliation":[{"name":"Microsoft, Redmond, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-2865-2583","authenticated-orcid":false,"given":"Paola","family":"Pedrelli","sequence":"additional","affiliation":[{"name":"Harvard Medical School, Boston, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-7646-5563","authenticated-orcid":false,"given":"Jina","family":"Suh","sequence":"additional","affiliation":[{"name":"University of Washington, Seattle, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-3710-4076","authenticated-orcid":false,"given":"Hoda","family":"Heidari","sequence":"additional","affiliation":[{"name":"Carnegie Mellon University, Pittsburgh, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-5364-3718","authenticated-orcid":false,"given":"Hong","family":"Shen","sequence":"additional","affiliation":[{"name":"Carnegie Mellon University, Pittsburgh, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"320","published-online":{"date-parts":[[2026,6,25]]},"reference":[{"key":"e_1_3_2_1_1_1","volume-title":"Seeking late night life lines: Experiences of conversational AI use in mental health crisis. arXiv [cs.HC] (Dec","author":"Ajmani Leah Hope","year":"2025","unstructured":"Leah Hope Ajmani, Arka Ghosh, Benjamin Kaveladze, Eugenia Kim, Keertana Namuduri, Theresa Nguyen, Ebele Okoli, Jessica Schleider, Denae Ford, and Jina Suh. 2025. Seeking late night life lines: Experiences of conversational AI use in mental health crisis. arXiv [cs.HC] (Dec. 2025)."},{"key":"e_1_3_2_1_2_1","volume-title":"Associated Press and Leah Willingham","author":"Selsky Andrew","year":"2022","unstructured":"Andrew Selsky, Associated Press and Leah Willingham, Associated Press. 2022. How some encounters between police and people with mental illness can turn tragic. https:\/\/www.pbs.org\/newshour\/health\/how-some-encounters-between-police-and-people-with-mental-illness-can-turn-tragic. Accessed: 2026-1-8."},{"key":"e_1_3_2_1_3_1","unstructured":"Yuntao Bai Saurav Kadavath Sandipan Kundu Amanda Askell Jackson Kernion Andy Jones Anna Chen Anna Goldie Azalia Mirhoseini Cameron McKinnon Carol Chen Catherine Olsson Christopher Olah Danny Hernandez Dawn Drain Deep Ganguli Dustin Li Eli Tran-Johnson Ethan Perez Jamie Kerr Jared Mueller Jeffrey Ladish Joshua Landau Kamal Ndousse Kamile Lukosuite Liane Lovitt Michael Sellitto Nelson Elhage Nicholas Schiefer Noemi Mercado Nova DasSarma Robert Lasenby Robin Larson Sam Ringer Scott Johnston Shauna Kravec Sheer El Showk Stanislav Fort Tamera Lanham Timothy Telleen-Lawton Tom Conerly Tom Henighan Tristan Hume Samuel R Bowman Zac Hatfield-Dodds Ben Mann Dario Amodei Nicholas Joseph Sam McCandlish Tom Brown and Jared Kaplan. 2022. Constitutional AI: Harmlessness from AI Feedback. arXiv [cs.CL] (Dec. 2022)."},{"key":"e_1_3_2_1_4_1","volume-title":"The art of saying no: Contextual noncompliance in language models. arXiv [cs.CL] (July","author":"Brahman Faeze","year":"2024","unstructured":"Faeze Brahman, Sachin Kumar, Vidhisha Balachandran, Pradeep Dasigi, Valentina Pyatkin, Abhilasha Ravichander, Sarah Wiegreffe, Nouha Dziri, Khyathi Chandu, Jack Hessel, Yulia Tsvetkov, Noah A Smith, Yejin Choi, and Hannaneh Hajishirzi. 2024. The art of saying no: Contextual noncompliance in language models. arXiv [cs.CL] (July 2024)."},{"key":"e_1_3_2_1_5_1","doi-asserted-by":"publisher","DOI":"10.1191\/1478088706qp063oa"},{"key":"e_1_3_2_1_6_1","volume-title":"Thematic analysis","author":"Braun Virginia","unstructured":"Virginia Braun and Victoria Clarke. 2012. Thematic analysis. American Psychological Association."},{"key":"e_1_3_2_1_7_1","volume-title":"Reflecting on reflexive thematic analysis. Qualitative research in sport, exercise and health 11, 4","author":"Braun Virginia","year":"2019","unstructured":"Virginia Braun and Victoria Clarke. 2019. Reflecting on reflexive thematic analysis. Qualitative research in sport, exercise and health 11, 4 (2019), 589\u2013597."},{"key":"e_1_3_2_1_8_1","volume-title":"Senate Bill 243: Companion Chatbots. Approved by Governor","author":"Legislature California State","year":"2025","unstructured":"California State Legislature. 2025. Senate Bill 243: Companion Chatbots. Approved by Governor October 13, 2025. https:\/\/leginfo.legislature.ca.gov\/faces\/billNavClient.xhtml?bill_id=202520260SB243 Chapter 677, Statutes of 2025."},{"key":"e_1_3_2_1_9_1","doi-asserted-by":"publisher","DOI":"10.1145\/3715275.3732063"},{"key":"e_1_3_2_1_10_1","volume-title":"Mahsa Ershadi, Gonzalo Ramos, Javier Hernandez, Ananya Bhattacharjee, Shahed Warreth, and Jina Suh.","author":"Chandra Mohit","year":"2024","unstructured":"Mohit Chandra, Suchismita Naik, Denae Ford, Ebele Okoli, Munmun De Choudhury, Mahsa Ershadi, Gonzalo Ramos, Javier Hernandez, Ananya Bhattacharjee, Shahed Warreth, and Jina Suh. 2024. From lived experience to insight: Unpacking the psychological risks of using AI conversational agents. arXiv [cs.HC] (Dec. 2024)."},{"key":"e_1_3_2_1_11_1","volume-title":"JailbreakBench: An open robustness benchmark for jailbreaking large language models. arXiv [cs.CR] (March","author":"Chao Patrick","year":"2024","unstructured":"Patrick Chao, Edoardo Debenedetti, Alexander Robey, Maksym Andriushchenko, Francesco Croce, Vikash Sehwag, Edgar Dobriban, Nicolas Flammarion, George J Pappas, Florian Tramer, Hamed Hassani, and Eric Wong. 2024. JailbreakBench: An open robustness benchmark for jailbreaking large language models. arXiv [cs.CR] (March 2024)."},{"key":"e_1_3_2_1_12_1","volume-title":"Jackie Chi Kit Cheung, and Golnoosh Farnadi","author":"Chehbouni Khaoula","year":"2025","unstructured":"Khaoula Chehbouni, Mohammed Haddou, Jackie Chi Kit Cheung, and Golnoosh Farnadi. 2025. Neither valid nor reliable? Investigating the use of LLMs as judges. arXiv [cs.CL] (Aug. 2025)."},{"key":"e_1_3_2_1_13_1","doi-asserted-by":"crossref","first-page":"2","DOI":"10.1007\/s10676-025-09837-2","article-title":"Helpful, harmless, honest? Sociotechnical limits of AI alignment and safety through Reinforcement Learning from Human Feedback","volume":"27","author":"Lindstr\u00f6m Adam Dahlgren","year":"2025","unstructured":"Adam Dahlgren Lindstr\u00f6m, Leila Methnani, Lea Krause, Petter Ericson, \u00cd\u00f1igo Mart\u00ednez de Rituerto de Troya, Dimitri Coelho Mollo, and Roel Dobbe. 2025. Helpful, harmless, honest? Sociotechnical limits of AI alignment and safety through Reinforcement Learning from Human Feedback. Ethics Inf. Technol. 27, 2 (June 2025), 28.","journal-title":"Ethics Inf. Technol."},{"key":"e_1_3_2_1_14_1","doi-asserted-by":"publisher","DOI":"10.1609\/icwsm.v8i1.14526"},{"key":"e_1_3_2_1_15_1","doi-asserted-by":"publisher","DOI":"10.1186\/1745-6215-15-267"},{"key":"e_1_3_2_1_16_1","doi-asserted-by":"publisher","DOI":"10.1016\/j.psc.2024.04.003"},{"key":"e_1_3_2_1_17_1","doi-asserted-by":"publisher","DOI":"10.1111\/sltb.12128"},{"key":"e_1_3_2_1_18_1","unstructured":"Foundational Contributors Ahmed El-Kishky Daniel Selsam Francis Song Giambattista Parascandolo Hongyu Ren Hunter Lightman Hyung Won Ilge Akkaya I Sutskever Jason Wei Jonathan Gordon K Cobbe Kevin Yu Lukasz Kondraciuk Max Schwarzer Mostafa Rohaninejad Noam Brown Shengjia Zhao Trapit Bansal Vineet Kosaraju Wenda Zhou Leadership J Pachocki Jerry Tworek L Fedus \u0141ukasz Kaiser Mark Chen Szymon Sidor Wojciech Zaremba Alex Karpenko Alexander Wei Allison Tam Ananya Kumar Andre Saraiva A Kondrich Andrey Mishchenko Ashvin Nair B Ghorbani Brandon McKinzie Chak Bry-Don Eastman Ming Li Chris Koch Dan Roberts David Dohan David M\u00e9ly Dimitris Tsipras Enoch Cheung Eric Wallace Hadi Salman Haim-Ing Bao Hessam Bagher-inezhad Ilya Kostrikov Jiacheng Feng John Rizzo Karina Nguyen Kevin Lu Kevin Stone Lorenz Kuhn Mason Meyer Mikhail Pavlov Nat McAleese Oleg Boiko O Murk Peter Zhokhov Randall Lin Raz Gaon Rhythm Garg Roshan James Rui Shu Scott McKinney Shibani Santurkar S Balaji Taylor Gordon Thomas Dimson Weiyi Zheng Aaron Jaech Adam Lerer Aiden Low Alex Carney Alexander Neitz Alexander Prokofiev Benjamin Sokolowsky Boaz Barak Borys Minaiev Botao Hao Bowen Baker Brandon Houghton Camillo Lugaresi Chelsea Voss Chen Shen Chris Orsinger Daniel Kappler Daniel Levy Doug Li Eben Freeman Edmund Wong Fan Wang F Such Foivos Tsimpourlas Geoff Salmon Gildas Chabot Guillaume Leclerc Hart Andrin Ian O'Connell Ignasi Ian Osband Clavera Gilaberte Jean Harb Jiahui Yu Jiayi Weng Joe Palermo John Hallman Jonathan Ward Julie Wang Kai Chen Katy Shi Keren Gu-Lemberg Kevin Liu Leo Liu Linden Li Luke Metz Maja Trebacz Manas R Joglekar Marko Tintor Melody Y Guan Mengyuan Yan Mia Glaese Michael Malek Michelle Fradin Mo Bavarian N Tezak Ofir Nachum Paul Ashbourne Pavel Izmailov Raphael Gontijo Lopes Reah Miyara R Leike Robin Brown Ryan Cheu Ryan Greene Saachi Jain Scottie Yan Shengli Hu Shuyuan Zhang Siyuan Fu Spencer Papay Suvansh Sanjeev Tao Wang Ted Sanders Tejal A Patwardhan Thibault Sottiaux Tianhao Zheng T Garipov Valerie Qi Vitchyr H Pong Vlad Fomenko Yinghai Lu Yining Chen Yu Bai Yuchen He Yuchen Zhang Zheng Shao Zhuohan Li Lauren Yang Mianna Chen Aidan Clark Jieqi Yu Kai Xiao Sam Toizer Sandhini Agarwal Safety Research Andrea Vallone Chong Zhang I Kivlichan Meghan Shah S Toyer Shraman Ray Chaudhuri Stephanie Lin Adam Richardson Andrew Duberstein C D Bourcy Dragos Oprica Florencia Leoni Made-Laine Boyd Matt Jones Matt Kaufer M Yatbaz Mengyuan Xu Mike McClay Mingxuan Wang Trevor Creech Vinnie Monaco Erik Ritter Evan Mays Joel Parish Jonathan Uesato Leon Maksin Michele Wang Miles Wang Neil Chowdhury Olivia Watkins Patrick Chao Rachel Dias Samuel Miserendino Red Teaming Lama Ahmad Michael Lampe Troy Peterson and Joost Huizinga. 2024. OpenAI o1 System Card. ArXiv abs\/2412.16720 (Dec. 2024)."},{"key":"e_1_3_2_1_19_1","series-title":"June 2025","volume-title":"I cannot write this because it violates our content policy\u201d: Understanding content moderation policies and user experiences in generative AI products. arXiv [cs.HC]","author":"Gao Lan","unstructured":"Lan Gao, Oscar Chen, Rachel Lee, Nick Feamster, Chenhao Tan, and Marshini Chetty. 2025. \u201cI cannot write this because it violates our content policy\u201d: Understanding content moderation policies and user experiences in generative AI products. arXiv [cs.HC] (June 2025)."},{"key":"e_1_3_2_1_20_1","unstructured":"Clifford Geertz. 1973. The impact of the concept of culture on the concept of man."},{"key":"e_1_3_2_1_21_1","doi-asserted-by":"publisher","DOI":"10.2196\/jmir.7082"},{"key":"e_1_3_2_1_22_1","volume-title":"Sam Toyer, Johannes Heidecke, Alex Beutel, and Amelia Glaese.","author":"Guan Melody Y","year":"2024","unstructured":"Melody Y Guan, Manas Joglekar, Eric Wallace, Saachi Jain, Boaz Barak, Alec Helyar, Rachel Dias, Andrea Vallone, Hongyu Ren, Jason Wei, Hyung Won Chung, Sam Toyer, Johannes Heidecke, Alex Beutel, and Amelia Glaese. 2024. Deliberative Alignment: Reasoning enables safer language models. arXiv [cs.CL] (Dec. 2024)."},{"key":"e_1_3_2_1_23_1","volume-title":"Sam Toyer, Johannes Heidecke, Alex Beutel, and Amelia Glaese.","author":"Guan Melody Y","year":"2024","unstructured":"Melody Y Guan, Manas Joglekar, Eric Wallace, Saachi Jain, Boaz Barak, Alec Helyar, Rachel Dias, Andrea Vallone, Hongyu Ren, Jason Wei, Hyung Won Chung, Sam Toyer, Johannes Heidecke, Alex Beutel, and Amelia Glaese. 2024. Deliberative Alignment: Reasoning enables safer language models. arXiv [cs.CL] (Dec. 2024)."},{"key":"e_1_3_2_1_24_1","volume-title":"Nathan Lambert, Yejin Choi, and Nouha Dziri.","author":"Han Seungju","year":"2024","unstructured":"Seungju Han, Kavel Rao, Allyson Ettinger, Liwei Jiang, Bill Yuchen Lin, Nathan Lambert, Yejin Choi, and Nouha Dziri. 2024. WildGuard: Open one-stop moderation tools for safety risks, jailbreaks, and refusals of LLMs. arXiv [cs.CL] (June 2024)."},{"key":"e_1_3_2_1_25_1","volume-title":"Proc. ACM Hum. Comput. Interact. 6, CSCW2 (Nov.","author":"Haque Md Romael","year":"2022","unstructured":"Md Romael Haque and Sabirat Rubya. 2022. \u201cfor an app supposed to make its users feel better, it sure is a joke\u201d - an analysis of user reviews of mobile mental health applications. Proc. ACM Hum. Comput. Interact. 6, CSCW2 (Nov. 2022), 1\u201329."},{"key":"e_1_3_2_1_26_1","doi-asserted-by":"publisher","DOI":"10.1145\/3359318"},{"key":"e_1_3_2_1_27_1","volume-title":"A Teen Was Suicidal. ChatGPT Was the Friend He Confided In. The New York Times (Aug","author":"Hill Kashmir","year":"2025","unstructured":"Kashmir Hill. 2025. A Teen Was Suicidal. ChatGPT Was the Friend He Confided In. The New York Times (Aug. 2025)."},{"key":"e_1_3_2_1_28_1","volume-title":"Proceedings of the AAAI\/ACM Conference on AI, Ethics, and Society","volume":"8","author":"Ibrahim Lujain","year":"2025","unstructured":"Lujain Ibrahim, Saffron Huang, Lama Ahmad, Umang Bhatt, and Markus Anderljung. 2025. Towards interactive evaluations for interaction harms in human-AI systems. In Proceedings of the AAAI\/ACM Conference on AI, Ethics, and Society, Vol. 8. 1302\u20131310."},{"key":"e_1_3_2_1_29_1","doi-asserted-by":"publisher","DOI":"10.1609\/aies.v8i2.36632"},{"key":"e_1_3_2_1_30_1","volume-title":"House Bill 1806: Wellness and Oversight for Psychological Resources Act. Signed into law","author":"Assembly Illinois General","year":"2025","unstructured":"Illinois General Assembly. 2025. House Bill 1806: Wellness and Oversight for Psychological Resources Act. Signed into law August 1, 2025. https:\/\/ilga.gov\/Legislation\/BillStatus?DocNum=1806&GAID=18&DocTypeID=HB&LegId=159219&SessionID=114 Public Act 104\u20130054."},{"key":"e_1_3_2_1_31_1","doi-asserted-by":"publisher","DOI":"10.1177\/1525822X05282260"},{"key":"e_1_3_2_1_32_1","doi-asserted-by":"publisher","DOI":"10.1177\/1468794109356737"},{"key":"e_1_3_2_1_33_1","volume-title":"Safiya Umoja Noble, and Benjamin Shestakofsky","author":"Joyce Kelly","year":"2021","unstructured":"Kelly Joyce, Laurel Smith-Doerr, Sharla Alegria, Susan Bell, Taylor Cruz, Steve G Hoffman, Safiya Umoja Noble, and Benjamin Shestakofsky. 2021. Toward a sociology of artificial intelligence: A call for research on inequalities and structural change. Socius 7 (Jan. 2021), 237802312199958."},{"key":"e_1_3_2_1_34_1","series-title":"April 2025","volume-title":"I've talked to ChatGPT about my issues last night.\u201d: Examining Mental Health Conversations with Large Language Models through Reddit Analysis. arXiv [cs.HC]","author":"Jung Kyuha","unstructured":"Kyuha Jung, Gyuho Lee, Yuanhui Huang, and Yunan Chen. 2025. \u201cI've talked to ChatGPT about my issues last night.\u201d: Examining Mental Health Conversations with Large Language Models through Reddit Analysis. arXiv [cs.HC] (April 2025)."},{"key":"e_1_3_2_1_35_1","doi-asserted-by":"publisher","DOI":"10.1145\/3635636.3664263"},{"key":"e_1_3_2_1_36_1","doi-asserted-by":"publisher","DOI":"10.1057\/s41599-025-04532-5"},{"key":"e_1_3_2_1_37_1","doi-asserted-by":"publisher","DOI":"10.1145\/3485874"},{"key":"e_1_3_2_1_38_1","doi-asserted-by":"publisher","DOI":"10.1126\/science.adi8982"},{"key":"e_1_3_2_1_39_1","doi-asserted-by":"publisher","DOI":"10.1016\/j.ijhcs.2025.103555"},{"key":"e_1_3_2_1_40_1","doi-asserted-by":"publisher","DOI":"10.1145\/3512899"},{"key":"e_1_3_2_1_41_1","volume-title":"HarmBench: A standardized evaluation framework for automated red teaming and robust refusal. arXiv [cs.LG] (Feb","author":"Mazeika Mantas","year":"2024","unstructured":"Mantas Mazeika, Long Phan, Xuwang Yin, Andy Zou, Zifan Wang, Norman Mu, Elham Sakhaee, Nathaniel Li, Steven Basart, Bo Li, David Forsyth, and Dan Hendrycks. 2024. HarmBench: A standardized evaluation framework for automated red teaming and robust refusal. arXiv [cs.LG] (Feb. 2024)."},{"key":"e_1_3_2_1_42_1","unstructured":"Miles McCain Ryn Linthicum Chloe Lubinski Alex Tamkin Saffron Huang Michael Stern Kunal Handa Esin Durmus Tyler Neylon Stuart Ritchie Kamya Jagadish Paruul Maheshwary Sarah Heck Alexandra Sanderford and Deep Ganguli. 2025. How People Use Claude for Support Advice and Companionship. https:\/\/www.anthropic.com\/news\/how-people-use-claude-for-support-advice-and-companionship."},{"key":"e_1_3_2_1_43_1","volume-title":"TikTok. In Proceedings of the 2023 CHI Conference on Human Factors in Computing Systems","volume":"16","author":"Milton Ashlee","year":"2023","unstructured":"Ashlee Milton, Leah Ajmani, Michael Ann DeVito, and Stevie Chancellor. 2023. \u201cI see me here\u201d: Mental health content, community, and algorithmic curation on TikTok. In Proceedings of the 2023 CHI Conference on Human Factors in Computing Systems, Vol. 16. ACM, New York, NY, USA, 1\u201317."},{"key":"e_1_3_2_1_44_1","doi-asserted-by":"publisher","DOI":"10.1145\/3715275.3732039"},{"key":"e_1_3_2_1_45_1","volume-title":"Towards a non-ideal methodological framework for Responsible ML. arXiv [cs.HC] (Jan","author":"Mothilal Ramaravind Kommiya","year":"2024","unstructured":"Ramaravind Kommiya Mothilal, Shion Guha, and Syed Ishtiaque Ahmed. 2024. Towards a non-ideal methodological framework for Responsible ML. arXiv [cs.HC] (Jan. 2024)."},{"key":"e_1_3_2_1_46_1","doi-asserted-by":"publisher","DOI":"10.1145\/3772318.3791285"},{"key":"e_1_3_2_1_47_1","unstructured":"Alondra Nelson. 2023. Thick Alignment."},{"key":"e_1_3_2_1_48_1","first-page":"2025","article-title":"An Act to amend the general business law, in relation to artificial intelligence companion models","volume":"6767","author":"State Assembly New York","year":"2025","unstructured":"New York State Assembly. 2025. An Act to amend the general business law, in relation to artificial intelligence companion models. Assembly Bill A6767, 2025-2026 Regular Sessions. https:\/\/www.nysenate.gov\/legislation\/bills\/2025\/A6767 Introduced by M. of A. Vanel; referred to the Committee on Consumer Affairs and Protection.","journal-title":"Assembly Bill"},{"key":"e_1_3_2_1_49_1","unstructured":"OpenAI. 2025. Strengthening ChatGPT's responses in sensitive conversations. https:\/\/openai.com\/index\/strengthening-chatgpt-responses-in-sensitive-conversations\/. Accessed: 2025-11-24."},{"key":"e_1_3_2_1_50_1","unstructured":"OpenAI Josh Achiam Steven Adler Sandhini Agarwal Lama Ahmad Ilge Akkaya Florencia Leoni Aleman Diogo Almeida Janko Altenschmidt Sam Altman Shyamal Anadkat Red Avila Igor Babuschkin Suchir Balaji Valerie Balcom Paul Baltescu Haiming Bao Mohammad Bavarian Jeff Belgum Irwan Bello Jake Berdine Gabriel Bernadett-Shapiro Christopher Berner Lenny Bogdonoff Oleg Boiko Madelaine Boyd Anna-Luisa Brakman Greg Brockman Tim Brooks Miles Brundage Kevin Button Trevor Cai Rosie Campbell Andrew Cann Brittany Carey Chelsea Carlson Rory Carmichael Brooke Chan Che Chang Fotis Chantzis Derek Chen Sully Chen Ruby Chen Jason Chen Mark Chen Ben Chess Chester Cho Casey Chu Hyung Won Chung Dave Cummings Jeremiah Currier Yunxing Dai Cory Decareaux Thomas Degry Noah Deutsch Damien Deville Arka Dhar David Dohan Steve Dowling Sheila Dunning Adrien Ecoffet Atty Eleti Tyna Eloundou David Farhi Liam Fedus Niko Felix Sim\u00f3n Posada Fishman Juston Forte Isabella Fulford Leo Gao Elie Georges Christian Gibson Vik Goel Tarun Gogineni Gabriel Goh Rapha Gontijo-Lopes Jonathan Gordon Morgan Grafstein Scott Gray Ryan Greene Joshua Gross Shixiang Shane Gu Yufei Guo Chris Hallacy Jesse Han Jeff Harris Yuchen He Mike Heaton Johannes Heidecke Chris Hesse Alan Hickey Wade Hickey Peter Hoeschele Brandon Houghton Kenny Hsu Shengli Hu Xin Hu Joost Huizinga Shantanu Jain Shawn Jain Joanne Jang Angela Jiang Roger Jiang Haozhun Jin Denny Jin Shino Jomoto Billie Jonn Heewoo Jun Tomer Kaftan \u0141ukasz Kaiser Ali Kamali Ingmar Kanitscheider Nitish Shirish Keskar Tabarak Khan Logan Kilpatrick Jong Wook Kim Christina Kim Yongjik Kim Jan Hendrik Kirchner Jamie Kiros Matt Knight Daniel Kokotajlo \u0141ukasz Kondraciuk Andrew Kondrich Aris Konstantinidis Kyle Kosic Gretchen Krueger Vishal Kuo Michael Lampe Ikai Lan Teddy Lee Jan Leike Jade Leung Daniel Levy Chak Ming Li Rachel Lim Molly Lin Stephanie Lin Mateusz Litwin Theresa Lopez Ryan Lowe Patricia Lue Anna Makanju Kim Malfacini Sam Manning Todor Markov Yaniv Markovski Bianca Martin Katie Mayer Andrew Mayne Bob McGrew Scott Mayer McKinney Christine McLeavey Paul McMillan Jake McNeil David Medina Aalok Mehta Jacob Menick Luke Metz Andrey Mishchenko Pamela Mishkin Vinnie Monaco Evan Morikawa Daniel Mossing Tong Mu Mira Murati Oleg Murk David M\u00e9ly Ashvin Nair Reiichiro Nakano Rajeev Nayak Arvind Neelakantan Richard Ngo Hyeonwoo Noh Long Ouyang Cullen O'Keefe Jakub Pachocki Alex Paino Joe Palermo Ashley Pantuliano Giambattista Parascandolo Joel Parish Emy Parparita Alex Passos Mikhail Pavlov Andrew Peng Adam Perelman Filipe de Avila Belbute Peres Michael Petrov Henrique Ponde de Oliveira Pinto Michael Pokorny Michelle Pokrass Vitchyr H Pong Tolly Powell Alethea Power Boris Power Elizabeth Proehl Raul Puri Alec Radford Jack Rae Aditya Ramesh Cameron Raymond Francis Real Kendra Rimbach Carl Ross Bob Rotsted Henri Roussez Nick Ryder Mario Saltarelli Ted Sanders Shibani Santurkar Girish Sastry Heather Schmidt David Schnurr John Schulman Daniel Selsam Kyla Sheppard Toki Sherbakov Jessica Shieh Sarah Shoker Pranav Shyam Szymon Sidor Eric Sigler Maddie Simens Jordan Sitkin Katarina Slama Ian Sohl Benjamin Sokolowsky Yang Song Natalie Staudacher Felipe Petroski Such Natalie Summers Ilya Sutskever Jie Tang Nikolas Tezak Madeleine B Thompson Phil Tillet Amin Tootoonchian Elizabeth Tseng Preston Tuggle Nick Turley Jerry Tworek Juan Felipe Cer\u00f3n Uribe Andrea Vallone Arun Vijayvergiya Chelsea Voss Carroll Wainwright Justin Jay Wang Alvin Wang Ben Wang Jonathan Ward Jason Wei C J Weinmann Akila Welihinda Peter Welinder Jiayi Weng Lilian Weng Matt Wiethoff Dave Willner Clemens Winter Samuel Wolrich Hannah Wong Lauren Workman Sherwin Wu Jeff Wu Michael Wu Kai Xiao Tao Xu Sarah Yoo Kevin Yu Qiming Yuan Wojciech Zaremba Rowan Zellers Chong Zhang Marvin Zhang Shengjia Zhao Tianhao Zheng Juntang Zhuang William Zhuk and Barret Zoph. 2023. GPT-4 Technical Report. arXiv [cs.CL] (March 2023)."},{"key":"e_1_3_2_1_51_1","doi-asserted-by":"publisher","DOI":"10.1007\/s10488-013-0528-y"},{"key":"e_1_3_2_1_52_1","doi-asserted-by":"publisher","DOI":"10.1145\/3411764.3445410"},{"key":"e_1_3_2_1_53_1","first-page":"1","article-title":"Exploring the ethical challenges of conversational AI in mental health care: Scoping review","volume":"12","author":"Meadi Mehrdad Rahsepar","year":"2025","unstructured":"Mehrdad Rahsepar Meadi, Tomas Sillekens, Suzanne Metselaar, Anton van Balkom, Justin Bernstein, and Neeltje Batelaan. 2025. Exploring the ethical challenges of conversational AI in mental health care: Scoping review. JMIR Ment. Health 12, 1 (Feb. 2025), e60432.","journal-title":"JMIR Ment. Health"},{"key":"e_1_3_2_1_54_1","volume-title":"Safetywashing: Do AI safety benchmarks actually measure safety progress? arXiv [cs.LG] (July","author":"Ren Richard","year":"2024","unstructured":"Richard Ren, Steven Basart, Adam Khoja, Alice Gatti, Long Phan, Xuwang Yin, Mantas Mazeika, Alexander Pan, Gabriel Mukobi, Ryan H Kim, Stephen Fitz, and Dan Hendrycks. 2024. Safetywashing: Do AI safety benchmarks actually measure safety progress? arXiv [cs.LG] (July 2024)."},{"key":"e_1_3_2_1_55_1","volume-title":"January Session, A.D.","author":"General Assembly Rhode Island","year":"2026","unstructured":"Rhode Island General Assembly. 2026. An Act Relating to Commercial Law\u2014General Regulatory Provisions\u2014Artificial Intelligence Companion Models. Senate Bill S2195, January Session, A.D. 2026. https:\/\/webserver.rilegislature.gov\/BillText\/BillText26\/SenateText26\/ S2195.pdf Introduced by Senators Urso, Gu, DiPalma, Paolino, Zurier, Murray, and Appollonio; referred to Senate Artificial Intelligence & Emerging Tech; LC003227."},{"key":"e_1_3_2_1_56_1","doi-asserted-by":"publisher","DOI":"10.1145\/3544548.3581559"},{"key":"e_1_3_2_1_57_1","doi-asserted-by":"publisher","DOI":"10.1145\/3287560.3287598"},{"key":"e_1_3_2_1_58_1","volume-title":"How RLHF Amplifies Sycophancy. arXiv [cs.AI] (Feb","author":"Shapira Itai","year":"2026","unstructured":"Itai Shapira, Gerdus Benade, and Ariel D Procaccia. 2026. How RLHF Amplifies Sycophancy. arXiv [cs.AI] (Feb. 2026)."},{"key":"e_1_3_2_1_59_1","doi-asserted-by":"publisher","DOI":"10.1145\/3600211.3604673"},{"key":"e_1_3_2_1_60_1","volume-title":"A Family's Call To 911 Turns Tragic. NPR (Oct.","author":"Sholtis Brett","year":"2020","unstructured":"Brett Sholtis. 2020. During A Mental Health Crisis, A Family's Call To 911 Turns Tragic. NPR (Oct. 2020)."},{"key":"e_1_3_2_1_61_1","doi-asserted-by":"publisher","DOI":"10.1038\/s44184-024-00097-4"},{"key":"e_1_3_2_1_62_1","doi-asserted-by":"publisher","DOI":"10.1145\/3613904.3642624"},{"key":"e_1_3_2_1_63_1","volume-title":"The typing cure: Experiences with Large Language Model chatbots for mental health support. arXiv [cs.HC] (Jan","author":"Song Inhwa","year":"2024","unstructured":"Inhwa Song, Sachin R Pendse, Neha Kumar, and Munmun De Choudhury. 2024. The typing cure: Experiences with Large Language Model chatbots for mental health support. arXiv [cs.HC] (Jan. 2024)."},{"key":"e_1_3_2_1_64_1","doi-asserted-by":"crossref","unstructured":"Alexandra Souly Qingyuan Lu Dillon Bowen Tu Trinh Elvis Hsieh Sana Pandey Pieter Abbeel Justin Svegliato Scott Emmons Olivia Watkins and Sam Toyer. 2024. A StrongREJECT for Empty Jailbreaks. arXiv:2402.10260 [cs.LG] https:\/\/arxiv.org\/abs\/2402.10260","DOI":"10.52202\/079017-3984"},{"key":"e_1_3_2_1_65_1","doi-asserted-by":"publisher","DOI":"10.1145\/3772318.3790910"},{"key":"e_1_3_2_1_66_1","doi-asserted-by":"publisher","DOI":"10.1145\/3613904.3642509"},{"key":"e_1_3_2_1_67_1","volume-title":"Regulating AI in mental health: Ethics of care perspective. JMIR Ment. Health 11 (Sept","author":"Tavory Tamar","year":"2024","unstructured":"Tamar Tavory. 2024. Regulating AI in mental health: Ethics of care perspective. JMIR Ment. Health 11 (Sept. 2024), e58493."},{"key":"e_1_3_2_1_68_1","volume-title":"Rebecca Qian, Anand Kannappan, Scott A Hale, and Paul R\u00f6ttger.","author":"Vidgen Bertie","year":"2023","unstructured":"Bertie Vidgen, Nino Scherrer, Hannah Rose Kirk, Rebecca Qian, Anand Kannappan, Scott A Hale, and Paul R\u00f6ttger. 2023. SimpleSafetyTests: A test suite for identifying critical safety risks in large language models. arXiv [cs.CL] (Nov. 2023)."},{"key":"e_1_3_2_1_69_1","unstructured":"Hanna Wallach Meera Desai A Feder Cooper Angelina Wang Chad Atalla Solon Barocas Su Lin Blodgett Alexandra Chouldechova Emily Corvi P Alex Dow Jean Garcia-Gathright Alexandra Olteanu Nicholas Pangakis Stefanie Reed Emily Sheng Dan Vann Jennifer Wortman Vaughan Matthew Vogel Hannah Washington and Abigail Z Jacobs. 2025. Position: Evaluating generative AI systems is a social science measurement challenge. arXiv [cs.CY] (Feb. 2025)."},{"key":"e_1_3_2_1_70_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/2024.findings-eacl.61"},{"key":"e_1_3_2_1_71_1","volume-title":"Hanna Wallach, Margaret Mitchell, Angelina Wang, Olawale Salaudeen, Rishi Bommasani, Deep Ganguli, Sanmi Koyejo, and William Isaac.","author":"Weidinger Laura","year":"2025","unstructured":"Laura Weidinger, Inioluwa Deborah Raji, Hanna Wallach, Margaret Mitchell, Angelina Wang, Olawale Salaudeen, Rishi Bommasani, Deep Ganguli, Sanmi Koyejo, and William Isaac. 2025. Toward an evaluation science for generative AI systems. arXiv [cs.AI] (March 2025)."},{"key":"e_1_3_2_1_72_1","volume-title":"Juan Mateos-Garcia, Stevie Bergman, Jackie Kay, Conor Griffin, Ben Bariach, Iason Gabriel, Verena Rieser, and William Isaac.","author":"Weidinger Laura","year":"2023","unstructured":"Laura Weidinger, Maribeth Rauh, Nahema Marchal, Arianna Manzini, Lisa Anne Hendricks, Juan Mateos-Garcia, Stevie Bergman, Jackie Kay, Conor Griffin, Ben Bariach, Iason Gabriel, Verena Rieser, and William Isaac. 2023. Sociotechnical Safety Evaluation of Generative AI Systems. arXiv [cs.AI] (Oct. 2023)."},{"key":"e_1_3_2_1_73_1","doi-asserted-by":"publisher","DOI":"10.1162\/tacl_a_00754"},{"key":"e_1_3_2_1_74_1","volume-title":"Proceedings of the CHI Conference on Human Factors in Computing Systems (CHI '24","author":"Wester Joel","year":"2024","unstructured":"Joel Wester, Tim Schrills, Henning Pohl, and Niels van Berkel. 2024. \u201cAs an AI language model, I cannot\u201d: Investigating LLM Denials of User Requests. In Proceedings of the CHI Conference on Human Factors in Computing Systems (CHI '24, Article 979). Association for Computing Machinery, New York, NY, USA, 1\u201314."},{"key":"e_1_3_2_1_75_1","doi-asserted-by":"publisher","DOI":"10.1145\/3479499"},{"key":"e_1_3_2_1_76_1","volume-title":"Kaixuan Huang, Luxi He, Boyi Wei, Dacheng Li, Ying Sheng, Ruoxi Jia, Bo Li, Kai Li, Danqi Chen, Peter Henderson, and Prateek Mittal.","author":"Xie Tinghao","year":"2024","unstructured":"Tinghao Xie, Xiangyu Qi, Yi Zeng, Yangsibo Huang, Udari Madhushani Sehwag, Kaixuan Huang, Luxi He, Boyi Wei, Dacheng Li, Ying Sheng, Ruoxi Jia, Bo Li, Kai Li, Danqi Chen, Peter Henderson, and Prateek Mittal. 2024. SORRY-bench: Systematically evaluating large language model safety refusal. arXiv [cs.AI] (June 2024)."},{"key":"e_1_3_2_1_77_1","volume-title":"Violeta J Rodriguez, and Koustuv Saha.","author":"Yoo Dong Whi","year":"2025","unstructured":"Dong Whi Yoo, Jiayue Melissa Shi, Violeta J Rodriguez, and Koustuv Saha. 2025. AI chatbots for mental health: Values and harms from lived experiences of depression. arXiv [cs.HC] (April 2025)."},{"key":"e_1_3_2_1_78_1","doi-asserted-by":"publisher","DOI":"10.1007\/s10676-019-09497-z"},{"key":"e_1_3_2_1_79_1","volume-title":"From hard refusals to safe-completions: Toward output-centric safety training. arXiv [cs.CY] (Aug","author":"Yuan Yuan","year":"2025","unstructured":"Yuan Yuan, Tina Sriskandarajah, Anna-Luisa Brakman, Alec Helyar, Alex Beutel, Andrea Vallone, and Saachi Jain. 2025. From hard refusals to safe-completions: Toward output-centric safety training. arXiv [cs.CY] (Aug. 2025)."},{"key":"e_1_3_2_1_80_1","doi-asserted-by":"publisher","DOI":"10.1145\/3706598.3713453"},{"key":"e_1_3_2_1_81_1","volume-title":"Universal and transferable adversarial attacks on aligned language models. arXiv [cs.CL] (July","author":"Zou Andy","year":"2023","unstructured":"Andy Zou, Zifan Wang, Nicholas Carlini, Milad Nasr, J Zico Kolter, and Matt Fredrikson. 2023. Universal and transferable adversarial attacks on aligned language models. arXiv [cs.CL] (July 2023)."}],"event":{"name":"FAccT '26: The 2026 ACM Conference on Fairness, Accountability, and Transparency","location":"Montreal QC Canada","acronym":"FAccT '26","sponsor":["ACM\/SIG"]},"container-title":["Proceedings of the 2026 ACM Conference on Fairness, Accountability, and Transparency"],"original-title":[],"link":[{"URL":"https:\/\/dl.acm.org\/doi\/abs\/10.1145\/3805689.3806723","content-type":"text\/html","content-version":"vor","intended-application":"syndication"},{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/3805689.3806723","content-type":"application\/pdf","content-version":"vor","intended-application":"syndication"},{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/3805689.3806723","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,6,29]],"date-time":"2026-06-29T18:10:02Z","timestamp":1782756602000},"score":1,"resource":{"primary":{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/3805689.3806723"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,6,25]]},"references-count":81,"alternative-id":["10.1145\/3805689.3806723","10.1145\/3805689"],"URL":"https:\/\/doi.org\/10.1145\/3805689.3806723","relation":{},"subject":[],"published":{"date-parts":[[2026,6,25]]},"assertion":[{"value":"2026-06-25","order":3,"name":"published","label":"Published","group":{"name":"publication_history","label":"Publication History"}}]}}