{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,30]],"date-time":"2026-06-30T06:58:30Z","timestamp":1782802710353,"version":"3.54.5"},"publisher-location":"New York, NY, USA","reference-count":110,"publisher":"ACM","license":[{"start":{"date-parts":[[2025,6,23]],"date-time":"2025-06-23T00:00:00Z","timestamp":1750636800000},"content-version":"vor","delay-in-days":0,"URL":"http:\/\/www.acm.org\/publications\/policies\/copyright_policy#Background"}],"funder":[{"DOI":"10.13039\/100000001","name":"NSF (National Science Foundation)","doi-asserted-by":"publisher","award":["2316768"],"award-info":[{"award-number":["2316768"]}],"id":[{"id":"10.13039\/100000001","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100008483","name":"NOMIS Stiftung","doi-asserted-by":"publisher","id":[{"id":"10.13039\/501100008483","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["dl.acm.org"],"crossmark-restriction":true},"short-container-title":[],"published-print":{"date-parts":[[2025,6,23]]},"DOI":"10.1145\/3715275.3732212","type":"proceedings-article","created":{"date-parts":[[2025,6,23]],"date-time":"2025-06-23T17:01:18Z","timestamp":1750698078000},"page":"3303-3325","update-policy":"https:\/\/doi.org\/10.1145\/crossmark-policy","source":"Crossref","is-referenced-by-count":9,"title":["Actions Speak Louder than Words: Agent Decisions Reveal Implicit Biases in Language Models"],"prefix":"10.1145","author":[{"ORCID":"https:\/\/orcid.org\/0009-0002-2885-5337","authenticated-orcid":false,"given":"Yuxuan","family":"Li","sequence":"first","affiliation":[{"name":"School of Computer Science, Carnegie Mellon University, Pittsburgh, Pennsylvania, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-4545-7859","authenticated-orcid":false,"given":"Hirokazu","family":"Shirado","sequence":"additional","affiliation":[{"name":"School of Computer Science, Carnegie Mellon University, Pittsburgh, Pennsylvania, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-9073-8054","authenticated-orcid":false,"given":"Sauvik","family":"Das","sequence":"additional","affiliation":[{"name":"School of Computer Science, Carnegie Mellon University, Pittsburgh, Pennsylvania, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"320","published-online":{"date-parts":[[2025,6,23]]},"reference":[{"key":"e_1_3_3_2_2_2","unstructured":"Josh Achiam Steven Adler Sandhini Agarwal Lama Ahmad Ilge Akkaya Florencia\u00a0Leoni Aleman Diogo Almeida Janko Altenschmidt Sam Altman Shyamal Anadkat et\u00a0al. 2023. Gpt-4 technical report. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2303.08774 (2023)."},{"key":"e_1_3_3_2_3_2","doi-asserted-by":"crossref","unstructured":"Xuechunzi Bai Angelina Wang Ilia Sucholutsky and Thomas\u00a0L Griffiths. 2025. Explicitly unbiased large language models still form biased associations. Proceedings of the National Academy of Sciences 122 8 (2025) e2416228122.","DOI":"10.1073\/pnas.2416228122"},{"key":"e_1_3_3_2_4_2","doi-asserted-by":"crossref","unstructured":"Delia Baldassarri and Maria Abascal. 2020. Diversity and prosocial behavior. Science 369 6508 (2020) 1183\u20131187.","DOI":"10.1126\/science.abb2432"},{"key":"e_1_3_3_2_5_2","doi-asserted-by":"publisher","DOI":"10.1145\/3442188.3445875"},{"key":"e_1_3_3_2_6_2","doi-asserted-by":"crossref","unstructured":"Soumya Barikeri Anne Lauscher Ivan Vuli\u0107 and Goran Glava\u0161. 2021. RedditBias: A real-world resource for bias evaluation and debiasing of conversational language models. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2106.03521 (2021).","DOI":"10.18653\/v1\/2021.acl-long.151"},{"key":"e_1_3_3_2_7_2","volume-title":"Fairness and machine learning: Limitations and opportunities","author":"Barocas Solon","year":"2023","unstructured":"Solon Barocas, Moritz Hardt, and Arvind Narayanan. 2023. Fairness and machine learning: Limitations and opportunities. MIT press."},{"key":"e_1_3_3_2_8_2","doi-asserted-by":"crossref","unstructured":"Julie\u00a0M Bateman and Bob Edwards. 2002. Gender and evacuation: A closer look at why women are more likely to evacuate for hurricanes. Natural Hazards Review 3 3 (2002) 107\u2013117.","DOI":"10.1061\/(ASCE)1527-6988(2002)3:3(107)"},{"key":"e_1_3_3_2_9_2","doi-asserted-by":"publisher","DOI":"10.1145\/3442188.3445922"},{"key":"e_1_3_3_2_10_2","doi-asserted-by":"crossref","unstructured":"Marianne Bertrand and Sendhil Mullainathan. 2004. Are Emily and Greg more employable than Lakisha and Jamal? A field experiment on labor market discrimination. American economic review 94 4 (2004) 991\u20131013.","DOI":"10.1257\/0002828042002561"},{"key":"e_1_3_3_2_11_2","doi-asserted-by":"crossref","unstructured":"Su\u00a0Lin Blodgett Solon Barocas Hal Daum\u00e9\u00a0III and Hanna Wallach. 2020. Language (technology) is power: A critical survey of\" bias\" in nlp. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2005.14050 (2020).","DOI":"10.18653\/v1\/2020.acl-main.485"},{"key":"e_1_3_3_2_12_2","unstructured":"Tolga Bolukbasi Kai-Wei Chang James\u00a0Y Zou Venkatesh Saligrama and Adam\u00a0T Kalai. 2016. Man is to computer programmer as woman is to homemaker? debiasing word embeddings. Advances in neural information processing systems 29 (2016)."},{"key":"e_1_3_3_2_13_2","doi-asserted-by":"crossref","unstructured":"Conrad Borchers Dalia\u00a0Sara Gala Benjamin Gilburt Eduard Oravkin Wilfried Bounsi Yuki\u00a0M Asano and Hannah\u00a0Rose Kirk. 2022. Looking for a handsome carpenter! debiasing GPT-3 job advertisements. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2205.11374 (2022).","DOI":"10.18653\/v1\/2022.gebnlp-1.22"},{"key":"e_1_3_3_2_14_2","doi-asserted-by":"crossref","unstructured":"Shikha Bordia and Samuel\u00a0R Bowman. 2019. Identifying and reducing gender bias in word-level language models. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/1904.03035 (2019).","DOI":"10.18653\/v1\/N19-3002"},{"key":"e_1_3_3_2_15_2","doi-asserted-by":"crossref","unstructured":"Thomas Buser Muriel Niederle and Hessel Oosterbeek. 2014. Gender competitiveness and career choices. The quarterly journal of economics 129 3 (2014) 1409\u20131447.","DOI":"10.1093\/qje\/qju009"},{"key":"e_1_3_3_2_16_2","doi-asserted-by":"crossref","unstructured":"Ignatius Cahyanto Lori Pennington-Gray Brijesh Thapa Siva Srinivasan Jorge Villegas Corene Matyas and Spiro Kiousis. 2014. An empirical evaluation of the determinants of tourist\u2019s hurricane evacuation decision making. Journal of Destination Marketing & Management 2 4 (2014) 253\u2013265.","DOI":"10.1016\/j.jdmm.2013.10.003"},{"key":"e_1_3_3_2_17_2","volume-title":"Proceedings of the 2019 Conference on Empirical Methods in Natural Language Processing and the 9th International Joint Conference on Natural Language Processing (EMNLP-IJCNLP): Tutorial Abstracts","author":"Chang Kai-Wei","year":"2019","unstructured":"Kai-Wei Chang, Vinodkumar Prabhakaran, and Vicente Ordonez. 2019. Bias and fairness in natural language processing. In Proceedings of the 2019 Conference on Empirical Methods in Natural Language Processing and the 9th International Joint Conference on Natural Language Processing (EMNLP-IJCNLP): Tutorial Abstracts."},{"key":"e_1_3_3_2_18_2","doi-asserted-by":"crossref","unstructured":"Myra Cheng Esin Durmus and Dan Jurafsky. 2023. Marked personas: Using natural language prompts to measure stereotypes in language models. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2305.18189 (2023).","DOI":"10.18653\/v1\/2023.acl-long.84"},{"key":"e_1_3_3_2_19_2","doi-asserted-by":"publisher","DOI":"10.1017\/CBO9780511802843"},{"key":"e_1_3_3_2_20_2","doi-asserted-by":"crossref","unstructured":"Hengfang Deng Daniel\u00a0P Aldrich Michael\u00a0M Danziger Jianxi Gao Nolan\u00a0E Phillips Sean\u00a0P Cornelius and Qi\u00a0Ryan Wang. 2021. High-resolution human mobility data reveal race and wealth disparities in disaster evacuation patterns. Humanities and Social Sciences Communications 8 1 (2021) 1\u20138.","DOI":"10.1057\/s41599-021-00824-8"},{"key":"e_1_3_3_2_21_2","doi-asserted-by":"publisher","DOI":"10.1145\/3442188.3445924"},{"key":"e_1_3_3_2_22_2","doi-asserted-by":"crossref","unstructured":"Thomas\u00a0P Dick and Sharon\u00a0F Rallis. 1991. Factors and influences on high school students\u2019 career choices. Journal for research in mathematics education 22 4 (1991) 281\u2013292.","DOI":"10.5951\/jresematheduc.22.4.0281"},{"key":"e_1_3_3_2_23_2","doi-asserted-by":"crossref","unstructured":"Danica Dillion Niket Tandon Yuling Gu and Kurt Gray. 2023. Can AI language models replace human participants? Trends in Cognitive Sciences 27 7 (2023) 597\u2013600.","DOI":"10.1016\/j.tics.2023.04.008"},{"key":"e_1_3_3_2_24_2","unstructured":"Abhimanyu Dubey Abhinav Jauhri Abhinav Pandey Abhishek Kadian Ahmad Al-Dahle Aiesha Letman Akhil Mathur Alan Schelten Amy Yang Angela Fan et\u00a0al. 2024. The llama 3 herd of models. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2407.21783 (2024)."},{"key":"e_1_3_3_2_25_2","unstructured":"Tyna Eloundou Alex Beutel David\u00a0G Robinson Keren Gu-Lemberg Anna-Luisa Brakman Pamela Mishkin Meghan Shah Johannes Heidecke Lilian Weng and Adam\u00a0Tauman Kalai. 2024. First-Person Fairness in Chatbots. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2410.19803 (2024)."},{"key":"e_1_3_3_2_26_2","doi-asserted-by":"publisher","DOI":"10.1145\/3593013.3594094"},{"key":"e_1_3_3_2_27_2","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/2022.emnlp-main.651"},{"key":"e_1_3_3_2_28_2","doi-asserted-by":"publisher","unstructured":"Isabel\u00a0O. Gallegos Ryan\u00a0A. Rossi Joe Barrow Md\u00a0Mehrab Tanjim Sungchul Kim Franck Dernoncourt Tong Yu Ruiyi Zhang and Nesreen\u00a0K. Ahmed. 2024. Bias and Fairness in Large Language Models: A Survey. Computational Linguistics 50 3 (09 2024) 1097\u20131179. 10.1162\/coli_a_00524 arXiv:https:\/\/direct.mit.edu\/coli\/article-pdf\/50\/3\/1097\/2471010\/coli_a_00524.pdf","DOI":"10.1162\/coli_a_00524"},{"key":"e_1_3_3_2_29_2","doi-asserted-by":"crossref","unstructured":"Chen Gao Xiaochong Lan Nian Li Yuan Yuan Jingtao Ding Zhilun Zhou Fengli Xu and Yong Li. 2024. Large language models empowered agent-based modeling and simulation: A survey and perspectives. Humanities and Social Sciences Communications 11 1 (2024) 1\u201324.","DOI":"10.1057\/s41599-024-03611-3"},{"key":"e_1_3_3_2_30_2","doi-asserted-by":"crossref","unstructured":"Chen Gao Xiaochong Lan Zhihong Lu Jinzhu Mao Jinghua Piao Huandong Wang Depeng Jin and Yong Li. 2023. S3: Social-network simulation system with large language model-empowered agents. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2307.14984 (2023).","DOI":"10.2139\/ssrn.4607026"},{"key":"e_1_3_3_2_31_2","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/2022.aacl-short.38"},{"key":"e_1_3_3_2_32_2","doi-asserted-by":"crossref","unstructured":"Samuel Gehman Suchin Gururangan Maarten Sap Yejin Choi and Noah\u00a0A Smith. 2020. Realtoxicityprompts: Evaluating neural toxic degeneration in language models. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2009.11462 (2020).","DOI":"10.18653\/v1\/2020.findings-emnlp.301"},{"key":"e_1_3_3_2_33_2","unstructured":"Brian Glassman. 2023. Financial insecurity and hardship in the pulse: an in-depth look at gender identity and sexual orientation. US Census Bureau. URL: https:\/\/tinyurl. com\/sszn8kym [accessed 2023-11-30] (2023)."},{"key":"e_1_3_3_2_34_2","doi-asserted-by":"publisher","DOI":"10.1145\/3630106.3658933"},{"key":"e_1_3_3_2_35_2","doi-asserted-by":"publisher","DOI":"10.1145\/3593013.3594067"},{"key":"e_1_3_3_2_36_2","doi-asserted-by":"crossref","unstructured":"Ariel Hasell and Brian\u00a0E Weeks. 2016. Partisan provocation: The role of partisan news use and emotional responses in political information sharing in social media. Human Communication Research 42 4 (2016) 641\u2013661.","DOI":"10.1111\/hcre.12092"},{"key":"e_1_3_3_2_37_2","doi-asserted-by":"crossref","unstructured":"Phil\u00a0JM Heiligers. 2012. Gender differences in medical students\u2019 motives and career choice. BMC medical education 12 (2012) 1\u201311.","DOI":"10.1186\/1472-6920-12-82"},{"key":"e_1_3_3_2_38_2","doi-asserted-by":"publisher","DOI":"10.1145\/3531146.3533184"},{"key":"e_1_3_3_2_39_2","doi-asserted-by":"crossref","unstructured":"Valentin Hofmann Pratyusha\u00a0Ria Kalluri Dan Jurafsky and Sharese King. 2024. AI generates covertly racist decisions about people based on their dialect. Nature 633 8028 (2024) 147\u2013154.","DOI":"10.1038\/s41586-024-07856-5"},{"key":"e_1_3_3_2_40_2","doi-asserted-by":"crossref","unstructured":"Michael\u00a0A Hogg and Scott\u00a0A Reid. 2006. Social identity self-categorization and the communication of group norms. Communication theory 16 1 (2006) 7\u201330.","DOI":"10.1111\/j.1468-2885.2006.00003.x"},{"key":"e_1_3_3_2_41_2","unstructured":"Jen-tse Huang Eric\u00a0John Li Man\u00a0Ho Lam Tian Liang Wenxuan Wang Youliang Yuan Wenxiang Jiao Xing Wang Zhaopeng Tu and Michael\u00a0R Lyu. 2024. How Far Are We on the Decision-Making of LLMs? Evaluating LLMs\u2019 Gaming Ability in Multi-Agent Environments. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2403.11807 (2024)."},{"key":"e_1_3_3_2_42_2","doi-asserted-by":"crossref","unstructured":"Shih-Kai Huang Michael\u00a0K Lindell Carla\u00a0S Prater Hao-Che Wu and Laura\u00a0K Siebeneck. 2012. Household evacuation decision making in response to Hurricane Ike. Natural Hazards Review 13 4 (2012) 283\u2013296.","DOI":"10.1061\/(ASCE)NH.1527-6996.0000074"},{"key":"e_1_3_3_2_43_2","unstructured":"Yue Huang Qihui Zhang Lichao Sun et\u00a0al. 2023. Trustgpt: A benchmark for trustworthy and responsible large language models. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2306.11507 (2023)."},{"key":"e_1_3_3_2_44_2","unstructured":"Nishtha Jain Maja Popovic Declan Groves and Eva Vanmassenhove. 2021. Generating gender augmented data for NLP. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2107.05987 (2021)."},{"key":"e_1_3_3_2_45_2","volume-title":"NeurIPS 2023 Foundation Models for Decision Making Workshop","author":"Jarrett Daniel","year":"2023","unstructured":"Daniel Jarrett, Miruna Pislar, Michiel\u00a0A Bakker, Michael\u00a0Henry Tessler, Raphael Koster, Jan Balaguer, Romuald Elie, Christopher Summerfield, and Andrea Tacchetti. 2023. Language agents as digital representatives in collective decision-making. In NeurIPS 2023 Foundation Models for Decision Making Workshop."},{"key":"e_1_3_3_2_46_2","unstructured":"Wonje Jeung Dongjae Jeon Ashkan Yousefpour and Jonghyun Choi. 2024. Large Language Models Still Exhibit Bias in Long Text. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2410.17519 (2024)."},{"key":"e_1_3_3_2_47_2","unstructured":"Jiaming Ji Mickel Liu Josef Dai Xuehai Pan Chi Zhang Ce Bian Boyuan Chen Ruiyang Sun Yizhou Wang and Yaodong Yang. 2024. Beavertails: Towards improved safety alignment of llm via a human-preference dataset. Advances in Neural Information Processing Systems 36 (2024)."},{"key":"e_1_3_3_2_48_2","unstructured":"Jingru Jia Zehua Yuan Junhao Pan Paul\u00a0E McNamara and Deming Chen. 2024. Decision-making behavior evaluation framework for llms under uncertain context. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2406.05972 (2024)."},{"key":"e_1_3_3_2_49_2","unstructured":"Albert\u00a0Q Jiang Alexandre Sablayrolles Antoine Roux Arthur Mensch Blanche Savary Chris Bamford Devendra\u00a0Singh Chaplot Diego de\u00a0las Casas Emma\u00a0Bou Hanna Florian Bressand et\u00a0al. 2024. Mixtral of experts. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2401.04088 (2024)."},{"key":"e_1_3_3_2_50_2","unstructured":"Junfeng Jiao Saleh Afroogh Yiming Xu and Connor Phillips. 2024. Navigating llm ethics: Advancements challenges and future directions. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2406.18841 (2024)."},{"key":"e_1_3_3_2_51_2","doi-asserted-by":"crossref","unstructured":"Ryuichi Kawamoto Daisuke Ninomiya Yoshihisa Kasai Tomo Kusunoki Nobuyuki Ohtsuka Teru Kumagi and Masanori Abe. 2016. Gender difference in preference of specialty as a career choice among Japanese medical students. BMC medical education 16 (2016) 1\u20138.","DOI":"10.1186\/s12909-016-0811-1"},{"key":"e_1_3_3_2_52_2","doi-asserted-by":"crossref","unstructured":"Sven Kepes George\u00a0C Banks and In-Sue Oh. 2014. Avoiding bias in publication bias research: The value of \u201cnull\u201d findings. Journal of Business and Psychology 29 (2014) 183\u2013203.","DOI":"10.1007\/s10869-012-9279-0"},{"key":"e_1_3_3_2_53_2","unstructured":"Hyunwoo Kim Youngjae Yu Liwei Jiang Ximing Lu Daniel Khashabi Gunhee Kim Yejin Choi and Maarten Sap. 2022. Prosocialdialog: A prosocial backbone for conversational agents. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2205.12688 (2022)."},{"key":"e_1_3_3_2_54_2","unstructured":"Hannah\u00a0Rose Kirk Yennie Jun Filippo Volpin Haider Iqbal Elias Benussi Frederic Dreyer Aleksandar Shtedritski and Yuki Asano. 2021. Bias out-of-the-box: An empirical analysis of intersectional occupational biases in popular generative language models. Advances in neural information processing systems 34 (2021) 2611\u20132624."},{"key":"e_1_3_3_2_55_2","doi-asserted-by":"publisher","DOI":"10.1145\/3576840.3578295"},{"key":"e_1_3_3_2_56_2","doi-asserted-by":"crossref","unstructured":"Anne Lauscher Tobias Lueken and Goran Glava\u0161. 2021. Sustainable modular debiasing of language models. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2109.03646 (2021).","DOI":"10.18653\/v1\/2021.findings-emnlp.411"},{"key":"e_1_3_3_2_57_2","doi-asserted-by":"crossref","unstructured":"M\u00a0Asher Lawson and Hemant Kakkar. 2022. Of pandemics politics and personality: The role of conscientiousness and political ideology in the sharing of fake news. Journal of Experimental Psychology: General 151 5 (2022) 1154.","DOI":"10.1037\/xge0001120"},{"key":"e_1_3_3_2_58_2","doi-asserted-by":"crossref","unstructured":"Chang-Woo Lee. 2013. Gender difference and specialty preference in medical career choice. Korean journal of medical education 25 1 (2013) 15.","DOI":"10.3946\/kjme.2013.25.1.15"},{"key":"e_1_3_3_2_59_2","doi-asserted-by":"publisher","DOI":"10.1145\/3630106.3658975"},{"key":"e_1_3_3_2_60_2","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/2024.acl-long.829"},{"key":"e_1_3_3_2_61_2","doi-asserted-by":"crossref","unstructured":"Tao Li Tushar Khot Daniel Khashabi Ashish Sabharwal and Vivek Srikumar. 2020. UNQOVERing stereotyping biases via underspecified questions. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2010.02428 (2020).","DOI":"10.18653\/v1\/2020.findings-emnlp.311"},{"key":"e_1_3_3_2_62_2","unstructured":"Percy Liang Rishi Bommasani Tony Lee Dimitris Tsipras Dilara Soylu Michihiro Yasunaga Yian Zhang Deepak Narayanan Yuhuai Wu Ananya Kumar et\u00a0al. 2022. Holistic evaluation of language models. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2211.09110 (2022)."},{"key":"e_1_3_3_2_63_2","first-page":"6565","volume-title":"International Conference on Machine Learning","author":"Liang Paul\u00a0Pu","year":"2021","unstructured":"Paul\u00a0Pu Liang, Chiyu Wu, Louis-Philippe Morency, and Ruslan Salakhutdinov. 2021. Towards understanding and mitigating social biases in language models. In International Conference on Machine Learning. PMLR, 6565\u20136576."},{"key":"e_1_3_3_2_64_2","unstructured":"Michael\u00a0K Lindell Ronald\u00a0W Perry and Marjorie\u00a0R Greene. 1980. Race and disaster warning response. Research paper. Columbus Ohio: Battelle Human Affairs Research Center (1980)."},{"key":"e_1_3_3_2_65_2","unstructured":"Yang Liu Yuanshun Yao Jean-Francois Ton Xiaoying Zhang Ruocheng Guo\u00a0Hao Cheng Yegor Klochkov Muhammad\u00a0Faaiz Taufiq and Hang Li. 2023. Trustworthy LLMs: A survey and guideline for evaluating large language models\u2019 alignment. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2308.05374 (2023)."},{"key":"e_1_3_3_2_66_2","doi-asserted-by":"crossref","unstructured":"Elisa\u00a0F Long M\u00a0Keith Chen and Ryne Rohla. 2020. Political storms: Emergent partisan skepticism of hurricane risks. Science advances 6 37 (2020) eabb7906.","DOI":"10.1126\/sciadv.abb7906"},{"key":"e_1_3_3_2_67_2","doi-asserted-by":"crossref","unstructured":"Kaiji Lu Piotr Mardziel Fangjing Wu Preetam Amancharla and Anupam Datta. 2020. Gender bias in neural natural language processing. Logic language and security: essays dedicated to Andre Scedrov on the occasion of his 65th birthday (2020) 189\u2013202.","DOI":"10.1007\/978-3-030-62077-6_14"},{"key":"e_1_3_3_2_68_2","doi-asserted-by":"crossref","unstructured":"Michael\u00a0W Macy and Robert Willer. 2002. From factors to actors: Computational sociology and agent-based modeling. Annual review of sociology 28 1 (2002) 143\u2013166.","DOI":"10.1146\/annurev.soc.28.110601.141117"},{"key":"e_1_3_3_2_69_2","doi-asserted-by":"publisher","DOI":"10.1145\/3630106.3658967"},{"key":"e_1_3_3_2_70_2","doi-asserted-by":"crossref","unstructured":"Najib\u00a0A Mozahem Dana\u00a0K Kozbar Ahmad\u00a0W Al\u00a0Hassan and Laila\u00a0A Mozahem. 2020. Gender differences in career choices among students in secondary school. International Journal of School & Educational Psychology 8 3 (2020) 184\u2013198.","DOI":"10.1080\/21683603.2018.1521759"},{"key":"e_1_3_3_2_71_2","doi-asserted-by":"crossref","unstructured":"David Neumark Roy\u00a0J Bank and Kyle\u00a0D Van\u00a0Nort. 1996. Sex discrimination in restaurant hiring: An audit study. The Quarterly journal of economics 111 3 (1996) 915\u2013941.","DOI":"10.2307\/2946676"},{"key":"e_1_3_3_2_72_2","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/2021.naacl-main.191"},{"key":"e_1_3_3_2_73_2","unstructured":"Abiodun\u00a0Finbarrs Oketunji Muhammad Anas and Deepthi Saina. 2023. Large Language Model (LLM) Bias Index\u2013LLMBI. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2312.14769 (2023)."},{"key":"e_1_3_3_2_74_2","doi-asserted-by":"crossref","unstructured":"Mathias Osmundsen Alexander Bor Peter\u00a0Bjerregaard Vahlstrup Anja Bechmann and Michael\u00a0Bang Petersen. 2021. Partisan polarization is the primary psychological motivation behind political fake news sharing on Twitter. American Political Science Review 115 3 (2021) 999\u20131015.","DOI":"10.1017\/S0003055421000290"},{"key":"e_1_3_3_2_75_2","doi-asserted-by":"crossref","unstructured":"Devah Pager and Lincoln Quillian. 2005. Walking the talk? What employers say versus what they do. American sociological review 70 3 (2005) 355\u2013380.","DOI":"10.1177\/000312240507000301"},{"key":"e_1_3_3_2_76_2","doi-asserted-by":"publisher","DOI":"10.1145\/3586183.3606763"},{"key":"e_1_3_3_2_77_2","unstructured":"Joon\u00a0Sung Park Carolyn\u00a0Q Zou Aaron Shaw Benjamin\u00a0Mako Hill Carrie Cai Meredith\u00a0Ringel Morris Robb Willer Percy Liang and Michael\u00a0S Bernstein. 2024. Generative agent simulations of 1 000 people. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2411.10109 (2024)."},{"key":"e_1_3_3_2_78_2","unstructured":"Alicia Parrish Angelica Chen Nikita Nangia Vishakh Padmakumar Jason Phang Jana Thompson Phu\u00a0Mon Htut and Samuel\u00a0R Bowman. 2021. BBQ: A hand-built bias benchmark for question answering. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2110.08193 (2021)."},{"key":"e_1_3_3_2_79_2","doi-asserted-by":"crossref","unstructured":"Rebecca Qian Candace Ross Jude Fernandes Eric Smith Douwe Kiela and Adina Williams. 2022. Perturbation augmentation for fairer nlp. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2205.12586 (2022).","DOI":"10.18653\/v1\/2022.emnlp-main.646"},{"key":"e_1_3_3_2_80_2","doi-asserted-by":"crossref","unstructured":"Yusu Qian Urwa Muaz Ben Zhang and Jae\u00a0Won Hyun. 2019. Reducing gender bias in word-level language models with a gender-equalizing loss function. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/1905.12801 (2019).","DOI":"10.18653\/v1\/P19-2031"},{"key":"e_1_3_3_2_81_2","doi-asserted-by":"crossref","unstructured":"Jasmin\u00a0K Riad Fran\u00a0H Norris and R\u00a0Barry Ruback. 1999. Predicting evacuation in two major disasters: Risk perception social influence and access to resources 1. Journal of applied social Psychology 29 5 (1999) 918\u2013934.","DOI":"10.1111\/j.1559-1816.1999.tb00132.x"},{"key":"e_1_3_3_2_82_2","volume-title":"CHI Conference on Human Factors in Computing Systems","author":"Robinson Katherine-Marie","year":"2024","unstructured":"Katherine-Marie Robinson, Violet Turri, Carol\u00a0J Smith, and Shannon\u00a0K Gallagher. 2024. Tales from the Wild West: Crafting Scenarios to Audit Bias in LLMs. In CHI Conference on Human Factors in Computing Systems."},{"key":"e_1_3_3_2_83_2","doi-asserted-by":"crossref","unstructured":"Marlene\u00a0M Rosenkoetter Eleanor\u00a0Krassen Covan Brenda\u00a0K Cobb Sheila Bunting and Martin Weinrich. 2007. Perceptions of older adults regarding evacuation in the event of a natural disaster. Public Health Nursing 24 2 (2007) 160\u2013168.","DOI":"10.1111\/j.1525-1446.2007.00620.x"},{"key":"e_1_3_3_2_84_2","unstructured":"Lydia Saad. 2022. US political ideology steady; conservatives moderates tie. Gallup News. Available at: https:\/\/news. gallup. com\/poll\/388988\/political-ideology-steady-conservatives-moderates-tie. aspx (2022)."},{"key":"e_1_3_3_2_85_2","doi-asserted-by":"crossref","unstructured":"Wanyun Shao and Feng Hao. 2020. Confidence in political leaders can slant risk perceptions of COVID\u201319 in a highly polarized environment. Social science & medicine 261 (2020) 113235.","DOI":"10.1016\/j.socscimed.2020.113235"},{"key":"e_1_3_3_2_86_2","unstructured":"Hua Shen Nicholas Clark and Tanushree Mitra. 2025. Mind the Value-Action Gap: Do LLMs Act in Alignment with Their Values? arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2501.15463 (2025)."},{"key":"e_1_3_3_2_87_2","doi-asserted-by":"crossref","unstructured":"Jieun Shin and Kjerstin Thorson. 2017. Partisan selective sharing: The biased diffusion of fact-checking messages on social media. Journal of communication 67 2 (2017) 233\u2013255.","DOI":"10.1111\/jcom.12284"},{"key":"e_1_3_3_2_88_2","unstructured":"Eric\u00a0Michael Smith Melissa Hall Melanie Kambadur Eleonora Presani and Adina Williams. 2022. \" I\u2019m sorry to hear that\": Finding New Biases in Language Models with a Holistic Descriptor Dataset. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2205.09209 (2022)."},{"key":"e_1_3_3_2_89_2","doi-asserted-by":"crossref","unstructured":"Stanley\u00a0K Smith and Chris McCarty. 2009. Fleeing the storm (s): An examination of evacuation behavior during Florida\u2019s 2004 hurricane season. Demography 46 (2009) 127\u2013145.","DOI":"10.1353\/dem.0.0048"},{"key":"e_1_3_3_2_90_2","volume-title":"Measuring occupational prestige on the 2012 general social survey","author":"Smith Tom\u00a0William","year":"2014","unstructured":"Tom\u00a0William Smith and Jaesok Son. 2014. Measuring occupational prestige on the 2012 general social survey. Vol.\u00a04. NORC at the University of Chicago Chicago."},{"key":"e_1_3_3_2_91_2","doi-asserted-by":"crossref","unstructured":"Jan\u00a0E Stets and Peter\u00a0J Burke. 2000. Identity theory and social identity theory. Social psychology quarterly (2000) 224\u2013237.","DOI":"10.2307\/2695870"},{"key":"e_1_3_3_2_92_2","unstructured":"Theodore\u00a0R Sumers Shunyu Yao Karthik Narasimhan and Thomas\u00a0L Griffiths. 2023. Cognitive architectures for language agents. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2309.02427 (2023)."},{"key":"e_1_3_3_2_93_2","volume-title":"Reinforcement learning: An introduction","author":"Sutton Richard\u00a0S","year":"2018","unstructured":"Richard\u00a0S Sutton and Andrew\u00a0G Barto. 2018. Reinforcement learning: An introduction. MIT press."},{"key":"e_1_3_3_2_94_2","unstructured":"Carmen Tanner Adrian Br\u00fcgger Susan van Schie and Carmen Lebherz. 2015. Actions speak louder than words. Zeitschrift f\u00fcr Psychologie\/Journal of Psychology (2015)."},{"key":"e_1_3_3_2_95_2","unstructured":"Gemini Team Rohan Anil Sebastian Borgeaud Jean-Baptiste Alayrac Jiahui Yu Radu Soricut Johan Schalkwyk Andrew\u00a0M Dai Anja Hauth Katie Millican et\u00a0al. 2023. Gemini: a family of highly capable multimodal models. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2312.11805 (2023)."},{"key":"e_1_3_3_2_96_2","unstructured":"Vishesh Thakur. 2023. Unveiling gender bias in terms of profession across LLMs: Analyzing and addressing sociological implications. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2307.09162 (2023)."},{"key":"e_1_3_3_2_97_2","unstructured":"Ewoenam\u00a0Kwaku Tokpo and Toon Calders. 2022. Text style transfer for bias mitigation using masked language modeling. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2201.08643 (2022)."},{"key":"e_1_3_3_2_98_2","unstructured":"U.S. Census Bureau. n. d.. Data.census.gov Table. https:\/\/data.census.gov\/table Accessed: 2025-01-13."},{"key":"e_1_3_3_2_99_2","doi-asserted-by":"crossref","unstructured":"Eva Vanmassenhove Chris Emmery and Dimitar Shterionov. 2021. Neutral rewriter: A rule-based and neural approach to automatic rewriting into gender-neutral alternatives. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2109.06105 (2021).","DOI":"10.18653\/v1\/2021.emnlp-main.704"},{"key":"e_1_3_3_2_100_2","unstructured":"Rita Vine. 2006. Google scholar. Journal of the Medical Library Association 94 1 (2006) 97."},{"key":"e_1_3_3_2_101_2","doi-asserted-by":"crossref","unstructured":"Lei Wang Chen Ma Xueyang Feng Zeyu Zhang Hao Yang Jingsen Zhang Zhiyuan Chen Jiakai Tang Xu Chen Yankai Lin et\u00a0al. 2024. A survey on large language model based autonomous agents. Frontiers of Computer Science 18 6 (2024) 186345.","DOI":"10.1007\/s11704-024-40231-1"},{"key":"e_1_3_3_2_102_2","doi-asserted-by":"crossref","unstructured":"Brian\u00a0E Weeks Daniel\u00a0S Lane Dam\u00a0Hee Kim Slgi\u00a0S Lee and Nojin Kwak. 2017. Incidental exposure selective exposure and political information sharing: Integrating online exposure patterns and expression on social media. Journal of computer-mediated communication 22 6 (2017) 363\u2013379.","DOI":"10.1111\/jcc4.12199"},{"key":"e_1_3_3_2_103_2","doi-asserted-by":"crossref","unstructured":"Jason Weismueller Richard\u00a0L Gruner Paul Harrigan Kristof Coussement and Shasha Wang. 2024. Information sharing and political polarisation on social media: The role of falsehood and partisanship. Information Systems Journal 34 3 (2024) 854\u2013893.","DOI":"10.1111\/isj.12453"},{"key":"e_1_3_3_2_104_2","unstructured":"Michael Williams and Tami Moser. 2019. The art of coding and thematic exploration in qualitative research. International management review 15 1 (2019) 45\u201355."},{"key":"e_1_3_3_2_105_2","unstructured":"Zhiheng Xi Wenxiang Chen Xin Guo Wei He Yiwen Ding Boyang Hong Ming Zhang Junzhe Wang Senjie Jin Enyu Zhou et\u00a0al. 2023. The rise and potential of large language model based agents: A survey. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2309.07864 (2023)."},{"key":"e_1_3_3_2_106_2","unstructured":"An Yang Baosong Yang Beichen Zhang Binyuan Hui Bo Zheng Bowen Yu Chengyuan Li Dayiheng Liu Fei Huang Haoran Wei et\u00a0al. 2024. Qwen2. 5 Technical Report. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2412.15115 (2024)."},{"key":"e_1_3_3_2_107_2","unstructured":"Sherry Yang Ofir Nachum Yilun Du Jason Wei Pieter Abbeel and Dale Schuurmans. 2023. Foundation models for decision making: Problems methods and opportunities. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2303.04129 (2023)."},{"key":"e_1_3_3_2_108_2","unstructured":"Ziyi Yang Zaibin Zhang Zirui Zheng Yuxian Jiang Ziyue Gan Zhiyu Wang Zijian Ling Jinsong Chen Martz Ma Bowen Dong et\u00a0al. 2024. Oasis: Open agents social interaction simulations on one million agents. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2411.11581 (2024)."},{"key":"e_1_3_3_2_109_2","doi-asserted-by":"crossref","unstructured":"Ilker Yildirim and LA Paul. 2024. From task structures to world models: what do LLMs know? Trends in Cognitive Sciences (2024).","DOI":"10.1016\/j.tics.2024.02.008"},{"key":"e_1_3_3_2_110_2","unstructured":"Shuyan Zhou Frank\u00a0F Xu Hao Zhu Xuhui Zhou Robert Lo Abishek Sridhar Xianyi Cheng Tianyue Ou Yonatan Bisk Daniel Fried et\u00a0al. 2023. Webarena: A realistic web environment for building autonomous agents. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2307.13854 (2023)."},{"key":"e_1_3_3_2_111_2","unstructured":"Xuhui Zhou Hao Zhu Leena Mathur Ruohong Zhang Haofei Yu Zhengyang Qi Louis-Philippe Morency Yonatan Bisk Daniel Fried Graham Neubig et\u00a0al. 2023. Sotopia: Interactive evaluation for social intelligence in language agents. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2310.11667 (2023)."}],"event":{"name":"FAccT '25: The 2025 ACM Conference on Fairness, Accountability, and Transparency","location":"Athens Greece","acronym":"FAccT '25"},"container-title":["Proceedings of the 2025 ACM Conference on Fairness, Accountability, and Transparency"],"original-title":[],"link":[{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/3715275.3732212","content-type":"application\/pdf","content-version":"vor","intended-application":"syndication"},{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/3715275.3732212","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,6,24]],"date-time":"2025-06-24T11:06:47Z","timestamp":1750763207000},"score":1,"resource":{"primary":{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/3715275.3732212"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,6,23]]},"references-count":110,"alternative-id":["10.1145\/3715275.3732212","10.1145\/3715275"],"URL":"https:\/\/doi.org\/10.1145\/3715275.3732212","relation":{},"subject":[],"published":{"date-parts":[[2025,6,23]]},"assertion":[{"value":"2025-06-23","order":3,"name":"published","label":"Published","group":{"name":"publication_history","label":"Publication History"}}]}}