{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,27]],"date-time":"2026-07-27T14:00:25Z","timestamp":1785160825308,"version":"3.55.0"},"reference-count":205,"publisher":"Springer Science and Business Media LLC","issue":"6","license":[{"start":{"date-parts":[[2025,10,13]],"date-time":"2025-10-13T00:00:00Z","timestamp":1760313600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,10,13]],"date-time":"2025-10-13T00:00:00Z","timestamp":1760313600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/100000001","name":"National Science Foundation","doi-asserted-by":"publisher","award":["2125858"],"award-info":[{"award-number":["2125858"]}],"id":[{"id":"10.13039\/100000001","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["AI Ethics"],"published-print":{"date-parts":[[2025,12]]},"DOI":"10.1007\/s43681-025-00814-5","type":"journal-article","created":{"date-parts":[[2025,10,13]],"date-time":"2025-10-13T09:03:15Z","timestamp":1760346195000},"page":"5795-5819","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":49,"title":["Navigating LLM ethics: advancements, challenges, and future directions"],"prefix":"10.1007","volume":"5","author":[{"given":"Junfeng","family":"Jiao","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0004-4883-0844","authenticated-orcid":false,"given":"Saleh","family":"Afroogh","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yiming","family":"Xu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Connor","family":"Phillips","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2025,10,13]]},"reference":[{"issue":"2","key":"814_CR1","doi-asserted-by":"publisher","DOI":"10.1097\/CR9.0000000000000041","volume":"3","author":"PS Hinds","year":"2023","unstructured":"Hinds, P.S., Miller, A.B.: Our words and the words of artificial intelligence: the accountability belongs to us. Cancer Care Res. Online 3(2), e041 (2023). https:\/\/doi.org\/10.1097\/CR9.0000000000000041","journal-title":"Cancer Care Res. Online"},{"key":"814_CR2","doi-asserted-by":"publisher","unstructured":"Wei, C., Wang, Y.-C., Wang, B., and Kuo, C.-C. J., \u201cAn overview on language models: recent developments and outlook,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2303.05759.","DOI":"10.48550\/arXiv.2303.05759"},{"key":"814_CR3","doi-asserted-by":"publisher","unstructured":"Bender, E. M., Gebru, T., McMillan-Major, A., and Shmitchell, S., \u201cOn the dangers of stochastic parrots: can language models be too big?,\u201d in FAccT \u201921: 2021 ACM Conference on Fairness, Accountability, and Transparency, ACM, 2021, pp. 610\u2013623. https:\/\/doi.org\/10.1145\/3442188.3445922.","DOI":"10.1145\/3442188.3445922"},{"key":"814_CR4","doi-asserted-by":"publisher","unstructured":"Wei, J. et al., \u201cEmergent abilities of large language models,\u201d 2022, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2206.07682.","DOI":"10.48550\/arXiv.2206.07682"},{"key":"814_CR5","unstructured":"Wei, J. et al., \u201cChain-of-thought prompting elicits reasoning in large language models\u201d, [Online]. Available: files\/9927\/Wei et al. - Chain-of-Thought Prompting Elicits Reasoning in La.pdf"},{"issue":"9","key":"814_CR6","doi-asserted-by":"publisher","first-page":"866","DOI":"10.1001\/jama.2023.14217","volume":"330","author":"NH Shah","year":"2023","unstructured":"Shah, N.H., Entwistle, D., Pfeffer, M.A.: Creation and adoption of large language models in medicine. JAMA 330(9), 866\u2013869 (2023). https:\/\/doi.org\/10.1001\/jama.2023.14217","journal-title":"JAMA"},{"key":"814_CR7","doi-asserted-by":"publisher","unstructured":"Kaddour, J., Harris, J., Mozes, M., Bradley, H., Raileanu, R., and McHardy, R., \u201cChallenges and applications of large language models,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2307.10169.","DOI":"10.48550\/arXiv.2307.10169"},{"key":"814_CR8","doi-asserted-by":"publisher","unstructured":"Yu, H., Shen, Z., Miao, C., Leung, C., Lesser, V.R., and Yang, Q., \u201cBuilding ethics into artificial intelligence,\u201d in Twenty-Seventh International Joint Conference on Artificial Intelligence {IJCAI-18}, International Joint Conferences on Artificial Intelligence Organization, 2018, pp. 5527\u20135533. https:\/\/doi.org\/10.24963\/ijcai.2018\/779.","DOI":"10.24963\/ijcai.2018\/779"},{"issue":"1","key":"814_CR9","doi-asserted-by":"publisher","first-page":"263","DOI":"10.1007\/s00146-018-0871-3","volume":"35","author":"WA Bauer","year":"2020","unstructured":"Bauer, W.A.: Virtuous vs. utilitarian artificial moral agents. AI Soc 35(1), 263\u2013271 (2020). https:\/\/doi.org\/10.1007\/s00146-018-0871-3","journal-title":"AI Soc"},{"key":"814_CR10","unstructured":"Grosz, B.J. et al., \u201cEmbedded ethics: integrating ethics broadly across computer science education,\u201d arXiv:1808.05686 [cs], 2018, [Online]. Available: http:\/\/arxiv.org\/abs\/1808.05686"},{"key":"814_CR11","doi-asserted-by":"publisher","DOI":"10.1136\/bmj.n71","author":"MJ Page","year":"2021","unstructured":"Page, M.J., McKenzie, J.E., Bossuyt, P.M., Boutron, I., Hoffmann, T.C., Mulrow, C.D., Shamseer, L., Tetzlaff, J.M., Akl, E.A., Brennan, S.E., Chou, R.: The PRISMA 2020 statement an updated guideline for reporting systematic reviews. bmj (2021). https:\/\/doi.org\/10.1136\/bmj.n71","journal-title":"bmj"},{"key":"814_CR12","doi-asserted-by":"publisher","DOI":"10.1016\/j.ijinfomgt.2023.102700","volume":"74","author":"BC Stahl","year":"2024","unstructured":"Stahl, B.C., Eke, D.: The ethics of ChatGPT \u2013 exploring the ethical issues of an emerging technology. Int. J. Inf. Manage 74, 102700 (2024). https:\/\/doi.org\/10.1016\/j.ijinfomgt.2023.102700","journal-title":"Int. J. Inf. Manage"},{"key":"814_CR13","doi-asserted-by":"publisher","unstructured":"Spennemann, D.H.R., \u201cExploring ethical boundaries: can ChatGPT be prompted to give advice on how to cheat in university assignments?,\u201d 2023, https:\/\/doi.org\/10.20944\/preprints202308.1271.v1.","DOI":"10.20944\/preprints202308.1271.v1"},{"key":"814_CR14","doi-asserted-by":"crossref","unstructured":"Whittlestone, J. and Clarke, S., \u201cAI challenges for society and ethics,\u201d 2022, [Online]. Available: http:\/\/arxiv.org\/abs\/2206.11068","DOI":"10.1093\/oxfordhb\/9780197579329.013.3"},{"key":"814_CR15","doi-asserted-by":"publisher","unstructured":"Weidinger, L. et al., \u201cEthical and social risks of harm from Language Models,\u201d 2021, https:\/\/doi.org\/10.48550\/arXiv.2112.04359.","DOI":"10.48550\/arXiv.2112.04359"},{"key":"814_CR16","doi-asserted-by":"publisher","unstructured":"Yang, J. et al., \u201cHarnessing the power of LLMs in practice: a survey on ChatGPT and beyond,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2304.13712.","DOI":"10.48550\/arXiv.2304.13712"},{"key":"814_CR17","doi-asserted-by":"publisher","unstructured":"Wei, M. and Zhou, Z., \u201cAI ethics issues in real world: evidence from AI incident database,\u201d 2022, https:\/\/doi.org\/10.48550\/arXiv.2206.07635.","DOI":"10.48550\/arXiv.2206.07635"},{"key":"814_CR18","doi-asserted-by":"publisher","unstructured":"Cortiz, D. and Zubiaga, A., \u201cEthical and technical challenges of AI in tackling hate speech,\u201d The International Review of Information Ethics, vol. 29, 2020, https:\/\/doi.org\/10.29173\/irie416.","DOI":"10.29173\/irie416"},{"key":"814_CR19","doi-asserted-by":"publisher","unstructured":"Bender, E.M., Gebru, T., McMillan-Major, A., and Shmitchell, S., \u201cOn the dangers of stochastic parrots: can language models be too big?,\u201d pp. 610\u2013623, 2021, https:\/\/doi.org\/10.1145\/3442188.3445922.","DOI":"10.1145\/3442188.3445922"},{"issue":"10","key":"814_CR20","doi-asserted-by":"publisher","first-page":"60","DOI":"10.1080\/15265161.2023.2250292","volume":"23","author":"S Laacke","year":"2023","unstructured":"Laacke, S., Gauckler, C.: Why personalized large language models fail to do what ethics is all about. Am. J. Bioeth. 23(10), 60\u201363 (2023). https:\/\/doi.org\/10.1080\/15265161.2023.2250292","journal-title":"Am. J. Bioeth."},{"key":"814_CR21","doi-asserted-by":"publisher","unstructured":"Zelch, I., Hagen, M., and Potthast, M., \u201cCommercialized generative AI: a critical study of the feasibility and ethics of generating native advertising using large language models in conversational web search,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2310.04892.","DOI":"10.48550\/arXiv.2310.04892"},{"key":"814_CR22","doi-asserted-by":"publisher","unstructured":"Henderson, P. et al., \u201cEthical challenges in data-driven dialogue systems,\u201d pp. 123\u2013129, 2018, https:\/\/doi.org\/10.1145\/3278721.3278777.","DOI":"10.1145\/3278721.3278777"},{"key":"814_CR23","doi-asserted-by":"publisher","unstructured":"Bommasani, R. et al., \u201cOn the opportunities and risks of foundation models,\u201d 2022, https:\/\/doi.org\/10.48550\/arXiv.2108.07258.","DOI":"10.48550\/arXiv.2108.07258"},{"issue":"2","key":"814_CR24","doi-asserted-by":"publisher","first-page":"553","DOI":"10.1007\/s43681-022-00188-y","volume":"3","author":"T Hagendorff","year":"2023","unstructured":"Hagendorff, T., Danks, D.: Ethical and methodological challenges in building morally informed AI systems. AI and Ethics 3(2), 553\u2013566 (2023). https:\/\/doi.org\/10.1007\/s43681-022-00188-y","journal-title":"AI and Ethics"},{"key":"814_CR25","doi-asserted-by":"publisher","unstructured":"Obreja, D.M. and Rughini\u015f, R., \u201cThe moral status of artificial intelligence: exploring users\u2019 anticipatory ethics in the controversy regarding LaMDA\u2019s sentience,\u201d in 2023 24th International Conference on Control Systems and Computer Science (CSCS), pp. 411\u2013417, 2023, https:\/\/doi.org\/10.1109\/CSCS59211.2023.00071.","DOI":"10.1109\/CSCS59211.2023.00071"},{"key":"814_CR26","doi-asserted-by":"publisher","unstructured":"Datta, T. and Dickerson, J.P., \u201cWho\u2019s thinking? a push for human-centered evaluation of LLMs using the XAI playbook,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2303.06223.","DOI":"10.48550\/arXiv.2303.06223"},{"key":"814_CR27","doi-asserted-by":"publisher","unstructured":"Chen, J. et al., \u201cWhen large language models meet personalization: perspectives of challenges and opportunities,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2307.16376.","DOI":"10.48550\/arXiv.2307.16376"},{"key":"814_CR28","doi-asserted-by":"publisher","unstructured":"Liu, Y. et al., \u201cTrustworthy LLMs: a survey and guideline for evaluating large language models\u2019 alignment,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2308.05374.","DOI":"10.48550\/arXiv.2308.05374"},{"key":"814_CR29","doi-asserted-by":"publisher","unstructured":"Prabhumoye, S., Boldt, B., Salakhutdinov, R., and Black, A.W., \u201cCase study: deontological ethics in NLP,\u201d 2021, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2010.04658.","DOI":"10.48550\/arXiv.2010.04658"},{"issue":"25","key":"814_CR30","volume":"26","author":"E Fournier-Tombs","year":"2023","unstructured":"Fournier-Tombs, E., McHardy, J.: A medical ethics framework for conversational artificial intelligence. J. Med. Int. Res. 26(25), e43068 (2023)","journal-title":"J. Med. Int. Res."},{"key":"814_CR31","doi-asserted-by":"publisher","DOI":"10.1007\/s43681-022-00148-6","author":"A Chan","year":"2023","unstructured":"Chan, A.: GPT-3 and InstructGPT: technological dystopianism, utopianism, and \u2018contextual\u2019 perspectives in AI ethics and industry. AI Ethics (2023). https:\/\/doi.org\/10.1007\/s43681-022-00148-6","journal-title":"AI Ethics"},{"key":"814_CR32","doi-asserted-by":"publisher","unstructured":"Zhou, J. et al., \u201cRethinking machine ethics\u2014can LLMs perform moral reasoning through the lens of moral theories?,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2308.15399.","DOI":"10.48550\/arXiv.2308.15399"},{"key":"814_CR33","doi-asserted-by":"publisher","unstructured":"Goanta, C., Aletras, N., Chalkidis, I., Ranchordas, S., and Spanakis, G., \u201cRegulation and NLP (RegNLP): taming large language models,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2310.05553.","DOI":"10.48550\/arXiv.2310.05553"},{"key":"814_CR34","doi-asserted-by":"publisher","unstructured":"Kiritchenko, S. and Nejadgholi, I., \u201cTowards ethics by design in online abusive content detection,\u201d 2020, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2010.14952.","DOI":"10.48550\/arXiv.2010.14952"},{"key":"814_CR35","doi-asserted-by":"publisher","unstructured":"Leidner, J.L. and Plachouras, V., \u201cEthical by design: ethics best practices for natural language processing,\u201d in EthNLP 2017, D. Hovy, S. Spruit, M. Mitchell, E. M. Bender, M. Strube, and H. Wallach, Eds., Association for Computational Linguistics, 2017, pp. 30\u201340. https:\/\/doi.org\/10.18653\/v1\/W17-1604.","DOI":"10.18653\/v1\/W17-1604"},{"key":"814_CR36","doi-asserted-by":"publisher","DOI":"10.1007\/s43681-023-00309-1","author":"S Afroogh","year":"2023","unstructured":"Afroogh, S., et al.: Embedded ethics for responsible artificial intelligence systems (EE-RAIS) in disaster management: a conceptual model and its deployment. AI Ethics (2023). https:\/\/doi.org\/10.1007\/s43681-023-00309-1","journal-title":"AI Ethics"},{"issue":"2","key":"814_CR37","doi-asserted-by":"publisher","first-page":"553","DOI":"10.1007\/s43681-022-00188-y","volume":"3","author":"T Hagendorff","year":"2023","unstructured":"Hagendorff, T., Danks, D.: Ethical and methodological challenges in building morally informed AI systems. AI Ethics 3(2), 553\u2013566 (2023). https:\/\/doi.org\/10.1007\/s43681-022-00188-y","journal-title":"AI Ethics"},{"key":"814_CR38","doi-asserted-by":"publisher","unstructured":"Caliskan, A., \u201cArtificial intelligence, bias, and ethics,\u201d in Thirty-Second International Joint Conference on Artificial Intelligence {IJCAI-23}, International Joint Conferences on Artificial Intelligence Organization, 2023, pp. 7007\u20137013. https:\/\/doi.org\/10.24963\/ijcai.2023\/799.","DOI":"10.24963\/ijcai.2023\/799"},{"key":"814_CR39","doi-asserted-by":"publisher","unstructured":"Yang, C., Rustogi, R., Brower-Sinning, R., Lewis, G.A., K\u00e4stner, C., and Wu, T., \u201cBeyond testers\u2019 biases: guiding model testing with knowledge bases using LLMs,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2310.09668.","DOI":"10.48550\/arXiv.2310.09668"},{"issue":"8","key":"814_CR40","doi-asserted-by":"publisher","DOI":"10.3390\/socsci12080435","volume":"12","author":"N Gross","year":"2023","unstructured":"Gross, N.: What ChatGPT tells us about gender: a cautionary tale about performativity and gender biases in AI. Soc Sci 12(8), 435 (2023). https:\/\/doi.org\/10.3390\/socsci12080435","journal-title":"Soc Sci"},{"key":"814_CR41","doi-asserted-by":"publisher","unstructured":"Venkit, P.N., Gautam, S., Panchanadikar, R., \u201cKenneth\u201d Huang, T.-H., and Wilson, S., \u201cNationality bias in text generation,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2302.02463.","DOI":"10.48550\/arXiv.2302.02463"},{"key":"814_CR42","doi-asserted-by":"publisher","unstructured":"Fang, X., Che, S., Mao, M., Zhang, H., Zhao, M., and Zhao, X., \u201cBias of AI-generated content: an examination of news produced by large language models,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2309.09825.","DOI":"10.48550\/arXiv.2309.09825"},{"key":"814_CR43","doi-asserted-by":"publisher","unstructured":"Haller, P., Aynetdinov, A., and Akbik, A., \u201cOpinionGPT: modelling explicit biases in instruction-tuned LLMs,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2309.03876.","DOI":"10.48550\/arXiv.2309.03876"},{"key":"814_CR44","doi-asserted-by":"publisher","unstructured":"Dai, S. et al., \u201cLLMs may Dominate Information Access: Neural Retrievers are Biased Towards LLM-Generated Texts,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2310.20501.","DOI":"10.48550\/arXiv.2310.20501"},{"key":"814_CR45","doi-asserted-by":"publisher","unstructured":"Urman, A. and Makhortykh, M., \u201cThe silence of the LLMs: cross-lingual analysis of political bias and false information prevalence in ChatGPT, google bard, and bing chat,\u201d 2023, https:\/\/doi.org\/10.31219\/osf.io\/q9v8f.","DOI":"10.31219\/osf.io\/q9v8f"},{"key":"814_CR46","doi-asserted-by":"publisher","unstructured":"Wan, Y., Pu, G., Sun, J., Garimella, A., Chang, K.-W., and Peng, N., \u201c\u2018Kelly is a warm person, joseph is a role model\u2019: gender biases in LLM-generated reference letters,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2310.09219.","DOI":"10.48550\/arXiv.2310.09219"},{"key":"814_CR47","doi-asserted-by":"publisher","unstructured":"Salewski, L., Alaniz, S., Rio-Torto, I., Schulz, E., and Akata, Z., \u201cIn-context impersonation reveals large language models\u2019 strengths and biases,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2305.14930.","DOI":"10.48550\/arXiv.2305.14930"},{"key":"814_CR48","doi-asserted-by":"publisher","unstructured":"Kotek, H., Dockum, R., and Sun, D.Q., \u201cGender bias and stereotypes in large language models,\u201d pp. 12\u201324, 2023, https:\/\/doi.org\/10.1145\/3582269.3615599.","DOI":"10.1145\/3582269.3615599"},{"key":"814_CR49","doi-asserted-by":"publisher","unstructured":"Gallegos, I.O. et al., \u201cBias and fairness in large language models: a survey,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2309.00770.","DOI":"10.48550\/arXiv.2309.00770"},{"key":"814_CR50","doi-asserted-by":"publisher","unstructured":"Kamruzzaman, M., Shovon, M.M.I., and Kim, G.L., \u201cInvestigating subtler biases in LLMs: ageism, beauty, institutional, and nationality bias in generative models,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2309.08902.","DOI":"10.48550\/arXiv.2309.08902"},{"key":"814_CR51","doi-asserted-by":"publisher","unstructured":"Huang, D., Bu, Q., Zhang, J., Xie, X., Chen, J., and Cui, H., \u201cBias assessment and mitigation in LLM-based code generation,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2309.14345.","DOI":"10.48550\/arXiv.2309.14345"},{"key":"814_CR52","doi-asserted-by":"publisher","unstructured":"Li, Y., Du, M., Song, R., Wang, X., and Wang, Y., \u201cA survey on fairness in large language models,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2308.10149.","DOI":"10.48550\/arXiv.2308.10149"},{"key":"814_CR53","doi-asserted-by":"publisher","unstructured":"Li, Y. and Zhang, Y., \u201cFairness of ChatGPT,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2305.18569.","DOI":"10.48550\/arXiv.2305.18569"},{"key":"814_CR54","doi-asserted-by":"publisher","unstructured":"Zhang, J., Bao, K., Zhang, Y., Wang, W., Feng, F., and He, X., \u201cIs ChatGPT fair for recommendation? evaluating fairness in large language model recommendation,\u201d pp. 993\u2013999, 2023, https:\/\/doi.org\/10.1145\/3604915.3608860.","DOI":"10.1145\/3604915.3608860"},{"key":"814_CR55","doi-asserted-by":"publisher","unstructured":"Hua, W., Ge, Y., Xu, S., Ji, J., and Zhang, Y., \u201cUP5: unbiased foundation model for fairness-aware recommendation,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2305.12090.","DOI":"10.48550\/arXiv.2305.12090"},{"key":"814_CR56","doi-asserted-by":"publisher","unstructured":"Fryer, Z., Axelrod, V., Packer, B., Beutel, A., Chen, J., and Webster, K., \u201cFlexible text generation for counterfactual fairness probing,\u201d 2022, https:\/\/doi.org\/10.48550\/arXiv.2206.13757.","DOI":"10.48550\/arXiv.2206.13757"},{"key":"814_CR57","doi-asserted-by":"publisher","unstructured":"Deldjoo, Y., \u201cFairness of ChatGPT and the role of explainable-guided prompts,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2307.11761.","DOI":"10.48550\/arXiv.2307.11761"},{"key":"814_CR58","doi-asserted-by":"publisher","unstructured":"Liu, Y., Gautam, S., Ma, J., and Lakkaraju, H., \u201cInvestigating the fairness of large language models for predictions on tabular data,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2310.14607.","DOI":"10.48550\/arXiv.2310.14607"},{"key":"814_CR59","unstructured":"Wang, R., Cheng, P., and Henao, R., \u201cToward fairness in text generation via mutual information minimization based on importance sampling,\u201d International Conference on Artificial Intelligence and Statistics, pp. 4473\u20134485, 2023, [Online]. Available: https:\/\/proceedings.mlr.press\/v206\/wang23c.html"},{"key":"814_CR60","doi-asserted-by":"publisher","unstructured":"Ma, H. et al., \u201cFairness-guided few-shot prompting for large language models,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2303.13217.","DOI":"10.48550\/arXiv.2303.13217"},{"key":"814_CR61","doi-asserted-by":"publisher","unstructured":"Li, H. et al., \u201cMulti-step jailbreaking privacy attacks on ChatGPT,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2304.05197.","DOI":"10.48550\/arXiv.2304.05197"},{"key":"814_CR62","doi-asserted-by":"publisher","DOI":"10.1016\/j.jiixd.2023.10.007","author":"X Wu","year":"2023","unstructured":"Wu, X., Duan, R., Ni, J.: Unveiling security, privacy, and ethical concerns of ChatGPT. J. Inf. Intell. (2023). https:\/\/doi.org\/10.1016\/j.jiixd.2023.10.007","journal-title":"J. Inf. Intell."},{"key":"814_CR63","doi-asserted-by":"publisher","unstructured":"Li, Y., Tan, Z., and Liu, Y., \u201cPrivacy-preserving prompt tuning for large language model services,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2305.06212.","DOI":"10.48550\/arXiv.2305.06212"},{"key":"814_CR64","doi-asserted-by":"publisher","unstructured":"Carranza, A. G., R. Farahani, N. Ponomareva, A. Kurakin, M. Jagielski, and M. Nasr, \u201cPrivacy-preserving recommender systems with synthetic query generation using differentially private large language models,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2305.05973.","DOI":"10.48550\/arXiv.2305.05973"},{"key":"814_CR65","doi-asserted-by":"crossref","unstructured":"Khowaja, S.A., Khuwaja, P., and Dev, K., \u201cChatGPT NEEDS SPADE (Sustainability, PrivAcy, Digital divide, and Ethics) Evaluation: A Review,\u201d Apr. 2023, [Online]. Available: http:\/\/arxiv.org\/abs\/2305.03123","DOI":"10.36227\/techrxiv.22619932.v1"},{"key":"814_CR66","doi-asserted-by":"publisher","unstructured":"Mai, P., Yan, R., Huang, Z., Yang, Y., and Pang, Y., \u201cSplit-and-denoise: protect large language model inference with local differential privacy,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2310.09130.","DOI":"10.48550\/arXiv.2310.09130"},{"key":"814_CR67","doi-asserted-by":"publisher","unstructured":"Mireshghallah, F., Inan, H.A., Hasegawa, M., R\u00fchle, V., Berg-Kirkpatrick, T., and Sim, R., \u201cPrivacy regularization: joint privacy-utility optimization in language models,\u201d 2021, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2103.07567.","DOI":"10.48550\/arXiv.2103.07567"},{"key":"814_CR68","doi-asserted-by":"publisher","unstructured":"Raeini, M., \u201cPrivacy-preserving large language models (PPLLMs),\u201d 2023, Rochester, NY. https:\/\/doi.org\/10.2139\/ssrn.4512071.","DOI":"10.2139\/ssrn.4512071"},{"key":"814_CR69","doi-asserted-by":"publisher","unstructured":"Montagna, S., Ferretti, S., Klopfenstein, L.C., Florio, A., and Pengo, M.F., \u201cData decentralisation of LLM-based chatbot systems in chronic disease self-management,\u201d in GoodIT \u201923. Association for Computing Machinery, 2023, pp. 205\u2013212. https:\/\/doi.org\/10.1145\/3582515.3609536.","DOI":"10.1145\/3582515.3609536"},{"key":"814_CR70","doi-asserted-by":"publisher","unstructured":"Vats, A. et al., \u201cRecovering from privacy-preserving masking with large language models,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2309.08628.","DOI":"10.48550\/arXiv.2309.08628"},{"key":"814_CR71","doi-asserted-by":"publisher","unstructured":"Urman, A. and Makhortykh, M., \u201cThe silence of the LLMs: Cross-Lingual Analysis of Political Bias and False Information Prevalence in ChatGPT, Google Bard, and Bing Chat,\u201d 2023, OSF Preprints. https:\/\/doi.org\/10.31219\/osf.io\/q9v8f.","DOI":"10.31219\/osf.io\/q9v8f"},{"key":"814_CR72","doi-asserted-by":"publisher","unstructured":"Chen, C. and Shu, K., \u201cCan LLM-generated misinformation be detected?,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2309.13788.","DOI":"10.48550\/arXiv.2309.13788"},{"key":"814_CR73","doi-asserted-by":"publisher","unstructured":"Jiang, B., Tan, Z., Nirmal, A., and Liu, H., \u201cDisinformation detection: an evolving challenge in the age of LLMs,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2309.15847.","DOI":"10.48550\/arXiv.2309.15847"},{"key":"814_CR74","doi-asserted-by":"publisher","unstructured":"Leite, J.A., Razuvayevskaya, O., Bontcheva, K., and Scarton, C., \u201cDetecting misinformation with LLM-predicted credibility signals and weak supervision,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2309.07601.","DOI":"10.48550\/arXiv.2309.07601"},{"key":"814_CR75","doi-asserted-by":"publisher","unstructured":"Wu, J. and Hooi, B., \u201cFake news in sheep\u2019s clothing: robust fake news detection against LLM-empowered style attacks,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2310.10830.","DOI":"10.48550\/arXiv.2310.10830"},{"key":"814_CR76","unstructured":"Anonymous, \u201cThe earth is flat because...: investigating LLMs\u2019 belief towards misinformation via persuasive conversation,\u201d 2023, [Online]. Available: https:\/\/openreview.net\/forum?id=DJXifFF2_M"},{"key":"814_CR77","doi-asserted-by":"publisher","unstructured":"Guo, W. and Caliskan, A., \u201cDetecting emergent intersectional biases: contextualized word embeddings contain a distribution of human-like biases,\u201d in AIES \u201921. Association for Computing Machinery, 2021, pp. 122\u2013133. https:\/\/doi.org\/10.1145\/3461702.3462536.","DOI":"10.1145\/3461702.3462536"},{"key":"814_CR78","doi-asserted-by":"crossref","unstructured":"Choi, E.C. and Ferrara, E., \u201cAutomated claim matching with large language models: empowering fact-checkers in the fight against misinformation,\u201d Oct. 2023, [Online]. Available: http:\/\/arxiv.org\/abs\/2310.09223","DOI":"10.2139\/ssrn.4614239"},{"key":"814_CR79","doi-asserted-by":"publisher","unstructured":"Lin, S., Hilton, J., and Evans, O., \u201cTruthfulQA: measuring how models mimic human falsehoods,\u201d 2022, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2109.07958.","DOI":"10.48550\/arXiv.2109.07958"},{"key":"814_CR80","doi-asserted-by":"publisher","unstructured":"Lucas, J., Uchendu, A., Yamashita, M., Lee, J., Rohatgi, S., and Lee, D., \u201cFighting fire with fire: the dual role of LLMs in crafting and detecting elusive disinformation,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2310.15515.","DOI":"10.48550\/arXiv.2310.15515"},{"key":"814_CR81","doi-asserted-by":"publisher","unstructured":"Su, J., Zhuo, T.Y., Mansurov, J., Wang, D., and Nakov, P., \u201cFake news detectors are biased against texts generated by large language models,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2309.08674.","DOI":"10.48550\/arXiv.2309.08674"},{"key":"814_CR82","doi-asserted-by":"publisher","unstructured":"Yang, K.-C. and Menczer. F., \u201cLarge language models can rate news outlet credibility,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2304.00228.","DOI":"10.48550\/arXiv.2304.00228"},{"key":"814_CR83","unstructured":"Chen, C. and Shu, K., Combating misinformation in the age of LLMs: opportunities and challenges. 2023. [Online]. Available: files\/9663\/Chen and Shu\u20142023\u2013Combating Misinformation in the Age of LLMs Oppor.pdf"},{"key":"814_CR84","doi-asserted-by":"publisher","unstructured":"Xia, B., Lu, Q., Zhu, L., Lee, S. U., Liu, Y., and Xing, Z., \u201cFrom principles to practice: an accountability metrics catalogue for managing AI risks,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2311.13158.","DOI":"10.48550\/arXiv.2311.13158"},{"key":"814_CR85","doi-asserted-by":"publisher","unstructured":"He, K. et al., \u201cA survey of large language models for healthcare: from data, technology, and applications to accountability and ethics,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2310.05694.","DOI":"10.48550\/arXiv.2310.05694"},{"key":"814_CR86","doi-asserted-by":"crossref","unstructured":"Li, Y., Du, Y., Zhou, K., Wang, J., Zhao, W.X., and Wen, J.-R., \u201cEvaluating object hallucination in large vision-language models,\u201d 2023, arXiv. [Online]. Available: http:\/\/arxiv.org\/abs\/2305.10355","DOI":"10.18653\/v1\/2023.emnlp-main.20"},{"issue":"4","key":"814_CR87","doi-asserted-by":"publisher","first-page":"e37432","DOI":"10.7759\/cureus.37432","volume":"15","author":"SA Athaluri","year":"2023","unstructured":"Athaluri, S.A., Manthena, S.V., Kesapragada, V.S.R.K.M., Yarlagadda, V., Dave, T., Duddumpudi, R.T.S.: Exploring the boundaries of reality: investigating the phenomenon of artificial intelligence hallucination in scientific writing through ChatGPT references. Cureus 15(4), e37432 (2023). https:\/\/doi.org\/10.7759\/cureus.37432","journal-title":"Cureus"},{"key":"814_CR88","doi-asserted-by":"publisher","unstructured":"Deroy, A., Ghosh, K., and Ghosh, S., \u201cHow ready are pre-trained abstractive models and LLMs for legal case judgement summarization?,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2306.01248.","DOI":"10.48550\/arXiv.2306.01248"},{"key":"814_CR89","doi-asserted-by":"publisher","first-page":"100021","DOI":"10.1016\/j.hfh.2022.100021","volume":"2","author":"A Choudhury","year":"2022","unstructured":"Choudhury, A., Asan, O.: Impact of accountability, training, and human factors on the use of artificial intelligence in healthcare: exploring the perceptions of healthcare practitioners in the us. Human Factors Healthcare 2, 100021 (2022)","journal-title":"Human Factors Healthcare"},{"key":"814_CR90","doi-asserted-by":"publisher","unstructured":"Rogers, A., \u201cChanging the world by changing the data,\u201d Proceedings of the 59th Annual Meeting of the Association for Computational Linguistics and the 11th International Joint Confer, 2021, https:\/\/doi.org\/10.18653\/v1\/2021.acl-long.170.","DOI":"10.18653\/v1\/2021.acl-long.170"},{"issue":"2","key":"814_CR91","doi-asserted-by":"publisher","DOI":"10.1016\/j.patter.2021.100205","volume":"2","author":"A Birhane","year":"2021","unstructured":"Birhane, A.: Algorithmic injustice: a relational ethics approach. Patterns 2(2), 100205 (2021). https:\/\/doi.org\/10.1016\/j.patter.2021.100205","journal-title":"Patterns"},{"issue":"11","key":"814_CR92","doi-asserted-by":"publisher","DOI":"10.1016\/j.patter.2021.100336","volume":"2","author":"A Paullada","year":"2021","unstructured":"Paullada, A., Raji, I.D., Bender, E.M., Denton, E., Hanna, A.: Data and its (dis)contents: a survey of dataset development and use in machine learning research. Patterns 2(11), 100336 (2021). https:\/\/doi.org\/10.1016\/j.patter.2021.100336","journal-title":"Patterns"},{"key":"814_CR93","doi-asserted-by":"publisher","unstructured":"Crisan, A., Drouhard, M., Vig, J., & Rajani, N, \u201cInteractive model cards: a human-centered approach to model documentation,\u201d In 2022 ACM Conference on Fairness, Accountability, and Transparency (FAccT \u201922). Association for Computing Machinery, New York, NY, USA, 427\u2013439. https: \/\/doi.org\/https:\/\/doi.org\/10.1145\/3531146.3533108 , 2022.","DOI":"10.1145\/3531146.3533108"},{"key":"814_CR94","doi-asserted-by":"publisher","unstructured":"Liesenfeld, A., Lopez, A., and Dingemanse, M., \u201cOpening up ChatGPT: tracking openness, transparency, and accountability in instruction-tuned text generators,\u201d in CUI \u201923. Association for Computing Machinery, 2023, pp. 1\u20136. https:\/\/doi.org\/10.1145\/3571884.3604316.","DOI":"10.1145\/3571884.3604316"},{"key":"814_CR95","doi-asserted-by":"publisher","unstructured":"Gudibande, A. et al., \u201cThe false promise of imitating proprietary LLMs,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2305.15717.","DOI":"10.48550\/arXiv.2305.15717"},{"key":"814_CR96","unstructured":"Mao, R., Chen, G., Zhang, X., Guerin, F., and Cambria, E., GPTEval: a survey on assessments of ChatGPT and GPT-4. 2023. [Online]. Available: files\/10027\/Mao et al.\u20132023\u2013GPTEval A Survey on Assessments of ChatGPT and GP.pdf"},{"key":"814_CR97","doi-asserted-by":"publisher","unstructured":"Huang, J. and Chen-Chuan Chang, K., \u201cCitation: a key to building responsible and accountable large language models,\u201d 2023. https:\/\/doi.org\/10.48550\/arXiv.2307.02185.","DOI":"10.48550\/arXiv.2307.02185"},{"key":"814_CR98","doi-asserted-by":"publisher","unstructured":"Guo, E. et al., \u201cneuroGPT-X: towards an accountable expert opinion tool for vestibular schwannoma,\u201d vol. 1, 2023, https:\/\/doi.org\/10.17632\/b9mck42r35.1.","DOI":"10.17632\/b9mck42r35.1"},{"key":"814_CR99","doi-asserted-by":"publisher","unstructured":"Khan, M., Hanna, A., \u201cThe subjects and stages of AI dataset development: a framework for dataset accountability,\u201d Forthcoming 19 Ohio St. Tech. L.J. (2023), Available at SSRN:https:\/\/ssrn.com\/abstract=4217148 or https:\/\/doi.org\/10.2139\/ssrn.4217148, 2022.","DOI":"10.2139\/ssrn.4217148"},{"key":"814_CR100","unstructured":"Cacciamani, G.E. et al., \u201cDevelopment of the ChatGPT, generative artificial intelligence and natural large language models for accountable reporting and use (CANGARU) guidelines,\u201d 2023, arXiv. [Online]. Available: http:\/\/arxiv.org\/abs\/2307.08974"},{"key":"814_CR101","unstructured":"Anderljung, M.E.T.S.J.O.L.S.B.B.E.B.J.S.R.T.L.S., \u201cTowards publicly accountable frontier LLMs: building an external scrutiny ecosystem under the ASPIRE framework.,\u201d arXiv preprintarXiv:2311.14711, 2023."},{"key":"814_CR102","doi-asserted-by":"publisher","unstructured":"Nabben, K., \u201cConstituting an AI: accountability lessons from an LLM experiment,\u201d 2023, Rochester, NY. https:\/\/doi.org\/10.2139\/ssrn.4561433.","DOI":"10.2139\/ssrn.4561433"},{"key":"814_CR103","doi-asserted-by":"publisher","DOI":"10.1136\/jme-2023-109347","author":"JW Allen","year":"2023","unstructured":"Allen, J.W., Earp, B.D., Koplin, J., Wilkinson, D.: Consent-GPT: is it ethical to delegate procedural consent to conversational AI? J. Med. Ethics (2023). https:\/\/doi.org\/10.1136\/jme-2023-109347","journal-title":"J. Med. Ethics"},{"key":"814_CR104","doi-asserted-by":"publisher","DOI":"10.7759\/cureus.43262","author":"M Jeyaraman","year":"2023","unstructured":"Jeyaraman, M., Balaji, S., Jeyaraman, N., Yadav, S.: Unraveling the ethical enigma: artificial intelligence in healthcare. Cureus (2023). https:\/\/doi.org\/10.7759\/cureus.43262","journal-title":"Cureus"},{"issue":"10","key":"814_CR105","doi-asserted-by":"publisher","first-page":"17","DOI":"10.1080\/15265161.2023.2233358","volume":"23","author":"V Rahimzadeh","year":"2023","unstructured":"Rahimzadeh, V., Kostick-Quenet, K., Blumenthal Barby, J., McGuire, A.L.: Ethics education for healthcare professionals in the era of ChatGPT and other large language models: do we still need it? Am. J. Bioeth. 23(10), 17\u201327 (2023). https:\/\/doi.org\/10.1080\/15265161.2023.2233358","journal-title":"Am. J. Bioeth."},{"issue":"6","key":"814_CR106","doi-asserted-by":"publisher","first-page":"e333","DOI":"10.1016\/S2589-7500(23)00083-3","volume":"5","author":"H Li","year":"2023","unstructured":"Li, H., Moon, J.T., Purkayastha, S., Celi, L.A., Trivedi, H., Gichoya, J.W.: Ethics of large language models in medicine and medical research. Lancet Digit Health 5(6), e333\u2013e335 (2023). https:\/\/doi.org\/10.1016\/S2589-7500(23)00083-3","journal-title":"Lancet Digit Health"},{"issue":"1","key":"814_CR107","doi-asserted-by":"publisher","first-page":"e48009","DOI":"10.2196\/48009","volume":"25","author":"C Wang","year":"2023","unstructured":"Wang, C., Liu, S., Yang, H., Guo, J., Wu, Y., Liu, J.: Ethical considerations of using ChatGPT in health care. J. Med. Internet Res. 25(1), e48009 (2023). https:\/\/doi.org\/10.2196\/48009","journal-title":"J. Med. Internet Res."},{"key":"814_CR108","doi-asserted-by":"publisher","first-page":"71","DOI":"10.1016\/j.neuroscience.2023.02.008","volume":"515","author":"A Graf","year":"2023","unstructured":"Graf, A., Bernardi, R.E.: ChatGPT in research: balancing ethics, transparency and advancement. Neuroscience 515, 71\u201373 (2023). https:\/\/doi.org\/10.1016\/j.neuroscience.2023.02.008","journal-title":"Neuroscience"},{"issue":"5","key":"814_CR109","doi-asserted-by":"publisher","first-page":"570","DOI":"10.1002\/asi.24750","volume":"74","author":"BD Lund","year":"2023","unstructured":"Lund, B.D., Wang, T., Mannuru, N.R., Nie, B., Shimray, S., Wang, Z.: ChatGPT and a new academic reality: artificial Intelligence-written research papers and the ethics of the large language models in scholarly publishing. J. Assoc. Inf. Sci. Technol. 74(5), 570\u2013581 (2023). https:\/\/doi.org\/10.1002\/asi.24750","journal-title":"J. Assoc. Inf. Sci. Technol."},{"key":"814_CR110","doi-asserted-by":"publisher","first-page":"17","DOI":"10.3354\/esep00195","volume":"21","author":"N Dehouche","year":"2021","unstructured":"Dehouche, N.: Plagiarism in the age of massive generative pre-trained transformers (GPT-3). Ethics Sci. Environ. Polit. 21, 17\u201323 (2021). https:\/\/doi.org\/10.3354\/esep00195","journal-title":"Ethics Sci. Environ. Polit."},{"key":"814_CR111","doi-asserted-by":"publisher","DOI":"10.5125\/jkaoms.2023.49.3.105","author":"JY Park","year":"2023","unstructured":"Park, J.Y.: Could ChatGPT help you to write your next scientific paper?: concerns on research ethics related to usage of artificial intelligence tools. J. Korean Assoc. Oral Maxillofac. Surg. (2023). https:\/\/doi.org\/10.5125\/jkaoms.2023.49.3.105","journal-title":"J. Korean Assoc. Oral Maxillofac. Surg."},{"key":"814_CR112","doi-asserted-by":"publisher","unstructured":"Schintler, L.A., McNeely, C.L., and Witte, J., \u201cA critical examination of the ethics of AI-mediated peer review,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2309.12356.","DOI":"10.48550\/arXiv.2309.12356"},{"key":"814_CR113","doi-asserted-by":"publisher","unstructured":"McKee, H.A. and Porter, J.E., \u201cEthics for AI writing: the importance of rhetorical context,\u201d in AIES \u201920. Association for Computing Machinery, 2020, pp. 110\u2013116. https:\/\/doi.org\/10.1145\/3375627.3375811.","DOI":"10.1145\/3375627.3375811"},{"key":"814_CR114","doi-asserted-by":"publisher","unstructured":"Lindemann, N.F., \u201cSealed knowledges: a critical approach to the usage of LLMs as search engines,\u201d in AIES \u201923. Association for Computing Machinery, 2023, pp. 985\u2013986. https:\/\/doi.org\/10.1145\/3600211.3604737.","DOI":"10.1145\/3600211.3604737"},{"key":"814_CR115","doi-asserted-by":"publisher","unstructured":"Spennemann, D.H.R., \u201cExploring ethical boundaries: can ChatGPT be prompted to give advice on how to cheat in university assignments?,\u201d 2023, Preprints. https:\/\/doi.org\/10.20944\/preprints202308.1271.v1.","DOI":"10.20944\/preprints202308.1271.v1"},{"key":"814_CR116","doi-asserted-by":"publisher","DOI":"10.1016\/j.lindif.2023.102274","volume":"103","author":"E Kasneci","year":"2023","unstructured":"Kasneci, E., et al.: ChatGPT for good? On opportunities and challenges of large language models for education. Learning and Individual Differences 103, 102274 (2023). https:\/\/doi.org\/10.1016\/j.lindif.2023.102274","journal-title":"Learning and Individual Differences"},{"issue":"10","key":"814_CR117","doi-asserted-by":"publisher","first-page":"28","DOI":"10.1080\/15265161.2023.2233356","volume":"23","author":"S Porsdam Mann","year":"2023","unstructured":"Porsdam Mann, S., Earp, B.D., M\u00f8ller, N., Vynn, S., Savulescu, J.: Autogen: a personalized large language model for academic enhancement\u2014ethics and proof of principle. Am. J. Bioeth. 23(10), 28\u201341 (2023). https:\/\/doi.org\/10.1080\/15265161.2023.2233356","journal-title":"Am. J. Bioeth."},{"key":"814_CR118","doi-asserted-by":"publisher","unstructured":"Mhlanga, D., \u201cOpen AI in education, the responsible and ethical use of ChatGPT towards lifelong learning,\u201d 2023, Rochester, NY. https:\/\/doi.org\/10.2139\/ssrn.4354422.","DOI":"10.2139\/ssrn.4354422"},{"key":"814_CR119","unstructured":"Shakir, U., Hess, J.L., James, M., and Katz, A., \u201cPushing ethics assessment forward in engineering: NLP-assisted qualitative coding of student responses,\u201d in 2023 ASEE Annual Conference & Exposition, 2023. [Online]. Available: https:\/\/peer.asee.org\/pushing-ethics-assessment-forward-in-engineering-nlp-assisted-qualitative-coding-of-student-responses"},{"issue":"1","key":"814_CR120","doi-asserted-by":"publisher","first-page":"1239","DOI":"10.33395\/jmp.v12i1.12693","volume":"12","author":"A Basir","year":"2023","unstructured":"Basir, A., Puspitasari, E.D., Aristarini, C.C., Sulastri, P.D., Ausat, A.M.A.: Ethical use of ChatGPT in the context of leadership and strategic decisions. Jurnal Minfo Polgan 12(1), 1239\u20131246 (2023). https:\/\/doi.org\/10.33395\/jmp.v12i1.12693","journal-title":"Jurnal Minfo Polgan"},{"key":"814_CR121","doi-asserted-by":"publisher","DOI":"10.1007\/s00146-022-01430-1","author":"M Ryan","year":"2022","unstructured":"Ryan, M., Christodoulou, E., Antoniou, J., Iordanou, K.: An AI ethics \u2018David and Goliath\u2019: value conflicts between large tech companies and their employees. AI & SOCIETY (2022). https:\/\/doi.org\/10.1007\/s00146-022-01430-1","journal-title":"AI & SOCIETY"},{"issue":"15","key":"814_CR122","doi-asserted-by":"publisher","first-page":"13470","DOI":"10.1609\/aaai.v35i15.17589","volume":"35","author":"N Lourie","year":"2021","unstructured":"Lourie, N., Le Bras, R., Choi, Y.: Scruples: a corpus of community ethical judgments on 32,000 real-life anecdotes. Proc. AAAI Conf. Artif. Intell. 35(15), 13470\u201313479 (2021). https:\/\/doi.org\/10.1609\/aaai.v35i15.17589","journal-title":"Proc. AAAI Conf. Artif. Intell."},{"key":"814_CR123","doi-asserted-by":"publisher","unstructured":"Cabrera, J., Loyola, M.S., Maga\u00f1a, I., and Rojas, R., \u201cEthical dilemmas, mental health, artificial intelligence, and LLM-based chatbots,\u201d I. Rojas, O. Valenzuela, F. Rojas Ruiz, L. J. Herrera, and F. Ortu\u00f1o, Eds., in Lecture Notes in Computer Science. Springer Nature Switzerland, 2023, pp. 313\u2013326. https:\/\/doi.org\/10.1007\/978-3-031-34960-7_22.","DOI":"10.1007\/978-3-031-34960-7_22"},{"key":"814_CR124","doi-asserted-by":"publisher","unstructured":"Ashurst, C., Hine, E., Sedille, P., and Carlier, A., \u201cAI ethics statements: analysis and lessons learnt from NeurIPS broader impact statements,\u201d in FAccT \u201922. Association for Computing Machinery, 2022, pp. 2047\u20132056. https:\/\/doi.org\/10.1145\/3531146.3533780.","DOI":"10.1145\/3531146.3533780"},{"issue":"0","key":"814_CR125","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1080\/14746700.2023.2255944","volume":"0","author":"A-R Bhojani","year":"2023","unstructured":"Bhojani, A.-R., Schwarting, M.: Truth and regret: large language models, the Quran, and misinformation. Theol. Sci. 0(0), 1\u20137 (2023). https:\/\/doi.org\/10.1080\/14746700.2023.2255944","journal-title":"Theol. Sci."},{"key":"814_CR126","doi-asserted-by":"publisher","unstructured":"Su, Z. et al., \u201cInfoEntropy loss to mitigate bias of learning difficulties for generative language models,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2310.19531.","DOI":"10.48550\/arXiv.2310.19531"},{"key":"814_CR127","doi-asserted-by":"publisher","unstructured":"Yu, P. and Ji, H., \u201cSelf information update for large language models through mitigating exposure bias,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2305.18582.","DOI":"10.48550\/arXiv.2305.18582"},{"key":"814_CR128","unstructured":"Ernst, J.S. et al., \u201cBias mitigation for large language models using adversarial learning\u201d, [Online]. Available: files\/9597\/Ernst et al. - Bias Mitigation for Large Language Models using Ad.pdf"},{"key":"814_CR129","doi-asserted-by":"publisher","unstructured":"Xue, M. et al., \u201cOccuQuest: mitigating occupational bias for inclusive large language models,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2310.16517.","DOI":"10.48550\/arXiv.2310.16517"},{"key":"814_CR130","doi-asserted-by":"publisher","unstructured":"Zhang, Z., Lyu, L., Ma, X., Wang, C., and Sun, X., \u201cFine-mixing: mitigating backdoors in fine-tuned language models,\u201d 2022, https:\/\/doi.org\/10.48550\/arXiv.2210.09545.","DOI":"10.48550\/arXiv.2210.09545"},{"key":"814_CR131","doi-asserted-by":"publisher","unstructured":"Jang, J. et al., \u201cKnowledge Unlearning for Mitigating Privacy Risks in Language Models,\u201d 2022, https:\/\/doi.org\/10.48550\/arXiv.2210.01504.","DOI":"10.48550\/arXiv.2210.01504"},{"key":"814_CR132","doi-asserted-by":"publisher","unstructured":"He, Z., Deng, H., Zhao, H., Liu, N., and Du, M., \u201cMitigating shortcuts in language models with soft label encoding,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2309.09380.","DOI":"10.48550\/arXiv.2309.09380"},{"key":"814_CR133","doi-asserted-by":"publisher","unstructured":"Thakur, H., Jain, A., Vaddamanu, P., Liang, P.P., and Morency, L.-P., \u201cLanguage models get a gender makeover: mitigating gender bias with few-shot data interventions,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2306.04597.","DOI":"10.48550\/arXiv.2306.04597"},{"key":"814_CR134","doi-asserted-by":"publisher","unstructured":"Dolci, T., \u201cFine-tuning language models to mitigate gender bias in sentence encoders,\u201d 2022 IEEE Eighth International Conference on Big Data Computing Service and Applications (BigDataService), pp. 175\u2013176, 2022, https:\/\/doi.org\/10.1109\/BigDataService55688.2022.00036.","DOI":"10.1109\/BigDataService55688.2022.00036"},{"key":"814_CR135","doi-asserted-by":"publisher","unstructured":"Zhao, J., Fang, M., Shi, Z., Li, Y., Chen, L., and Pechenizkiy, M., \u201cCHBias: bias evaluation and mitigation of chinese conversational language models,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2305.11262.","DOI":"10.48550\/arXiv.2305.11262"},{"key":"814_CR136","doi-asserted-by":"publisher","unstructured":"Omrani, A. et al., \u201cSocial-group-agnostic bias mitigation via the stereotype content model,\u201d in Proceedings of the 61st Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers), Association for Computational Linguistics, 2023, pp. 4123\u20134139. https:\/\/doi.org\/10.18653\/v1\/2023.acl-long.227.","DOI":"10.18653\/v1\/2023.acl-long.227"},{"key":"814_CR137","doi-asserted-by":"publisher","unstructured":"Gupta, U. et al., \u201cMitigating gender bias in distilled language models via counterfactual role reversal,\u201d 2022, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2203.12574.","DOI":"10.48550\/arXiv.2203.12574"},{"key":"814_CR138","doi-asserted-by":"publisher","DOI":"10.1145\/3628602","author":"M Bozdag","year":"2023","unstructured":"Bozdag, M., Sevim, N., Ko\u00e7, A.: Measuring and mitigating gender bias in legal contextualized language models. ACM Trans. Knowl. Discov. Data (2023). https:\/\/doi.org\/10.1145\/3628602","journal-title":"ACM Trans. Knowl. Discov. Data"},{"key":"814_CR139","doi-asserted-by":"publisher","unstructured":"Varshney, N., Yao, W., Zhang, H., Chen, J., and Yu, D., \u201cA stitch in time saves nine: detecting and mitigating hallucinations of LLMs by validating low-confidence generation,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2307.03987.","DOI":"10.48550\/arXiv.2307.03987"},{"key":"814_CR140","doi-asserted-by":"crossref","unstructured":"Lee, H., Hong, S., Park, J., Kim, T., Kim, G., and Ha, J., \u201c[Industry] KoSBI: a dataset for mitigating social bias risks towards safer large language model applications,\u201d Proceedings of the 61st Annual Meeting of the Association for Computational Linguistics, 2023, [Online]. Available: https:\/\/virtual2023.aclweb.org\/paper_I55.html","DOI":"10.18653\/v1\/2023.acl-industry.21"},{"key":"814_CR141","doi-asserted-by":"publisher","unstructured":"Dong, X., Zhu, Z., Wang, Z., Teleki, M., and Caverlee, J., \u201cCo$^2$PT: mitigating bias in pre-trained language models through counterfactual contrastive prompt tuning,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2310.12490.","DOI":"10.48550\/arXiv.2310.12490"},{"key":"814_CR142","doi-asserted-by":"publisher","unstructured":"Huang, D., Bu, Q., Zhang, J., Xie, X., Chen, J., and Cui, H., \u201cBias assessment and mitigation in LLM-based code generation,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2309.14345.","DOI":"10.48550\/arXiv.2309.14345"},{"key":"814_CR143","doi-asserted-by":"publisher","unstructured":"Steed, R., Panda, S., Kobren, A., and Wick, M., \u201cUpstream mitigation is not all you need: testing the bias transfer hypothesis in pre-trained language models,\u201d in ACL 2022, S. Muresan, P. Nakov, and A. Villavicencio, Eds., Association for Computational Linguistics, 2022, pp. 3524\u20133542. https:\/\/doi.org\/10.18653\/v1\/2022.acl-long.247.","DOI":"10.18653\/v1\/2022.acl-long.247"},{"key":"814_CR144","doi-asserted-by":"publisher","unstructured":"Ngo, H. et al., \u201cMitigating harm in language models with conditional-likelihood filtration,\u201d 2021, https:\/\/doi.org\/10.48550\/arXiv.2108.07790.","DOI":"10.48550\/arXiv.2108.07790"},{"key":"814_CR145","doi-asserted-by":"publisher","first-page":"490","DOI":"10.2197\/ipsjjip.29.490","volume":"29","author":"S Moon","year":"2021","unstructured":"Moon, S., Okazaki, N.: Effects and mitigation of out-of-vocabulary in universal language models. J. Inf. Process. 29, 490\u2013503 (2021). https:\/\/doi.org\/10.2197\/ipsjjip.29.490","journal-title":"J. Inf. Process."},{"key":"814_CR146","doi-asserted-by":"publisher","unstructured":"Lu, J. et al., \u201cEvaluation and mitigation of agnosia in multimodal large language models,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2309.04041.","DOI":"10.48550\/arXiv.2309.04041"},{"key":"814_CR147","doi-asserted-by":"publisher","unstructured":"Van, H., \u201cMitigating data scarcity for large language models,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2302.01806.","DOI":"10.48550\/arXiv.2302.01806"},{"key":"814_CR148","doi-asserted-by":"publisher","unstructured":"Viswanath, H. and Zhang, T., \u201cFairPy: a toolkit for evaluation of social biases and their mitigation in large language models,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2302.05508.","DOI":"10.48550\/arXiv.2302.05508"},{"issue":"1","key":"814_CR149","doi-asserted-by":"publisher","DOI":"10.1016\/j.ipm.2022.103139","volume":"60","author":"T Shen","year":"2023","unstructured":"Shen, T., Li, J., Bouadjenek, M.R., Mai, Z., Sanner, S.: Towards understanding and mitigating unintended biases in language model-driven conversational recommendation. Inf. Process. Manag. 60(1), 103139 (2023). https:\/\/doi.org\/10.1016\/j.ipm.2022.103139","journal-title":"Inf. Process. Manag."},{"key":"814_CR150","doi-asserted-by":"publisher","unstructured":"Ji, Z., Yu, T., Y. Xu, N. Lee, E. Ishii, and P. Fung, \u201cTowards mitigating hallucination in large language models via self-reflection,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2310.06271.","DOI":"10.48550\/arXiv.2310.06271"},{"key":"814_CR151","doi-asserted-by":"publisher","unstructured":"Leiser, F., Eckhardt, S., Knaeble, M., Maedche, A., Schwabe, G., and Sunyaev, A., \u201cFrom ChatGPT to FactGPT: A Participatory Design Study to Mitigate the Effects of Large Language Model Hallucinations on Users,\u201d pp. 81\u201390, 2023, https:\/\/doi.org\/10.1145\/3603555.3603565.","DOI":"10.1145\/3603555.3603565"},{"key":"814_CR152","doi-asserted-by":"publisher","unstructured":"Mahabadi, R.K., Belinkov, Y., and Henderson, J., \u201cEnd-to-end bias mitigation by modelling biases in Corpora,\u201d 2020, https:\/\/doi.org\/10.48550\/arXiv.1909.06321.","DOI":"10.48550\/arXiv.1909.06321"},{"key":"814_CR153","doi-asserted-by":"publisher","unstructured":"Garimella, A. et al., \u201cHe is very intelligent, she is very beautiful? on mitigating social biases in language modelling and generation,\u201d in Findings 2021, C. Zong, F. Xia, W. Li, and R. Navigli, Eds., Association for Computational Linguistics, 2021, pp. 4534\u20134545. https:\/\/doi.org\/10.18653\/v1\/2021.findings-acl.397.","DOI":"10.18653\/v1\/2021.findings-acl.397"},{"key":"814_CR154","doi-asserted-by":"publisher","unstructured":"Ungless, E.L., Rafferty, A., Nag, H., and Ross, B., \u201cA robust bias mitigation procedure based on the stereotype content model,\u201d 2022, https:\/\/doi.org\/10.48550\/arXiv.2210.14552.","DOI":"10.48550\/arXiv.2210.14552"},{"key":"814_CR155","unstructured":"Martin, K., Ethics of data and analytics: concepts and cases. CRC Press, 2022. [Online]. Available: https:\/\/books.google.com\/books?id=E51kEAAAQBAJ"},{"key":"814_CR156","doi-asserted-by":"publisher","unstructured":"Liu, Z., Zhang, X., and Peng, F., \u201cMitigating unintended memorization in language models via alternating teaching,\u201d in ICASSP 2023 - 2023 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), 2023, pp. 1\u20135. https:\/\/doi.org\/10.1109\/ICASSP49357.2023.10096557.","DOI":"10.1109\/ICASSP49357.2023.10096557"},{"issue":"17","key":"814_CR157","doi-asserted-by":"publisher","first-page":"14857","DOI":"10.1609\/aaai.v35i17.17744","volume":"35","author":"R Liu","year":"2021","unstructured":"Liu, R., Jia, C., Wei, J., Xu, G., Wang, L., Vosoughi, S.: Mitigating political bias in language models through reinforced calibration. Proce. AAAI Confer. Artif. Intell. 35(17), 14857\u201314866 (2021). https:\/\/doi.org\/10.1609\/aaai.v35i17.17744","journal-title":"Proce. AAAI Confer. Artif. Intell."},{"key":"814_CR158","doi-asserted-by":"publisher","unstructured":"Jin, X., Barbieri, F., Kennedy, B., Davani, A.M., Neves, L., and Ren, X., \u201cOn transferability of bias mitigation effects in language model fine-tuning,\u201d 2021, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2010.12864.","DOI":"10.48550\/arXiv.2010.12864"},{"key":"814_CR159","doi-asserted-by":"publisher","unstructured":"Liao, Q.V. and Vaughan, J.W., \u201cAI transparency in the age of LLMs: a human-centered research roadmap,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2306.01941.","DOI":"10.48550\/arXiv.2306.01941"},{"key":"814_CR160","doi-asserted-by":"publisher","unstructured":"Wu, T., Terry, M., and Cai, C.J., \u201cAI chains: transparent and controllable human-AI interaction by chaining large language model prompts,\u201d pp. 1\u201322, 2022, https:\/\/doi.org\/10.1145\/3491102.3517582.","DOI":"10.1145\/3491102.3517582"},{"issue":"3","key":"814_CR161","doi-asserted-by":"publisher","first-page":"224","DOI":"10.1111\/1753-0407.13361","volume":"15","author":"N Musacchio","year":"2023","unstructured":"Musacchio, N., et al.: Transparent machine learning suggests a key driver in the decision to start insulin therapy in individuals with type 2 diabetes. J. Diabetes 15(3), 224\u2013236 (2023). https:\/\/doi.org\/10.1111\/1753-0407.13361","journal-title":"J. Diabetes"},{"key":"814_CR162","doi-asserted-by":"publisher","unstructured":"Huang, Z., Gutierrez, S., Kamana, H., and Macneil, S., \u201cMemory sandbox: transparent and interactive memory management for conversational agents,\u201d in UIST \u201923 Adjunct. Association for Computing Machinery, 2023, pp. 1\u20133. https:\/\/doi.org\/10.1145\/3586182.3615796.","DOI":"10.1145\/3586182.3615796"},{"key":"814_CR163","doi-asserted-by":"publisher","unstructured":"Bang, J., Lee, B.-T., and Park, P., \u201cExamination of ethical principles for LLM-based recommendations in conversational AI,\u201d in 2023 International Conference on Platform Technology and Service (PlatCon), 2023, pp. 109\u2013113. https:\/\/doi.org\/10.1109\/PlatCon60102.2023.10255221.","DOI":"10.1109\/PlatCon60102.2023.10255221"},{"key":"814_CR164","doi-asserted-by":"publisher","unstructured":"Liesenfeld, A., Lopez, A., and Dingemanse, M., \u201cOpening up ChatGPT: Tracking openness, transparency, and accountability in instruction-tuned text generators,\u201d pp. 1\u20136, 2023, https:\/\/doi.org\/10.1145\/3571884.3604316.","DOI":"10.1145\/3571884.3604316"},{"key":"814_CR165","doi-asserted-by":"publisher","unstructured":"Glukhov, D., Shumailov, I., Gal, Y., Papernot, N., and Papyan, V., \u201cLLM censorship: a machine learning challenge or a computer security problem?,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2307.10719.","DOI":"10.48550\/arXiv.2307.10719"},{"key":"814_CR166","doi-asserted-by":"crossref","unstructured":"Karamolegkou, A., Li, J., Zhou, L., and S\u00f8gaard, A., \u201cCopyright violations and large language models,\u201d Oct. 2023, [Online]. Available: http:\/\/arxiv.org\/abs\/2310.13771","DOI":"10.18653\/v1\/2023.emnlp-main.458"},{"key":"814_CR167","unstructured":"Rahman, N. and Santacana, E., \u201cBeyond fair use: legal risk evaluation for training LLMs on copyrighted text\u201d, [Online]. Available: files\/9726\/Rahman and Santacana - Beyond Fair Use Legal Risk Evaluation for Trainin.pdf"},{"key":"814_CR168","doi-asserted-by":"crossref","unstructured":"Peng, W. et al., \u201cAre you copying my model? protecting the copyright of large language models for EaaS via backdoor watermark,\u201d May 2023, [Online]. Available: http:\/\/arxiv.org\/abs\/2305.10036","DOI":"10.18653\/v1\/2023.acl-long.423"},{"key":"814_CR169","unstructured":"Chu, T., Song, Z., and Yang, C., \u201cHow to protect copyright data in optimization of large language models?,\u201d Aug. 2023, [Online]. Available: http:\/\/arxiv.org\/abs\/2308.12247"},{"key":"814_CR170","doi-asserted-by":"publisher","unstructured":"Liu, Y., Hu, H., Chen, X., Zhang, X., and Sun, L., \u201cWatermarking classification dataset for copyright protection,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2305.13257.","DOI":"10.48550\/arXiv.2305.13257"},{"key":"814_CR171","doi-asserted-by":"publisher","unstructured":"Waidelich, L., Lambert, M., Al-Washash, Z., Kroschwald, S., Schuster, T., and D\u00f6ring, N., \u201cUsing large language models for\u00a0the\u00a0enforcement of\u00a0consumer rights in\u00a0Germany,\u201d J. Ma\u015blankowski, B. Marcinkowski, and P. Rupino da Cunha, Eds., in Lecture Notes in Business Information Processing. Springer Nature Switzerland, 2023, pp. 1\u201315. https:\/\/doi.org\/10.1007\/978-3-031-43590-4_1.","DOI":"10.1007\/978-3-031-43590-4_1"},{"key":"814_CR172","doi-asserted-by":"publisher","first-page":"431","DOI":"10.1613\/jair.1.12590","volume":"71","author":"S Kiritchenko","year":"2021","unstructured":"Kiritchenko, S., Nejadgholi, I., Fraser, K.C.: Confronting abusive language online: a survey from the Ethical and Human Rights perspective. J. Artif. Intell. Res. 71, 431\u2013478 (2021). https:\/\/doi.org\/10.1613\/jair.1.12590","journal-title":"J. Artif. Intell. Res."},{"key":"814_CR173","doi-asserted-by":"publisher","unstructured":"Nguyen, T.T., Wilson, C., and Dalins, J., \u201cFine-Tuning llama 2 large language models for detecting online sexual predatory chats and abusive texts,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2308.14683.","DOI":"10.48550\/arXiv.2308.14683"},{"key":"814_CR174","doi-asserted-by":"publisher","unstructured":"Plaza-del-arco, F.M., Nozza, D., and Hovy, D., \u201cRespectful or toxic? using zero-shot learning with language models to detect hate speech,\u201d in WOAH 2023, Y. Chung, P. R\\textbackslash\"ottger, D. Nozza, Z. Talat, and A. Mostafazadeh Davani, Eds., Association for Computational Linguistics, 2023, pp. 60\u201368. https:\/\/doi.org\/10.18653\/v1\/2023.woah-1.6.","DOI":"10.18653\/v1\/2023.woah-1.6"},{"key":"814_CR175","doi-asserted-by":"publisher","unstructured":"Hartvigsen, T., Gabriel, S., Palangi, H., Sap, M., Ray, D., and Kamar, E., \u201cToxiGen: a large-scale machine-generated dataset for adversarial and implicit hate speech detection,\u201d 2022, https:\/\/doi.org\/10.48550\/arXiv.2203.09509.","DOI":"10.48550\/arXiv.2203.09509"},{"key":"814_CR176","doi-asserted-by":"publisher","unstructured":"Felkner, V.K., Chang, H.-C. H., Jang, E., and May, J., \u201cWinoQueer: a community-in-the-loop benchmark for anti-LGBTQ+ bias in large language models,\u201d 2023, https:\/\/doi.org\/10.48550\/arXiv.2306.15087.","DOI":"10.48550\/arXiv.2306.15087"},{"key":"814_CR177","unstructured":"Ottosson, D., \u201cCyberbullying detection on social platforms using largelanguage models,\u201d 2023, [Online]. Available: https:\/\/urn.kb.se\/resolve?urn=urn:nbn:se:miun:diva-48990"},{"key":"814_CR178","doi-asserted-by":"publisher","DOI":"10.1007\/s43681-023-00289-2","author":"J M\u00f6kander","year":"2023","unstructured":"M\u00f6kander, J., Schuett, J., Kirk, H.R., Floridi, L.: Auditing large language models: a three-layered approach. AI and Ethics (2023). https:\/\/doi.org\/10.1007\/s43681-023-00289-2","journal-title":"AI and Ethics"},{"key":"814_CR179","unstructured":"Mireshghallah, F., \u201cAuditing and mitigating safety risks in large language models,\u201d 2023. [Online]. Available: https:\/\/escholarship.org\/uc\/item\/28f9b6px"},{"key":"814_CR180","unstructured":"Zhang, Y., Fitzgibbon, B., Garofolo, D., Kota, A., Papenhausen, E., and Mueller, K., \u201cAn explainable AI approach to large language model assisted causal model auditing and development\u201d, [Online]. Available: files\/9681\/Zhang et al.\u2013An Explainable AI Approach to Large Language Model.pdf"},{"key":"814_CR181","doi-asserted-by":"publisher","unstructured":"Jones, E., Dragan, A., Raghunathan, A., and Steinhardt, J., \u201cAutomatically auditing large language models via discrete optimization,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2303.04381.","DOI":"10.48550\/arXiv.2303.04381"},{"key":"814_CR182","doi-asserted-by":"publisher","unstructured":"Hasanbeig, H., Sharma, H., Betthauser, L., Frujeri, F.V., and Momennejad, I., \u201cALLURE: Auditing and improving LLM-based evaluation of text using iterative in-context-learning,\u201d 2023, arXiv. https:\/\/doi.org\/10.48550\/arXiv.2309.13701.","DOI":"10.48550\/arXiv.2309.13701"},{"key":"814_CR183","doi-asserted-by":"publisher","unstructured":"F\u00f6hr, T.L., Marten, K.-U., and Schreyer, M., \u201cDeep learning meets risk-based auditing: a holistic framework for leveraging foundation and task-specific models in audit procedures,\u201d 2023, Rochester, NY. https:\/\/doi.org\/10.2139\/ssrn.4488271.","DOI":"10.2139\/ssrn.4488271"},{"key":"814_CR184","doi-asserted-by":"publisher","DOI":"10.1145\/3600211.3604712","author":"C Rastogi","year":"2023","unstructured":"Rastogi, C., Tulio Ribeiro, M., King, N., Nori, H., Amershi, S.: \u201cSupporting human-AI collaboration in auditing LLMs with LLMs\u201d, in AIES \u201923. Assoc. Comput. Mach. (2023). https:\/\/doi.org\/10.1145\/3600211.3604712","journal-title":"Assoc. Comput. Mach."},{"key":"814_CR185","unstructured":"Kokhlikyan, N., Miglani, V., Martin, M., Wang, E., Alsallakh, B., Reynolds, J., ... & Reblitz-Richardson, O. (2020). Captum: A unified and generic model interpretability library for pytorch. arXiv preprint arXiv:2009.07896."},{"key":"814_CR186","doi-asserted-by":"crossref","unstructured":"Miglani, V., Yang, A., Markosyan, A.H., Garcia-Olano, D., & Kokhlikyan, N. (2023). Using captum to explain generative language models. arXiv preprint arXiv:2312.05491.","DOI":"10.18653\/v1\/2023.nlposs-1.19"},{"key":"814_CR187","doi-asserted-by":"crossref","unstructured":"Tufanov, I., Hambardzumyan, K., Ferrando, J., & Voita, E. (2024). LM transparency tool: Interactive tool for analyzing transformer language models. arXiv preprint arXiv:2404.07004.","DOI":"10.18653\/v1\/2024.acl-demos.6"},{"key":"814_CR188","doi-asserted-by":"crossref","unstructured":"Nadeem, M., Bethke, A., & Reddy, S. (2020). StereoSet: Measuring stereotypical bias in pretrained language models. arXiv preprint arXiv:2004.09456.","DOI":"10.18653\/v1\/2021.acl-long.416"},{"key":"814_CR189","doi-asserted-by":"crossref","unstructured":"Nangia, N., Vania, C., Bhalerao, R., & Bowman, S.R. (2020). CrowS-pairs: A challenge dataset for measuring social biases in masked language models. arXiv preprint arXiv:2010.00133.","DOI":"10.18653\/v1\/2020.emnlp-main.154"},{"key":"814_CR190","unstructured":"Wang, S., Li, R., Chen, X., Yuan, Y., Wong, D.F., & Yang, M. (2025). Exploring the impact of personality traits on llm bias and toxicity. arXiv preprint arXiv:2502.12566."},{"key":"814_CR191","doi-asserted-by":"crossref","unstructured":"Puttick, A., Rankwiler, L., Ikae, C., & Kurpicz-Briki, M. (2024). The BIAS Detection Framework: Bias Detection in Word Embeddings and Language Models for European Languages. arXiv preprint arXiv:2407.18689.","DOI":"10.18653\/v1\/2025.gebnlp-1.3"},{"key":"814_CR192","doi-asserted-by":"crossref","unstructured":"May, C., Wang, A., Bordia, S., Bowman, S.R., & Rudinger, R. (2019). On measuring social biases in sentence encoders. arXiv preprint arXiv:1903.10561.","DOI":"10.18653\/v1\/N19-1063"},{"key":"814_CR193","doi-asserted-by":"crossref","unstructured":"Binkyte, R. (2025). Interactional Fairness in LLM Multi-Agent Systems: An Evaluation Framework. arXiv preprint arXiv:2505.12001.","DOI":"10.1609\/aies.v8i1.36563"},{"key":"814_CR194","unstructured":"Achiam, J., Adler, S., Agarwal, S., Ahmad, L., Akkaya, I., Aleman, F.L., ... & McGrew, B. (2023). Gpt-4 technical report. arXiv preprint arXiv:2303.08774."},{"key":"814_CR195","unstructured":"Anthropic. (2024). The Claude 3 Model Family: Opus, Sonnet, Haiku"},{"key":"814_CR196","unstructured":"Team, G., Anil, R., Borgeaud, S., Alayrac, J.B., Yu, J., Soricut, R., ... & Blanco, L. (2023). Gemini: a family of highly capable multimodal models. arXiv preprint arXiv:2312.11805."},{"key":"814_CR197","doi-asserted-by":"crossref","unstructured":"Conneau, A., Lample, G., Rinott, R., Williams, A., Bowman, S.R., Schwenk, H., & Stoyanov, V. (2018). XNLI: Evaluating cross-lingual sentence representations. arXiv preprint arXiv:1809.05053.","DOI":"10.18653\/v1\/D18-1269"},{"key":"814_CR198","doi-asserted-by":"publisher","first-page":"522","DOI":"10.1162\/tacl_a_00474","volume":"10","author":"N Goyal","year":"2022","unstructured":"Goyal, N., Gao, C., Chaudhary, V., Chen, P.J., Wenzek, G., Ju, D., Fan, A.: The flores-101 evaluation benchmark for low-resource and multilingual machine translation. Transactions of the Association for Computational Linguistics 10, 522\u2013538 (2022)","journal-title":"Transactions of the Association for Computational Linguistics"},{"key":"814_CR199","doi-asserted-by":"publisher","first-page":"300","DOI":"10.1162\/tacl_a_00550","volume":"11","author":"AM Davani","year":"2023","unstructured":"Davani, A.M., Atari, M., Kennedy, B., Dehghani, M.: Hate speech classifiers learn normative social stereotypes. Transactions of the Association for Computational Linguistics 11, 300\u2013319 (2023)","journal-title":"Transactions of the Association for Computational Linguistics"},{"issue":"7","key":"814_CR200","doi-asserted-by":"publisher","DOI":"10.1093\/pnasnexus\/pgad210","volume":"2","author":"B Kennedy","year":"2023","unstructured":"Kennedy, B., Golazizian, P., Trager, J., Atari, M., Hoover, J., Mostafazadeh Davani, A., Dehghani, M.: The (moral) language of hate. PNAS Nexus 2(7), pgad210 (2023)","journal-title":"PNAS Nexus"},{"key":"814_CR201","doi-asserted-by":"crossref","unstructured":"Omrani Sabbaghi, S., Wolfe, R., & Caliskan, A. (2023, August). Evaluating biased attitude associations of language models in an intersectional context. In Proceedings of the 2023 AAAI\/ACM Conference on AI, Ethics, and Society (pp. 542\u2013553).","DOI":"10.1145\/3600211.3604666"},{"key":"814_CR202","unstructured":"Nicolas, G., & Caliskan, A. (2024). A taxonomy of stereotype content in large language models. arXiv preprint arXiv:2408.00162."},{"issue":"11","key":"814_CR203","doi-asserted-by":"publisher","DOI":"10.1093\/pnasnexus\/pgae493","volume":"3","author":"G Nicolas","year":"2024","unstructured":"Nicolas, G., Caliskan, A.: Directionality and representativeness are differentiable components of stereotypes in large language models. PNAS Nexus 3(11), pgae493 (2024)","journal-title":"PNAS Nexus"},{"key":"814_CR204","doi-asserted-by":"crossref","unstructured":"Omrani, A., Salkhordeh_Ziabari, A., Yu, C., Golazizian, P., Kennedy, B., Atari, M., ... & Dehghani, M. (2023, January). Social-group-agnostic bias mitigation via the stereotype content model. Association for Computational Linguistics.","DOI":"10.18653\/v1\/2023.acl-long.227"},{"key":"814_CR205","doi-asserted-by":"publisher","DOI":"10.1145\/3728881","author":"Y Xiao","year":"2025","unstructured":"Xiao, Y., Liu, A., Liang, S., Liu, X., Tao, D.: Fairness mediator: neutralize stereotype associations to mitigate bias in large language models. Proce. ACM Softw. Eng. (2025). https:\/\/doi.org\/10.1145\/3728881","journal-title":"Proce. ACM Softw. Eng."}],"container-title":["AI and Ethics"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s43681-025-00814-5.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s43681-025-00814-5\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s43681-025-00814-5.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,11,7]],"date-time":"2025-11-07T01:39:08Z","timestamp":1762479548000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s43681-025-00814-5"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,10,13]]},"references-count":205,"journal-issue":{"issue":"6","published-print":{"date-parts":[[2025,12]]}},"alternative-id":["814"],"URL":"https:\/\/doi.org\/10.1007\/s43681-025-00814-5","relation":{},"ISSN":["2730-5953","2730-5961"],"issn-type":[{"value":"2730-5953","type":"print"},{"value":"2730-5961","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,10,13]]},"assertion":[{"value":"31 July 2024","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"1 July 2025","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"13 October 2025","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}},{"value":"Not applicable.","order":3,"name":"Ethics","group":{"name":"EthicsHeading","label":"Informed consent"}}]}}