{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,8,25]],"date-time":"2026-08-25T11:52:25Z","timestamp":1787658745295,"version":"build-2736575974"},"reference-count":146,"publisher":"Springer Science and Business Media LLC","license":[{"start":{"date-parts":[[2026,8,24]],"date-time":"2026-08-24T00:00:00Z","timestamp":1787529600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,8,24]],"date-time":"2026-08-24T00:00:00Z","timestamp":1787529600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/100006754","name":"United States Department of Defense | United States Army | U.S. Army Research, Development and Engineering Command | Army Research Laboratory","doi-asserted-by":"publisher","award":["W911NF-23-2-0183"],"award-info":[{"award-number":["W911NF-23-2-0183"]}],"id":[{"id":"10.13039\/100006754","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/100000185","name":"United States Department of Defense | Defense Advanced Research Projects Agency","doi-asserted-by":"publisher","award":["HR001121C0165"],"award-info":[{"award-number":["HR001121C0165"]}],"id":[{"id":"10.13039\/100000185","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/100000181","name":"United States Department of Defense | United States Air Force | AFMC | Air Force Office of Scientific Research","doi-asserted-by":"publisher","award":["A9550-23-1-0463"],"award-info":[{"award-number":["A9550-23-1-0463"]}],"id":[{"id":"10.13039\/100000181","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Nat Hum Behav"],"DOI":"10.1038\/s41562-026-02550-0","type":"journal-article","created":{"date-parts":[[2026,8,24]],"date-time":"2026-08-24T15:04:08Z","timestamp":1787583848000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["The shrinking landscape of linguistic diversity in the age of large language models"],"prefix":"10.1038","author":[{"ORCID":"https:\/\/orcid.org\/0000-0003-2129-6165","authenticated-orcid":false,"given":"Zhivar","family":"Sourati","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-3381-7071","authenticated-orcid":false,"given":"Farzan","family":"Karimi-Malekabadi","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-5054-8264","authenticated-orcid":false,"given":"Meltem","family":"Ozcan","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-3880-164X","authenticated-orcid":false,"given":"Colin","family":"McDaniel","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Alireza","family":"Ziabari","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jackson","family":"Trager","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Ala N.","family":"Tak","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Meng","family":"Chen","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Fred","family":"Morstatter","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-9478-4365","authenticated-orcid":false,"given":"Morteza","family":"Dehghani","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,8,24]]},"reference":[{"key":"2550_CR1","unstructured":"Orwell, G. Nineteen Eighty-Four 69\u201370 (Penguin Books, 1949)."},{"key":"2550_CR2","doi-asserted-by":"publisher","first-page":"934","DOI":"10.1037\/pspp0000020","volume":"108","author":"G Park","year":"2015","unstructured":"Park, G. et al. Automatic personality assessment through social media language. J. Pers. Soc. Psychol. 108, 934\u2013952 (2015).","journal-title":"J. Pers. Soc. Psychol."},{"key":"2550_CR3","doi-asserted-by":"publisher","first-page":"239","DOI":"10.1207\/s15326950dp4203_1","volume":"42","author":"J Oberlander","year":"2006","unstructured":"Oberlander, J. & Gill, A. J. Language with character: a stratified corpus comparison of individual differences in e-mail communication. Discourse Process. 42, 239\u2013270 (2006).","journal-title":"Discourse Process."},{"key":"2550_CR4","doi-asserted-by":"publisher","first-page":"110818","DOI":"10.1016\/j.paid.2021.110818","volume":"177","author":"JD Moreno","year":"2021","unstructured":"Moreno, J. D., Martinez-Huertas, J. A., Olmos, R., Jorge-Botana, G. & Botella, J. Can personality traits be measured analyzing written language? A meta-analytic study on computational methods. Pers. Individ. Dif. 177, 110818 (2021).","journal-title":"Pers. Individ. Dif."},{"key":"2550_CR5","doi-asserted-by":"publisher","first-page":"457","DOI":"10.1613\/jair.2349","volume":"30","author":"F Mairesse","year":"2007","unstructured":"Mairesse, F., Walker, M. A., Mehl, M. R. & Moore, R. K. Using linguistic cues for the automatic recognition of personality in conversation and text. J. Artif. Intell. Res. 30, 457\u2013500 (2007).","journal-title":"J. Artif. Intell. Res."},{"key":"2550_CR6","doi-asserted-by":"publisher","first-page":"73791","DOI":"10.1371\/journal.pone.0073791","volume":"8","author":"HA Schwartz","year":"2013","unstructured":"Schwartz, H. A. et al. Personality, gender, and age in the language of social media: the open-vocabulary approach. PLoS ONE 8, 73791 (2013).","journal-title":"PLoS ONE"},{"key":"2550_CR7","doi-asserted-by":"publisher","first-page":"30","DOI":"10.1075\/aila.27.02kra","volume":"27","author":"C Kramsch","year":"2014","unstructured":"Kramsch, C. Language and culture. AILA Rev. 27, 30\u201355 (2014).","journal-title":"AILA Rev."},{"key":"2550_CR8","first-page":"381","volume":"9","author":"J Gumperz","year":"1968","unstructured":"Gumperz, J. The speech community. Int. Encycl. Soc. Sci. 9, 381\u2013386 (1968).","journal-title":"Int. Encycl. Soc. Sci."},{"key":"2550_CR9","unstructured":"Nguyen, D. & Ros\u00e9, C. P. Language use as a reflection of socialization in online communities. In Proc. Workshop on Language in Social Media (eds Nagarajan, M. & Gamon, M.) 76\u201385 (Association for Computational Linguistics, 2011)."},{"key":"2550_CR10","doi-asserted-by":"publisher","first-page":"135","DOI":"10.1111\/josl.12080","volume":"18","author":"D Bamman","year":"2014","unstructured":"Bamman, D., Eisenstein, J. & Schnoebelen, T. Gender identity and lexical variation in social media. J. Socioling. 18, 135\u2013160 (2014).","journal-title":"J. Socioling."},{"key":"2550_CR11","doi-asserted-by":"publisher","first-page":"42","DOI":"10.1016\/S0262-4079(11)62167-2","volume":"211","author":"JW Pennebaker","year":"2011","unstructured":"Pennebaker, J. W. The secret life of pronouns. New Sci. 211, 42\u201345 (2011).","journal-title":"New Sci."},{"key":"2550_CR12","doi-asserted-by":"publisher","first-page":"438","DOI":"10.1177\/0261927X16668376","volume":"36","author":"MD Robinson","year":"2017","unstructured":"Robinson, M. D., Boyd, R. L., Fetterman, A. K. & Persich, M. R. The mind versus the body in political (and nonpolitical) discourse: linguistic evidence for an ideological signature in US politics. J. Lang. Soc. Psychol. 36, 438\u2013461 (2017).","journal-title":"J. Lang. Soc. Psychol."},{"key":"2550_CR13","doi-asserted-by":"publisher","first-page":"244","DOI":"10.1016\/j.compenvurbsys.2015.12.003","volume":"59","author":"Y Huang","year":"2016","unstructured":"Huang, Y., Guo, D., Kasakoff, A. & Grieve, J. Understanding US regional linguistic variation with Twitter data analysis. Comput. Environ. Urban Syst. 59, 244\u2013255 (2016).","journal-title":"Comput. Environ. Urban Syst."},{"key":"2550_CR14","unstructured":"Eisenstein, J., O\u2019Connor, B., Smith, N. A. & Xing, E. A latent variable model for geographic lexical variation. In Proc. 2010 Conference on Empirical Methods in Natural Language Processing (eds Li, H. & M\u00e0rquez, L.) 1277\u20131287 (Association for Computational Linguistics, 2010)."},{"key":"2550_CR15","unstructured":"Peterson, K., Hohensee, M. & Xia, F. Email formality in the workplace: a case study on the enron corpus. In Proc. Workshop on Language in Social Media (LSM 2011) (eds Nagarajan, M. & Gamon, M.) 86\u201395 (Association for Computational Linguistics, 2011)."},{"key":"2550_CR16","doi-asserted-by":"publisher","first-page":"538","DOI":"10.1002\/asi.21001","volume":"60","author":"E Stamatatos","year":"2009","unstructured":"Stamatatos, E. A survey of modern authorship attribution methods. J. Assoc. Inf. Sci. Technol. 60, 538\u2013556 (2009).","journal-title":"J. Assoc. Inf. Sci. Technol."},{"key":"2550_CR17","doi-asserted-by":"publisher","first-page":"251","DOI":"10.1093\/llc\/fqm020","volume":"22","author":"J Grieve","year":"2007","unstructured":"Grieve, J. Quantitative authorship attribution: an evaluation of techniques. Lit. Ling. Comput. 22, 251\u2013270 (2007).","journal-title":"Lit. Ling. Comput."},{"key":"2550_CR18","first-page":"1027","volume":"10","author":"J Cassell","year":"2005","unstructured":"Cassell, J. & Tversky, D. The language of online intercultural community formation. J. Comput. Mediat. Commun. 10, 1027 (2005).","journal-title":"J. Comput. Mediat. Commun."},{"key":"2550_CR19","first-page":"9110","volume":"9","author":"B Danet","year":"2003","unstructured":"Danet, B. & Herring, S. C. Introduction: the multilingual internet. J. Comput. Mediat. Commun. 9, 9110 (2003).","journal-title":"J. Comput. Mediat. Commun."},{"key":"2550_CR20","doi-asserted-by":"publisher","first-page":"104696","DOI":"10.1016\/j.cognition.2021.104696","volume":"212","author":"B Kennedy","year":"2021","unstructured":"Kennedy, B. et al. Moral concerns are differentially observable in language. Cognition 212, 104696 (2021).","journal-title":"Cognition"},{"key":"2550_CR21","doi-asserted-by":"publisher","first-page":"244","DOI":"10.1038\/s41562-018-0516-z","volume":"3","author":"JC Jackson","year":"2019","unstructured":"Jackson, J. C., Gelfand, M., De, S. & Fox, A. The loosening of American culture over 200 years is associated with a creativity\u2013order trade-off. Nat. Hum. Behav. 3, 244\u2013250 (2019).","journal-title":"Nat. Hum. Behav."},{"key":"2550_CR22","doi-asserted-by":"publisher","first-page":"147","DOI":"10.1038\/s41586-024-07856-5","volume":"633","author":"V Hofmann","year":"2024","unstructured":"Hofmann, V., Kalluri, P. R., Jurafsky, D. & King, S. AI generates covertly racist decisions about people based on their dialect. Nature 633, 147\u2013154 (2024).","journal-title":"Nature"},{"key":"2550_CR23","doi-asserted-by":"publisher","first-page":"4714","DOI":"10.1044\/2024_JSLHR-24-00274","volume":"67","author":"AB Richard","year":"2024","unstructured":"Richard, A. B., Lelandais, M., Reilly, K. T. & Jacquin-Courtois, S. Linguistic markers of subtle cognitive impairment in connected speech: a systematic review. J. Speech Lang. Hear. Res. 67, 4714\u20134733 (2024).","journal-title":"J. Speech Lang. Hear. Res."},{"key":"2550_CR24","doi-asserted-by":"crossref","unstructured":"Eyigoz, E., Mathur, S., Santamaria, M., Cecchi, G. & Naylor, M. Linguistic markers predict onset of Alzheimer\u2019s disease. EClinicalMedicine 28, 100583 (2020).","DOI":"10.1016\/j.eclinm.2020.100583"},{"key":"2550_CR25","doi-asserted-by":"publisher","first-page":"2081","DOI":"10.1109\/TASL.2011.2112351","volume":"19","author":"B Roark","year":"2011","unstructured":"Roark, B., Mitchell, M., Hosom, J.-P., Hollingshead, K. & Kaye, J. Spoken language derived measures for detecting mild cognitive impairment. IEEE Trans. Audio Speech Lang. Process. 19, 2081\u20132090 (2011).","journal-title":"IEEE Trans. Audio Speech Lang. Process."},{"key":"2550_CR26","doi-asserted-by":"publisher","first-page":"1355734","DOI":"10.3389\/fpsyg.2024.1355734","volume":"15","author":"RN Trifu","year":"2024","unstructured":"Trifu, R. N. et al. Linguistic markers for major depressive disorder: a cross-sectional study using an automated procedure. Front. Psychol. 15, 1355734 (2024).","journal-title":"Front. Psychol."},{"key":"2550_CR27","doi-asserted-by":"crossref","unstructured":"Weerasinghe, J., Morales, K. & Greenstadt, R. \"Because\u2026 i was told\u2026 so much\u201d: linguistic indicators of mental health status on Twitter. Proc. Priv. Enhanc. Technol. 4, 152\u2013171 (2019).","DOI":"10.2478\/popets-2019-0063"},{"key":"2550_CR28","unstructured":"Whorf, B. L. Language, Thought, and Reality: Selected Writings of Benjamin Lee Whorf (MIT Press, 2012)."},{"key":"2550_CR29","doi-asserted-by":"publisher","first-page":"87","DOI":"10.1146\/annurev-anthro-092611-145828","volume":"41","author":"P Eckert","year":"2012","unstructured":"Eckert, P. Three waves of variation study: the emergence of meaning in the study of sociolinguistic variation. Annu. Rev. Anthropol. 41, 87\u2013100 (2012).","journal-title":"Annu. Rev. Anthropol."},{"key":"2550_CR30","unstructured":"Hofstede, G. Culture\u2019s Consequences: Comparing Values, Behaviors, Institutions and Organizations Across Nations 2nd edn (Sage, 2001)."},{"key":"2550_CR31","unstructured":"OpenAI. Introducing ChatGPT https:\/\/openai.com\/blog\/chatgpt (2022)."},{"key":"2550_CR32","doi-asserted-by":"publisher","unstructured":"Gemini Team et al. Gemini: a family of highly capable multimodal models. Preprint at https:\/\/doi.org\/10.48550\/arXiv.2312.11805 (2023).","DOI":"10.48550\/arXiv.2312.11805"},{"key":"2550_CR33","unstructured":"Bailyn, E. ChatGPT Usage Statistics: March 2026 https:\/\/firstpagesage.com\/seo-blog\/chatgpt-usage-statistics\/ (FirstPageSage, 2026)."},{"key":"2550_CR34","unstructured":"Nearly 1 in 3 College Students Have Used ChatGPT on Written Assignments https:\/\/www.intelligent.com\/nearly-1-in-3-college-students-have-used-chatgpt-on-written-assignments\/ (Intelligent, 2024)."},{"key":"2550_CR35","unstructured":"McClain, C. Americans\u2019 Use of ChatGPT is Ticking Up, But Few Trust Its Election Information https:\/\/www.pewresearch.org\/short-reads\/2024\/03\/26\/americans-use-of-chatgpt-is-ticking-up-but-few-trust-its-election-information\/ (Pew Research Center, 2024)."},{"key":"2550_CR36","doi-asserted-by":"publisher","unstructured":"Handa, K. et al. Which economic tasks are performed with AI? Evidence from millions of Claude conversations. Preprint at https:\/\/doi.org\/10.48550\/arXiv.2503.04761 (2025).","DOI":"10.48550\/arXiv.2503.04761"},{"key":"2550_CR37","doi-asserted-by":"publisher","first-page":"933","DOI":"10.1162\/tacl_a_00681","volume":"12","author":"M Mizrahi","year":"2024","unstructured":"Mizrahi, M. et al. State of what art? A call for multi-prompt LLM evaluation. Trans. Assoc. Comput. Linguist. 12, 933\u2013949 (2024).","journal-title":"Trans. Assoc. Comput. Linguist."},{"key":"2550_CR38","doi-asserted-by":"publisher","first-page":"1954","DOI":"10.1038\/s42256-025-01115-6","volume":"7","author":"G Serapio-Garc\u00eda","year":"2025","unstructured":"Serapio-Garc\u00eda, G. et al. A psychometric framework for evaluating and shaping personality traits in large language models. Nat. Mach. Intell. 7, 1954\u20131968 (2025).","journal-title":"Nat. Mach. Intell."},{"key":"2550_CR39","unstructured":"Ghosh, S. et al. A closer look at the limitations of instruction tuning. In Proc. 41st International Conference on Machine Learning 624 (JMLR, 2024)."},{"key":"2550_CR40","unstructured":"Santurkar, S. et al. Whose opinions do language models reflect? In International Conference on Machine Learning 29971\u201330004 (JMLR, 2023)."},{"key":"2550_CR41","doi-asserted-by":"publisher","unstructured":"Ireland, M.E. & Mehl, M. R. in The Oxford Handbook of Language and Social Psychology (ed. Holtgraves, T. M.) 201\u2013218 https:\/\/doi.org\/10.1093\/oxfordhb\/9780199838639.013.034 (Oxford Univ. Press, 2014).","DOI":"10.1093\/oxfordhb\/9780199838639.013.034"},{"key":"2550_CR42","doi-asserted-by":"publisher","first-page":"86","DOI":"10.1093\/schbul\/sbac215","volume":"49","author":"H Corona Hern\u00e1ndez","year":"2023","unstructured":"Corona Hern\u00e1ndez, H. et al. Natural language processing markers for psychosis and other psychiatric disorders: emerging themes and research agenda from a cross-linguistic workshop. Schizophr. Bull. 49, 86\u201392 (2023).","journal-title":"Schizophr. Bull."},{"key":"2550_CR43","doi-asserted-by":"publisher","first-page":"1121","DOI":"10.1080\/02699930441000030","volume":"18","author":"S Rude","year":"2004","unstructured":"Rude, S., Gortner, E.-M. & Pennebaker, J. Language use of depressed and depression-vulnerable college students. Cogn. Emot. 18, 1121\u20131133 (2004).","journal-title":"Cogn. Emot."},{"key":"2550_CR44","doi-asserted-by":"publisher","first-page":"117822261879286","DOI":"10.1177\/1178222618792860","volume":"10","author":"G Coppersmith","year":"2018","unstructured":"Coppersmith, G., Leary, R., Crutchley, P. & Fine, A. Natural language processing of social media as screening for suicide risk. Biomed. Inform. Insights 10, 1178222618792860 (2018).","journal-title":"Biomed. Inform. Insights"},{"key":"2550_CR45","doi-asserted-by":"publisher","first-page":"7","DOI":"10.1016\/j.cobeha.2017.05.009","volume":"18","author":"SC Matz","year":"2017","unstructured":"Matz, S. C. & Netzer, O. Using big data as a window into consumers\u2019 psychology. Curr. Opin. Behav. Sci. 18, 7\u201312 (2017).","journal-title":"Curr. Opin. Behav. Sci."},{"key":"2550_CR46","doi-asserted-by":"publisher","first-page":"106525","DOI":"10.1016\/j.chb.2020.106525","volume":"114","author":"S Winter","year":"2021","unstructured":"Winter, S., Maslowska, E. & Vos, A. L. The effects of trait-based personalization in social media advertising. Comput. Hum. Behav. 114, 106525 (2021).","journal-title":"Comput. Hum. Behav."},{"key":"2550_CR47","doi-asserted-by":"publisher","first-page":"298","DOI":"10.1111\/j.1468-2958.2010.01377.x","volume":"36","author":"SS Sundar","year":"2010","unstructured":"Sundar, S. S. & Marathe, S. S. Personalization versus customization: the importance of agency, privacy, and power usage. Hum. Commun. Res. 36, 298\u2013322 (2010).","journal-title":"Hum. Commun. Res."},{"key":"2550_CR48","doi-asserted-by":"crossref","unstructured":"Ryan, M. J., Held, W. & Yang, D. Unintended impacts of LLM alignment on global representation. In Proc. 62nd Annual Meeting of the Association for Computational Linguistics1, 16121\u201316140 (2024).","DOI":"10.18653\/v1\/2024.acl-long.853"},{"key":"2550_CR49","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1145\/3597307","volume":"15","author":"R Navigli","year":"2023","unstructured":"Navigli, R., Conia, S. & Ross, B. Biases in large language models: origins, inventory, and discussion. ACMJ. Data Inf. Qual. 15, 1\u201321 (2023).","journal-title":"ACMJ. Data Inf. Qual."},{"key":"2550_CR50","doi-asserted-by":"publisher","unstructured":"Atari, M., Xue, M. J., Park, P. S., Blasi, D. & Henrich, J. Which humans? Preprint at PsyArXiv https:\/\/doi.org\/10.31234\/osf.io\/5b26t (2023).","DOI":"10.31234\/osf.io\/5b26t"},{"key":"2550_CR51","doi-asserted-by":"publisher","first-page":"400","DOI":"10.1038\/s42256-025-00986-z","volume":"7","author":"A Wang","year":"2025","unstructured":"Wang, A., Morgenstern, J. & Dickerson, J. P. Large language models that replace human participants can harmfully misportray and flatten identity groups. Nat. Mach. Intell. 7, 400\u2013411 (2025).","journal-title":"Nat. Mach. Intell."},{"key":"2550_CR52","doi-asserted-by":"publisher","first-page":"0306621","DOI":"10.1371\/journal.pone.0306621","volume":"19","author":"D Rozado","year":"2024","unstructured":"Rozado, D. The political preferences of LLMs. PLoS ONE 19, 0306621 (2024).","journal-title":"PLoS ONE"},{"key":"2550_CR53","doi-asserted-by":"publisher","unstructured":"Pan, K. & Zeng, Y. Do LLMs possess a personality? Making the MBTI test an amazing evaluation for large language models. Preprint at https:\/\/doi.org\/10.48550\/arXiv.2307.16180 (2023).","DOI":"10.48550\/arXiv.2307.16180"},{"key":"2550_CR54","doi-asserted-by":"publisher","first-page":"245","DOI":"10.1093\/pnasnexus\/pgae245","volume":"3","author":"S Abdurahman","year":"2024","unstructured":"Abdurahman, S. et al. Perils and opportunities in using large language models in psychological research. PNAS Nexus 3, 245 (2024).","journal-title":"PNAS Nexus"},{"key":"2550_CR55","doi-asserted-by":"publisher","first-page":"3813","DOI":"10.1126\/sciadv.adt3813","volume":"11","author":"D Kobak","year":"2025","unstructured":"Kobak, D., Gonz\u00e1lez-M\u00e1rquez, R., Horv\u00e1t, E. -\u00c1 & Lause, J. Delving into LLM-assisted writing in biomedical publications through excess vocabulary. Sci. Adv. 11, 3813 (2025).","journal-title":"Sci. Adv."},{"key":"2550_CR56","doi-asserted-by":"publisher","first-page":"2599","DOI":"10.1038\/s41562-025-02273-8","volume":"9","author":"W Liang","year":"2025","unstructured":"Liang, W. et al. Quantifying large language model usage in scientific papers. Nat. Hum. Behav. 9, 2599\u20132609 (2025).","journal-title":"Nat. Hum. Behav."},{"key":"2550_CR57","doi-asserted-by":"crossref","unstructured":"Liang, W. et al. The widespread adoption of large language model-assisted writing across society. Patterns 6, 101366 (2025).","DOI":"10.1016\/j.patter.2025.101366"},{"key":"2550_CR58","doi-asserted-by":"publisher","first-page":"3597","DOI":"10.1007\/s11192-025-05341-y","volume":"130","author":"T Bao","year":"2025","unstructured":"Bao, T., Zhao, Y., Mao, J. & Zhang, C. Examining linguistic shifts in academic writing before and after the launch of chatGPT: a study on preprint papers. Scientometrics 130, 3597\u20133627 (2025).","journal-title":"Scientometrics"},{"key":"2550_CR59","doi-asserted-by":"crossref","unstructured":"Guo, Y., Shang, G., Vazirgiannis, M. & Clavel, C. The curious decline of linguistic diversity: training language models on synthetic text. In Findings of the Association for Computational Linguistics: NAACL 2024 3589\u20133604 (ACL, 2024).","DOI":"10.18653\/v1\/2024.findings-naacl.228"},{"key":"2550_CR60","doi-asserted-by":"publisher","first-page":"265","DOI":"10.1007\/s10462-024-10903-2","volume":"57","author":"A Mu\u00f1oz-Ortiz","year":"2024","unstructured":"Mu\u00f1oz-Ortiz, A., G\u00f3mez-Rodr\u00edguez, C. & Vilares, D. Contrasting linguistic patterns in human and LLM-generated news text. Artif. Intell. Rev. 57, 265 (2024).","journal-title":"Artif. Intell. Rev."},{"key":"2550_CR61","doi-asserted-by":"publisher","first-page":"2504966122","DOI":"10.1073\/pnas.2504966122","volume":"122","author":"W Xu","year":"2025","unstructured":"Xu, W., Jojic, N., Rao, S., Brockett, C. & Dolan, B. Echoes in AI: quantifying lack of plot diversity in LLM outputs. Proc. Natl Acad. Sci. USA 122, 2504966122 (2025).","journal-title":"Proc. Natl Acad. Sci. USA"},{"key":"2550_CR62","unstructured":"Padmakumar, V. & He, H. Does writing with language models reduce content diversity? In The Twelfth International Conference on Learning Representations. https:\/\/openreview.net\/forum?id=Feiz5HtCD0"},{"key":"2550_CR63","doi-asserted-by":"publisher","first-page":"5290","DOI":"10.1126\/sciadv.adn5290","volume":"10","author":"AR Doshi","year":"2024","unstructured":"Doshi, A. R. & Hauser, O. P. Generative AI enhances individual creativity but reduces the collective diversity of novel content. Sci. Adv. 10, 5290 (2024).","journal-title":"Sci. Adv."},{"key":"2550_CR64","doi-asserted-by":"crossref","unstructured":"Anderson, B. R., Shah, J. H. & Kreminski, M. Homogenization effects of large language models on human creative ideation. In Proc. 16th Conference on Creativity & Cognition 413\u2013425 (Association for Computing Machinery, 2024).","DOI":"10.1145\/3635636.3656204"},{"key":"2550_CR65","doi-asserted-by":"crossref","unstructured":"Agarwal, D., Naaman, M. & Vashistha, A. AI suggestions homogenize writing toward western styles and diminish cultural nuances. In Proc. 2025 CHI Conference on Human Factors in Computing Systems 1\u201321 (Association for Computing Machinery, 2025).","DOI":"10.1145\/3706598.3713564"},{"key":"2550_CR66","doi-asserted-by":"publisher","first-page":"100207","DOI":"10.1016\/j.chbah.2025.100207","volume":"6","author":"K Moon","year":"2025","unstructured":"Moon, K., Green, A. E. & Kushlev, K. Homogenizing effect of large language models (LLMs) on creative diversity: an empirical comparison of human and chatGPT writing. Comput. Hum. Behav. Artif. Hum. 6, 100207 (2025).","journal-title":"Comput. Hum. Behav. Artif. Hum."},{"key":"2550_CR67","doi-asserted-by":"publisher","DOI":"10.1186\/s40537-024-00986-7","volume":"11","author":"A Alvero","year":"2024","unstructured":"Alvero, A. et al. Large language models, social demography, and hegemony: comparing authorship in human and synthetic text. J. Big Data 11, 138 (2024).","journal-title":"J. Big Data"},{"key":"2550_CR68","doi-asserted-by":"publisher","unstructured":"Zhang, S., Xu, J. & Alvero, A. Generative AI meets open-ended survey responses: research participant use of AI and homogenization. Sociol. Methods Res. https:\/\/doi.org\/10.1177\/00491241251327130 (2025).","DOI":"10.1177\/00491241251327130"},{"key":"2550_CR69","unstructured":"Lee, J., Alvero, A., Joachims, T. & Kizilcec, R. Poor alignment and steerability of large language models: evidence from college admission essays. In Proc. Workshop on Social Simulation with LLMs at COLM (2025)."},{"key":"2550_CR70","unstructured":"Hans, A. et al. Spotting LLMs with binoculars: zero-shot detection of machine-generated text. In Proc. 41st International Conference on Machine Learning 698 (JMLR, 2024)."},{"key":"2550_CR71","doi-asserted-by":"publisher","unstructured":"Tufts, B., Zhao, X. & Li, L. A practical examination of AI-generated text detectors for large language models. In Findings of the Association for Computational Linguistics: NAACL 2025 (eds Chiruzzo, L. et al.) 4839\u20134856 https:\/\/doi.org\/10.18653\/v1\/2025.findings-naacl.271 (Association for Computational Linguistics, 2025).","DOI":"10.18653\/v1\/2025.findings-naacl.271"},{"key":"2550_CR72","doi-asserted-by":"crossref","unstructured":"Dugan, L. et al. Raid: a shared benchmark for robust evaluation of machine-generated text detectors. In Proc. 62nd Annual Meeting of the Association for Computational Linguistics1, 12463\u201312492 (ACL, 2024).","DOI":"10.18653\/v1\/2024.acl-long.674"},{"key":"2550_CR73","unstructured":"Buchert, J.-M. The 6 Best AI Detectors Based on Objective Studies and Usage https:\/\/intellectualead.com\/best-ai-detectors-guide\/ (Intellectual Lead, 2025)."},{"key":"2550_CR74","doi-asserted-by":"crossref","unstructured":"Singer, J. D., Willett, J. B. Applied Longitudinal Data Analysis: Modeling Change and Event Occurrence (Oxford Univ, 2003).","DOI":"10.1093\/acprof:oso\/9780195152968.001.0001"},{"key":"2550_CR75","doi-asserted-by":"publisher","first-page":"480","DOI":"10.1177\/1536867X1501500208","volume":"15","author":"A Linden","year":"2015","unstructured":"Linden, A. Conducting interrupted time-series analysis for single- and multiple-group comparisons. Stata J. 15, 480\u2013500 (2015).","journal-title":"Stata J."},{"key":"2550_CR76","doi-asserted-by":"publisher","first-page":"7","DOI":"10.1016\/0165-1765(95)00791-1","volume":"51","author":"D Bell","year":"1996","unstructured":"Bell, D., Kay, J. & Malley, J. A non-parametric approach to non-linear causality testing. Econ. Lett. 51, 7\u201318 (1996).","journal-title":"Econ. Lett."},{"key":"2550_CR77","doi-asserted-by":"publisher","unstructured":"Dubey, A. et al. The llama 3 herd of models. Preprint at https:\/\/doi.org\/10.48550\/arXiv.2407.21783 (2024).","DOI":"10.48550\/arXiv.2407.21783"},{"key":"2550_CR78","doi-asserted-by":"publisher","first-page":"5802","DOI":"10.1073\/pnas.1218772110","volume":"110","author":"M Kosinski","year":"2013","unstructured":"Kosinski, M., Stillwell, D. & Graepel, T. Private traits and attributes are predictable from digital records of human behavior. Proc. Natl Acad. Sci. USA 110, 5802\u20135805 (2013).","journal-title":"Proc. Natl Acad. Sci. USA"},{"key":"2550_CR79","doi-asserted-by":"publisher","unstructured":"Tufekci, Z. Engineering the public: big data, surveillance and computational politics. First Monday https:\/\/doi.org\/10.5210\/fm.v19i7.4901 (2014).","DOI":"10.5210\/fm.v19i7.4901"},{"key":"2550_CR80","doi-asserted-by":"publisher","first-page":"3635","DOI":"10.1073\/pnas.1720347115","volume":"115","author":"N Garg","year":"2018","unstructured":"Garg, N., Schiebinger, L., Jurafsky, D. & Zou, J. Word embeddings quantify 100 years of gender and ethnic stereotypes. Proc. Natl Acad. Sci. USA 115, 3635\u20133644 (2018).","journal-title":"Proc. Natl Acad. Sci. USA"},{"key":"2550_CR81","unstructured":"Kennedy, B., Ashokkumar, A., Boyd, R. L. & Dehghani, M. in Handbook of Language Analysis in Psychology (eds Dehghani, M. & Boyd, R. L.) 3\u201362 (Guilford Publications, 2022)."},{"key":"2550_CR82","doi-asserted-by":"publisher","first-page":"524","DOI":"10.1016\/j.jrp.2009.01.006","volume":"43","author":"JB Hirsh","year":"2009","unstructured":"Hirsh, J. B. & Peterson, J. B. Personality and language use in self-narratives. J. Res. Pers. 43, 524\u2013527 (2009).","journal-title":"J. Res. Pers."},{"key":"2550_CR83","doi-asserted-by":"publisher","first-page":"862","DOI":"10.1037\/0022-3514.90.5.862","volume":"90","author":"MR Mehl","year":"2006","unstructured":"Mehl, M. R., Gosling, S. D. & Pennebaker, J. W. Personality in its natural habitat: manifestations and implicit folk theories of personality in daily life. J. Pers. Soc. Psychol. 90, 862\u2013877 (2006).","journal-title":"J. Pers. Soc. Psychol."},{"key":"2550_CR84","unstructured":"Boyd, R. L., Ashokkumar, A., Seraj, S. & Pennebaker, J. W. The Development and Psychometric Properties of LIWC-22 Technical report https:\/\/www.liwc.app (Univ. Texas at Austin, 2022)."},{"key":"2550_CR85","unstructured":"Frimer, J. A., Boghrati, R., Haidt, J., Graham, J. & Dehgani, M. Moral Foundations Dictionary for Linguistic Analyses 2.0 (Provalis Research, 2019)."},{"key":"2550_CR86","doi-asserted-by":"publisher","first-page":"593","DOI":"10.1016\/j.sbspro.2015.06.078","volume":"192","author":"Y Ishikawa","year":"2015","unstructured":"Ishikawa, Y. Gender differences in vocabulary use in essay writing by university students. ProcediaSoc. Behav. Sci. 192, 593\u2013600 (2015).","journal-title":"ProcediaSoc. Behav. Sci."},{"key":"2550_CR87","doi-asserted-by":"publisher","first-page":"104035","DOI":"10.1016\/j.jrp.2020.104035","volume":"89","author":"J Chen","year":"2020","unstructured":"Chen, J., Qiu, L. & Ho, M.-H. R. A meta-analysis of linguistic markers of extraversion: positive emotion and social process words. J. Res. Pers. 89, 104035 (2020).","journal-title":"J. Res. Pers."},{"key":"2550_CR88","doi-asserted-by":"publisher","unstructured":"Li, L. & Tomasello, M. On the moral functions of language. Soc. Cogn. https:\/\/doi.org\/10.1521\/soco.2021.39.1.99 (2021).","DOI":"10.1521\/soco.2021.39.1.99"},{"key":"2550_CR89","doi-asserted-by":"publisher","first-page":"537","DOI":"10.1162\/COLI_a_00258","volume":"42","author":"D Nguyen","year":"2016","unstructured":"Nguyen, D., Do\u011fru\u00f6z, A. S., Ros\u00e9, C. P. & De Jong, F. Computational sociolinguistics: a survey. Comput. Linguist. 42, 537\u2013593 (2016).","journal-title":"Comput. Linguist."},{"key":"2550_CR90","doi-asserted-by":"publisher","first-page":"24","DOI":"10.1177\/0261927X09351676","volume":"29","author":"YR Tausczik","year":"2010","unstructured":"Tausczik, Y. R. & Pennebaker, J. W. The psychological meaning of words: LIWC and computerized text analysis methods. J. Lang. Soc. Psychol. 29, 24\u201354 (2010).","journal-title":"J. Lang. Soc. Psychol."},{"key":"2550_CR91","doi-asserted-by":"publisher","first-page":"429","DOI":"10.1017\/S0140525X0999094X","volume":"32","author":"N Evans","year":"2009","unstructured":"Evans, N. & Levinson, S. C. The myth of language universals: language diversity and its importance for cognitive science. Behav. Brain Sci. 32, 429\u2013448 (2009).","journal-title":"Behav. Brain Sci."},{"key":"2550_CR92","doi-asserted-by":"publisher","unstructured":"arXiv.org submitters. arXiv dataset. kaggle https:\/\/doi.org\/10.34740\/KAGGLE\/DSV\/7548853 (2024).","DOI":"10.34740\/KAGGLE\/DSV\/7548853"},{"key":"2550_CR93","doi-asserted-by":"publisher","unstructured":"Verma, V., Fleisig, E., Tomlin, N. & Klein, D. Ghostbuster: detecting text ghostwritten by large language models. In Proc. 2024 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies1 (eds Duh, K. et al.) 1702\u20131717 https:\/\/doi.org\/10.18653\/v1\/2024.naacl-long.95 (Association for Computational Linguistics, 2024).","DOI":"10.18653\/v1\/2024.naacl-long.95"},{"key":"2550_CR94","doi-asserted-by":"publisher","unstructured":"Adam, G. A. et al. GPTZero: robust detection of LLM-generated texts. Preprint at https:\/\/doi.org\/10.48550\/arXiv.2602.13042 (2026).","DOI":"10.48550\/arXiv.2602.13042"},{"key":"2550_CR95","doi-asserted-by":"publisher","unstructured":"Zhan, H., He, X., Xu, Q., Wu, Y. & Stenetorp, P. G3detector: general GPT-generated text detector. Preprint at https:\/\/doi.org\/10.48550\/arXiv.2305.12680 (2023).","DOI":"10.48550\/arXiv.2305.12680"},{"key":"2550_CR96","first-page":"9051","volume":"32","author":"R Zellers","year":"2019","unstructured":"Zellers, R. et al. Defending against neural fake news. Adv. Neural Inf. Process. Syst. 32, 9051\u20139062 (2019).","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"2550_CR97","doi-asserted-by":"publisher","first-page":"688","DOI":"10.1038\/163688a0","volume":"163","author":"EH SIMPSON","year":"1949","unstructured":"SIMPSON, E. H. Measurement of diversity. Nature 163, 688 (1949).","journal-title":"Nature"},{"key":"2550_CR98","doi-asserted-by":"publisher","first-page":"379","DOI":"10.1002\/j.1538-7305.1948.tb01338.x","volume":"27","author":"CE Shannon","year":"1948","unstructured":"Shannon, C. E. A mathematical theory of communication. Bell Syst. Tech. J. 27, 379\u2013423 (1948).","journal-title":"Bell Syst. Tech. J."},{"key":"2550_CR99","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1016\/S0010-0277(98)00034-1","volume":"68","author":"E Gibson","year":"1998","unstructured":"Gibson, E. Linguistic complexity: locality of syntactic dependencies. Cognition 68, 1\u201376 (1998).","journal-title":"Cognition"},{"key":"2550_CR100","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1037\/h0093508","volume":"56","author":"W Johnson","year":"1944","unstructured":"Johnson, W. Studies in language behavior: a program of research. Psychol. Monogr. 56, 1\u201315 (1944).","journal-title":"Psychol. Monogr."},{"key":"2550_CR101","doi-asserted-by":"publisher","first-page":"264","DOI":"10.1177\/1476993X11398845","volume":"10","author":"H Mardaga","year":"2012","unstructured":"Mardaga, H. Hapax legomena: a neglected field in biblical studies. Curr. Biblic. Res. 10, 264\u2013274 (2012).","journal-title":"Curr. Biblic. Res."},{"key":"2550_CR102","doi-asserted-by":"publisher","first-page":"297","DOI":"10.1007\/BF02310555","volume":"16","author":"LJ Cronbach","year":"1951","unstructured":"Cronbach, L. J. Coefficient alpha and the internal structure of tests. Psychometrika 16, 297\u2013334 (1951).","journal-title":"Psychometrika"},{"key":"2550_CR103","unstructured":"West, S. G., Finch, J. F. & Curran, P. J. in Structural Equation Modeling: Concepts, Issues, and Applications (ed. Hoyle, R. H.) 56\u201375 (Sage, 1995)."},{"key":"2550_CR104","doi-asserted-by":"publisher","first-page":"16\u201329","DOI":"10.1037\/1082-989X.1.1.16","volume":"1","author":"PJ Curran","year":"1996","unstructured":"Curran, P. J., West, S. G. & Finch, J. F. The robustness of test statistics to nonnormality and specification error in confirmatory factor analysis. Psychol. Methods 1, 16\u201329 (1996).","journal-title":"Psychol. Methods"},{"key":"2550_CR105","doi-asserted-by":"publisher","first-page":"961","DOI":"10.1017\/S0266466606000442","volume":"22","author":"L Kilian","year":"2006","unstructured":"Kilian, L. New introduction to multiple time series analysis, by Helmut L\u00fctkepohl, Springer, 2005. Econ. Theory 22, 961\u2013967 (2006).","journal-title":"Econ. Theory"},{"key":"2550_CR106","doi-asserted-by":"crossref","unstructured":"Ivanov, V. & Kilian, L. A practitioner\u2019s guide to lag order selection for VAR impulse response analysis. Stud. Nonlinear Dyn. Econom.9, 1 (2005).","DOI":"10.2202\/1558-3708.1219"},{"key":"2550_CR107","doi-asserted-by":"publisher","first-page":"797","DOI":"10.1007\/s00181-018-1446-3","volume":"56","author":"SB Bruns","year":"2019","unstructured":"Bruns, S. B. & Stern, D. I. Lag length selection and p-hacking in Granger causality testing: prevalence and performance of meta-regression models. Empir. Econ. 56, 797\u2013830 (2019).","journal-title":"Empir. Econ."},{"key":"2550_CR108","first-page":"427","volume":"74","author":"DA Dickey","year":"1979","unstructured":"Dickey, D. A. & Fuller, W. A. Distribution of the estimators for autoregressive time series with a unit root. J. Am. Stat. Assoc. 74, 427\u2013431 (1979).","journal-title":"J. Am. Stat. Assoc."},{"key":"2550_CR109","unstructured":"Box, G. E., Jenkins, G. M., Reinsel, G. C. & Ljung, G. M. Time Series Analysis: Forecasting and Control (John Wiley & Sons, 2015)."},{"key":"2550_CR110","doi-asserted-by":"publisher","unstructured":"Wood, S. N. Generalized Additive Models: An Introduction with R 2nd edn https:\/\/doi.org\/10.1201\/9781315370279 (Chapman and Hall\/CRC, 2017).","DOI":"10.1201\/9781315370279"},{"key":"2550_CR111","first-page":"297","volume":"1","author":"T Hastie","year":"1986","unstructured":"Hastie, T. & Tibshirani, R. Generalized additive models. Stat. Sci. 1, 297\u2013310 (1986).","journal-title":"Stat. Sci."},{"key":"2550_CR112","doi-asserted-by":"publisher","first-page":"86","DOI":"10.1016\/j.wocn.2018.03.002","volume":"70","author":"M Wieling","year":"2018","unstructured":"Wieling, M. Analyzing dynamic phonetic data using generalized additive mixed modeling: a tutorial focusing on articulatory differences between l1 and l2 speakers of English. J. Phon. 70, 86\u2013116 (2018).","journal-title":"J. Phon."},{"key":"2550_CR113","unstructured":"Greene, R. et al. New and Improved Embedding Model https:\/\/openai.com\/index\/new-and-improved-embedding-model (OpenAI, 2024)."},{"key":"2550_CR114","doi-asserted-by":"crossref","unstructured":"Conneau, A. & Kiela, D. SentEval: an evaluation toolkit for universal sentence representations. In Proc. Eleventh International Conference on Language Resources and Evaluation (LREC 2018) (eds Calzolari, N. et al.) https:\/\/aclanthology.org\/L18-1269\/ (European Language Resources Association (ELRA), 2018).","DOI":"10.63317\/2gscircifffd"},{"key":"2550_CR115","doi-asserted-by":"publisher","first-page":"2305016120","DOI":"10.1073\/pnas.2305016120","volume":"120","author":"F Gilardi","year":"2023","unstructured":"Gilardi, F., Alizadeh, M. & Kubli, M. ChatGPT outperforms crowd workers for text-annotation tasks. Proc. Natl Acad. Sci. USA 120, 2305016120 (2023).","journal-title":"Proc. Natl Acad. Sci. USA"},{"key":"2550_CR116","doi-asserted-by":"publisher","first-page":"17","DOI":"10.1007\/s42001-024-00345-9","volume":"8","author":"M Alizadeh","year":"2025","unstructured":"Alizadeh, M. et al. Open-source LLMs for text annotation: a practical guide for model setting and fine-tuning. J. Comput. Soc. Sci. 8, 17 (2025).","journal-title":"J. Comput. Soc. Sci."},{"key":"2550_CR117","doi-asserted-by":"publisher","first-page":"29","DOI":"10.1348\/000711006X126600","volume":"61","author":"KL Gwet","year":"2008","unstructured":"Gwet, K. L. Computing inter-rater reliability and its variance in the presence of high agreement. Br. J. Math. Stat. Psychol. 61, 29\u201348 (2008).","journal-title":"Br. J. Math. Stat. Psychol."},{"key":"2550_CR118","unstructured":"Levene, H. in Contributions to Probability and Statistics (ed. Olkin, I.) 278\u2013292 (Stanford Univ. Press, 1960)."},{"key":"2550_CR119","unstructured":"Gentzkow, M., Shapiro, J. M. & Taddy, M. Congressional Record for the 43rd\u2013114th Congresses: Parsed Speeches and Phrase Counts https:\/\/data.stanford.edu\/congress_text (Stanford Libraries, 2018)."},{"key":"2550_CR120","unstructured":"Silver, N. & Mehta, D. Both Republicans and Democrats Have an Age Problem https:\/\/web.archive.org\/web\/20221216043420\/; https:\/\/fivethirtyeight.com\/features\/both-republicans-and-democrats-have-an-age-problem\/ (2014)."},{"key":"2550_CR121","doi-asserted-by":"publisher","first-page":"1235348","DOI":"10.3389\/fneur.2023.1235348","volume":"14","author":"M Abu Raya","year":"2023","unstructured":"Abu Raya, M. et al. The reciprocal relationship between openness and creativity: from neurobiology to multicultural environments. Front. Neurol. 14, 1235348 (2023).","journal-title":"Front. Neurol."},{"key":"2550_CR122","unstructured":"Goldberg, L. R. in Personality and Personality Disorders 34\u201347 (Routledge, 2013)."},{"key":"2550_CR123","unstructured":"Goldberg, L. Standard Markers of the Big-five Structure (Oregon Research Institute, 1990)."},{"key":"2550_CR124","doi-asserted-by":"publisher","first-page":"1296","DOI":"10.1037\/0022-3514.77.6.1296","volume":"77","author":"JW Pennebaker","year":"1999","unstructured":"Pennebaker, J. W. & King, L. A. Linguistic styles: language use as an individual difference. J. Pers. Soc. Psychol. 77, 1296\u20131312 (1999).","journal-title":"J. Pers. Soc. Psychol."},{"key":"2550_CR125","doi-asserted-by":"crossref","unstructured":"John, O. P., Donahue, E. M. & Kentle, R. L. The Big Five Inventory \u2013versions 4a and 54 Technical report (Univ. California, Berkeley, Institute of Personality and Social Research, 1991).","DOI":"10.1037\/t07550-000"},{"key":"2550_CR126","doi-asserted-by":"publisher","first-page":"2","DOI":"10.1609\/icwsm.v7i2.14467","volume":"7","author":"F Celli","year":"2013","unstructured":"Celli, F., Pianesi, F., Stillwell, D. & Kosinski, M. Workshop on computational personality recognition: shared task. Proc. Int. AAAI Conf. Web Soc. Media 7, 2\u20135 (2013).","journal-title":"Proc. Int. AAAI Conf. Web Soc. Media"},{"key":"2550_CR127","first-page":"85","volume":"10","author":"MH Davis","year":"1980","unstructured":"Davis, M. H. A multidimensional approach to individual differences in empathy. JSAS Cat. Sel. Doc. Psychol. 10, 85 (1980).","journal-title":"JSAS Cat. Sel. Doc. Psychol."},{"key":"2550_CR128","doi-asserted-by":"publisher","unstructured":"Omitaomu, D. et al. Empathic conversations: a multi-level dataset of contextualized conversations. Preprint at https:\/\/doi.org\/10.48550\/arXiv.2205.12698 (2022).","DOI":"10.48550\/arXiv.2205.12698"},{"key":"2550_CR129","doi-asserted-by":"crossref","unstructured":"Barriere, V., Sedoc, J., Tafreshi, S. & Giorgi, S. Findings of WASSA 2023 shared task on empathy, emotion and personality detection in conversation and reactions to news articles. In Proc. 13th Workshop on Computational Approaches to Subjectivity, Sentiment, and Social Media Analysis 511\u2013525 (Association for Computational Linguistics, 2023).","DOI":"10.18653\/v1\/2023.wassa-1.44"},{"key":"2550_CR130","doi-asserted-by":"publisher","unstructured":"Graham, J. et al. Chapter two - moral foundations theory: the pragmatic validity of moral pluralism. Adv. Exp. Soc. Psychol. 47, 55\u2013130 https:\/\/doi.org\/10.1016\/B978-0-12-407236-7.00002-4 (2013).","DOI":"10.1016\/B978-0-12-407236-7.00002-4"},{"key":"2550_CR131","doi-asserted-by":"publisher","first-page":"55","DOI":"10.1162\/0011526042365555","volume":"133","author":"J Haidt","year":"2004","unstructured":"Haidt, J. & Joseph, C. Intuitive ethics: how innately prepared intuitions generate culturally variable virtues. Daedalus 133, 55\u201366 (2004).","journal-title":"Daedalus"},{"key":"2550_CR132","doi-asserted-by":"crossref","unstructured":"Atari, M. et al. Morality beyond the weird: how the nomological network of morality varies across cultures. J. Pers. Soc. Psychol. 125, 1157\u20131188 (2023).","DOI":"10.1037\/pspp0000470"},{"key":"2550_CR133","doi-asserted-by":"publisher","first-page":"366","DOI":"10.1037\/a0021847","volume":"101","author":"J Graham","year":"2011","unstructured":"Graham, J. et al. Mapping the moral domain. J. Pers. Soc. Psychol. 101, 366\u2013385 (2011).","journal-title":"J. Pers. Soc. Psychol."},{"key":"2550_CR134","doi-asserted-by":"publisher","unstructured":"Beltagy, I., Peters, M. E. & Cohan, A. Longformer: the long-document transformer. Preprint at https:\/\/doi.org\/10.48550\/arXiv.2004.05150 (2020).","DOI":"10.48550\/arXiv.2004.05150"},{"key":"2550_CR135","doi-asserted-by":"publisher","first-page":"826","DOI":"10.3758\/s13428-023-02072-x","volume":"56","author":"I Shatz","year":"2024","unstructured":"Shatz, I. Assumption-checking rather than (just) testing: the importance of visualization and effect size in statistical diagnostics. Behav. Res. Methods 56, 826\u2013845 (2024).","journal-title":"Behav. Res. Methods"},{"key":"2550_CR136","doi-asserted-by":"publisher","first-page":"2576","DOI":"10.3758\/s13428-021-01587-5","volume":"53","author":"U Knief","year":"2021","unstructured":"Knief, U. & Forstmeier, W. Violating the normality assumption may be the lesser of two evils. Behav. Res. Methods 53, 2576\u20132590 (2021).","journal-title":"Behav. Res. Methods"},{"key":"2550_CR137","first-page":"3","volume":"8","author":"C Bonferroni","year":"1936","unstructured":"Bonferroni, C. Teoria statistica delle classi e calcolo delle probabilita. Pubbl. R. Inst. Super. Sci. Econ. Commerc. Firenze 8, 3\u201362 (1936).","journal-title":"Pubbl. R. Inst. Super. Sci. Econ. Commerc. Firenze"},{"key":"2550_CR138","doi-asserted-by":"publisher","first-page":"1319","DOI":"10.2466\/pms.1976.43.3f.1319","volume":"43","author":"LL Havlicek","year":"1976","unstructured":"Havlicek, L. L. & Peterson, N. L. Robustness of the Pearson correlation against violations of assumptions. Percept. Mot. Skills 43, 1319\u20131334 (1976).","journal-title":"Percept. Mot. Skills"},{"key":"2550_CR139","doi-asserted-by":"publisher","first-page":"195","DOI":"10.1111\/1468-2389.00172","volume":"9","author":"RR Wilcox","year":"2001","unstructured":"Wilcox, R. R. Modern insights about Pearson\u2019s correlation and least squares regression. Int. J. Sel. Assess. 9, 195\u2013205 (2001).","journal-title":"Int. J. Sel. Assess."},{"key":"2550_CR140","doi-asserted-by":"publisher","first-page":"92","DOI":"10.5334\/irsp.82","volume":"30","author":"M Delacre","year":"2017","unstructured":"Delacre, M., Lakens, D. & Leys, C. Why psychologists should by default use Welch\u2019s t-test instead of Student\u2019s t-test. Int. Rev. Soc. Psychol. 30, 92\u2013101 (2017).","journal-title":"Int. Rev. Soc. Psychol."},{"key":"2550_CR141","doi-asserted-by":"publisher","first-page":"436","DOI":"10.1111\/j.1467-8640.2012.00460.x","volume":"29","author":"SM Mohammad","year":"2013","unstructured":"Mohammad, S. M. & Turney, P. D. Crowdsourcing a word\u2013emotion association lexicon. Comput. Intell. 29, 436\u2013465 (2013).","journal-title":"Comput. Intell."},{"key":"2550_CR142","unstructured":"Mohammad, S. & Turney, P. Emotions evoked by common words and phrases: using Mechanical Turk to create an emotion lexicon. In Proc. NAACL HLT 2010 Workshop on Computational Approaches to Analysis and Generation of Emotion in Text 26\u201334 https:\/\/aclanthology.org\/W10-0204 (Association for Computational Linguistics, 2010)."},{"key":"2550_CR143","unstructured":"Sedoc, J., Buechel, S., Nachmany, Y., Buffone, A. & Ungar, L. Learning word ratings for empathy and distress from document-level user responses. In Proc. Twelfth Language Resources and Evaluation Conference (eds Calzolari, N. et al.) 1664\u20131673 https:\/\/aclanthology.org\/2020.lrec-1.206\/ (European Language Resources Association, 2020)."},{"key":"2550_CR144","doi-asserted-by":"crossref","unstructured":"Baumgartner, J., Zannettou, S., Keegan, B., Squire, M. & Blackburn, J. The pushshift reddit dataset. In Proc. International AAAI Conference on Web and Social Media 14, 830\u2013839 (2020).","DOI":"10.1609\/icwsm.v14i1.7347"},{"key":"2550_CR145","doi-asserted-by":"publisher","unstructured":"Sourati, Z. Code for \"The shrinking landscape of linguistic diversity in the age of large language models\". Zenodo https:\/\/doi.org\/10.5281\/zenodo.21633457 (2026).","DOI":"10.5281\/zenodo.21633457"},{"key":"2550_CR146","doi-asserted-by":"publisher","first-page":"155","DOI":"10.1037\/0033-2909.112.1.155","volume":"112","author":"J Cohen","year":"1992","unstructured":"Cohen, J. A power primer. Psychol. Bull. 112, 155\u2013159 (1992).","journal-title":"Psychol. Bull."}],"container-title":["Nature Human Behaviour"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/www.nature.com\/articles\/s41562-026-02550-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/www.nature.com\/articles\/s41562-026-02550-0","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/www.nature.com\/articles\/s41562-026-02550-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,8,24]],"date-time":"2026-08-24T15:04:16Z","timestamp":1787583856000},"score":1,"resource":{"primary":{"URL":"https:\/\/www.nature.com\/articles\/s41562-026-02550-0"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,8,24]]},"references-count":146,"alternative-id":["2550"],"URL":"https:\/\/doi.org\/10.1038\/s41562-026-02550-0","relation":{},"ISSN":["2397-3374"],"issn-type":[{"value":"2397-3374","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,8,24]]},"assertion":[{"value":"17 February 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"10 July 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"24 August 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"The authors declare no competing interests.","order":1,"name":"Ethics","label":"Competing interests","group":{"name":"EthicsHeading","label":"Ethics"}}]}}