{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,22]],"date-time":"2026-07-22T16:31:29Z","timestamp":1784737889064,"version":"3.55.0"},"reference-count":55,"publisher":"Springer Science and Business Media LLC","issue":"2","license":[{"start":{"date-parts":[[2026,1,29]],"date-time":"2026-01-29T00:00:00Z","timestamp":1769644800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,1,29]],"date-time":"2026-01-29T00:00:00Z","timestamp":1769644800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int. J. Mach. Learn. &amp; Cyber."],"published-print":{"date-parts":[[2026,2]]},"DOI":"10.1007\/s13042-025-02868-7","type":"journal-article","created":{"date-parts":[[2026,1,29]],"date-time":"2026-01-29T14:39:58Z","timestamp":1769697598000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["Detecting AI-generated text in high-resource languages: developing a RoBERTa-CNN hybrid model for academic integrity challenge"],"prefix":"10.1007","volume":"17","author":[{"given":"Manish","family":"Prajapati","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Santos Kumar","family":"Baliarsingh","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Prabhu Prasad","family":"Dev","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,1,29]]},"reference":[{"key":"2868_CR1","unstructured":"Fitria TN (2021) Artificial intelligence (ai) in education: Using ai tools for teaching and learning process. In: Prosiding Seminar Nasional & Call for Paper STIE AAS, vol. 4, pp. 134\u2013147"},{"key":"2868_CR2","doi-asserted-by":"crossref","unstructured":"Bengesi S, El-Sayed H, Sarker MK, Houkpati Y, Irungu J, Oladunni T (2024) Advancements in generative ai: A comprehensive review of gans, gpt, autoencoders, diffusion model, and transformers. IEEE Access","DOI":"10.1109\/ACCESS.2024.3397775"},{"key":"2868_CR3","doi-asserted-by":"crossref","unstructured":"Gehrmann S, Strobelt H, Rush AM (2019) Gltr: Statistical detection and visualization of generated text. arXiv preprint arXiv:1906.04043","DOI":"10.18653\/v1\/P19-3019"},{"key":"2868_CR4","unstructured":"Mohamadi S, Mujtaba G, Le N, Doretto G, Adjeroh DA (2023) Chatgpt in the age of generative ai and large language models: a concise survey. arXiv preprint arXiv:2307.04251"},{"key":"2868_CR5","doi-asserted-by":"crossref","unstructured":"Prajapati M, Baliarsingh SK, Dora C, Bhoi A, Hota J, Mohanty JP (2024) Detection of ai-generated text using large language model. In: 2024 international conference on emerging systems and intelligent computing (ESIC), pp. 735\u2013740. IEEE","DOI":"10.1109\/ESIC60604.2024.10481602"},{"key":"2868_CR6","unstructured":"Liu Y, Ott M, Goyal N, Du J, Joshi M, Chen D, Levy O, Lewis M, Zettlemoyer L, Stoyanov V (2019) Roberta: A robustly optimized bert pretraining approach. arXiv preprint arXiv:1907.11692"},{"key":"2868_CR7","doi-asserted-by":"publisher","first-page":"112018","DOI":"10.1016\/j.asoc.2024.112018","volume":"164","author":"AJ Lak","year":"2024","unstructured":"Lak AJ, Boostani R, Alenizi FA, Mohammed AS, Fakhrahmad SM (2024) Roberta, resnext and bilstm with self-attention: the ultimate trio for customer sentiment analysis. Appl Soft Comput 164:112018","journal-title":"Appl Soft Comput"},{"key":"2868_CR8","doi-asserted-by":"publisher","first-page":"100050","DOI":"10.1016\/j.nlp.2023.100050","volume":"6","author":"A Singh","year":"2024","unstructured":"Singh A, Sharma D, Nandy A, Singh VK (2024) Towards a large sized curated and annotated corpus for discriminating between human written and ai generated texts: a case study of text sourced from wikipedia and chatgpt. Natural Language Process J 6:100050","journal-title":"Natural Language Process J"},{"key":"2868_CR9","doi-asserted-by":"crossref","unstructured":"Mihir TK, Harsha, KVS, Nitya SY, Krishna GB, Anamalamudi S, Sivarajan S (2024) Machine learning approaches to identify ai-generated text: a comparative analysis. In: 2024 international conference on intelligent computing and emerging communication technologies (ICEC), pp. 1\u20136. IEEE","DOI":"10.1109\/ICEC59683.2024.10837481"},{"issue":"1","key":"2868_CR10","doi-asserted-by":"publisher","first-page":"66","DOI":"10.37396\/jsc.v7i1.388","volume":"7","author":"DS Rahayu","year":"2024","unstructured":"Rahayu DS, Novita R, Ahsyar TK (2024) Sentiment analysis chatgpt using the multinominal na\u00efve bayes classifier (nbc) algorithm. J Sistem Cerdas 7(1):66\u201374","journal-title":"J Sistem Cerdas"},{"issue":"1","key":"2868_CR11","doi-asserted-by":"publisher","first-page":"2434301","DOI":"10.1080\/08839514.2024.2434301","volume":"38","author":"W Tao","year":"2024","unstructured":"Tao W, Wang L, Meng Q, Li R, Han P, Shi Y, Shan L, Geng X (2024) Text-to-text transfer transformer based method for generating startup scenarios for new equipment in power grids. Appl Artif Intell 38(1):2434301","journal-title":"Appl Artif Intell"},{"key":"2868_CR12","doi-asserted-by":"crossref","unstructured":"Fabregas AC, Arellano PBV, Pinili AND (2020) Long-short term memory (lstm) networks with time series and spatio-temporal approaches applied in forecasting earthquakes in the philippines. In: Proceedings of the 4th international conference on natural language processing and information retrieval, pp. 188\u2013193","DOI":"10.1145\/3443279.3443288"},{"key":"2868_CR13","doi-asserted-by":"crossref","unstructured":"Kollar J, Alshibli M (2024) An overview of artificial intelligence\u2019s accuracy. In: 2024 ieee long island systems, applications and technology conference (LISAT), pp. 1\u20138. IEEE","DOI":"10.1109\/LISAT63094.2024.10808042"},{"key":"2868_CR14","doi-asserted-by":"crossref","unstructured":"Shahriar A, Pandit D, Rahman MS (2024) Xlnet-cnn: Combining global context understanding of xlnet with local context capture through convolution for improved multi-label text classification. In: Proceedings of the 11th international conference on networking, systems, and security, pp. 24\u201331","DOI":"10.1145\/3704522.3704540"},{"key":"2868_CR15","doi-asserted-by":"crossref","unstructured":"Athiwaratkun B, Stokes JW (2017) Malware classification with lstm and gru language models and a character-level cnn. In: 2017 IEEE international conference on acoustics, speech and signal processing (ICASSP), pp. 2482\u20132486. IEEE","DOI":"10.1109\/ICASSP.2017.7952603"},{"key":"2868_CR16","doi-asserted-by":"crossref","unstructured":"Kumarage T, Sheth P, Moraffah R, Garland J, Liu H (2023) How reliable are ai-generated-text detectors? an assessment framework using evasive soft prompts. arXiv preprint arXiv:2310.05095","DOI":"10.18653\/v1\/2023.findings-emnlp.94"},{"key":"2868_CR17","doi-asserted-by":"crossref","unstructured":"Habibzadeh F (2023) Gptzero performance in identifying artificial intelligence-generated medical texts: a preliminary study. J Korean Med Sci 38(38)","DOI":"10.3346\/jkms.2023.38.e319"},{"key":"2868_CR18","unstructured":"Mitchell E, Lee Y, Khazatsky A, Manning CD, Finn C (2023) Detectgpt: Zero-shot machine-generated text detection using probability curvature. arXiv preprint arXiv:2301.11305"},{"issue":"1","key":"2868_CR19","first-page":"10","volume":"1","author":"LI Martins","year":"2024","unstructured":"Martins LI, Wonu N, Victor-Edema UA (2024) Evaluating the efficacy of ai-detection tools in assessing human and ai-generated content variants. Faculty Natural Appl Sci J Comput Appl 1(1):10\u201316","journal-title":"Faculty Natural Appl Sci J Comput Appl"},{"key":"2868_CR20","unstructured":"Bao G, Zhao Y, Teng Z, Yang L, Zhang Y (2023) Fast-detectgpt: Efficient zero-shot detection of machine-generated text via conditional probability curvature. arXiv preprint arXiv:2310.05130"},{"key":"2868_CR21","unstructured":"Sadasivan VS, Kumar A, Balasubramanian S, Wang W, Feizi S (2023) Can ai-generated text be reliably detected? arXiv preprint arXiv:2303.11156"},{"key":"2868_CR22","doi-asserted-by":"crossref","unstructured":"Goldberg Y, Hirst, G (2017) Neural network methods in natural language processing. morgan & claypool publishers (2017). zitiert auf Seite 69","DOI":"10.1007\/978-3-031-02165-7"},{"key":"2868_CR23","unstructured":"Mikolov T, Chen K, Corrado G, Dean J (2013) Efficient estimation of word representations in vector space. arXiv preprint arXiv:1301.3781"},{"key":"2868_CR24","doi-asserted-by":"crossref","unstructured":"Ling W, Dyer C, Black AW, Trancoso I (2015) Two\/too simple adaptations of word2vec for syntax problems. In: Proceedings of the 2015 Conference of the North American chapter of the association for computational linguistics: human language technologies, pp. 1299\u20131304","DOI":"10.3115\/v1\/N15-1142"},{"key":"2868_CR25","unstructured":"Neumann M, Iyyer M, Gardner M, Clark C, Lee K, Zettlemoyer L (2018) Deep contextualized word representations. arXiv preprint arXiv:1802.05365"},{"key":"2868_CR26","doi-asserted-by":"crossref","unstructured":"Howard J, Ruder S (2018) Universal language model fine-tuning for text classification. arXiv preprint arXiv:1801.06146","DOI":"10.18653\/v1\/P18-1031"},{"key":"2868_CR27","unstructured":"Vaswani A, Shazeer N, Parmar N, Uszkoreit J, Jones L, Gomez AN, Kaiser \u0141, Polosukhin I (2017) Attention is all you need. Adv Neural Inform Process Syst 30"},{"key":"2868_CR28","unstructured":"Devlin J, Chang M-W, Lee K, Toutanova K (2018) Bert: Pre-training of deep bidirectional transformers for language understanding. arXiv preprint arXiv:1810.04805"},{"issue":"140","key":"2868_CR29","first-page":"1","volume":"21","author":"C Raffel","year":"2020","unstructured":"Raffel C, Shazeer N, Roberts A, Lee K, Narang S, Matena M, Zhou Y, Li W, Liu PJ (2020) Exploring the limits of transfer learning with a unified text-to-text transformer. J Mach Learn Res 21(140):1\u201367","journal-title":"J Mach Learn Res"},{"key":"2868_CR30","unstructured":"Radford A, Narasimhan K, Salimans T, Sutskever I (2018) Improving language understanding with unsupervised learning. 2018. URL: https:\/\/openaicom\/research\/language-unsupervised"},{"key":"2868_CR31","doi-asserted-by":"crossref","unstructured":"Lewis M, Liu Y, Goyal N, Ghazvininejad M, Mohamed A, Levy O, Stoyanov V, Zettlemoyer L (2019) Bart: Denoising sequence-to-sequence pre-training for natural language generation, translation, and comprehension. arXiv preprint arXiv:1910.13461","DOI":"10.18653\/v1\/2020.acl-main.703"},{"key":"2868_CR32","unstructured":"Hou X, Zhao Y, Liu Y, Yang Z, Wang K, Li L, Luo X, Lo D, Grundy J, Wang H (2023) Large language models for software engineering: A systematic literature review. arXiv preprint arXiv:2308.10620"},{"key":"2868_CR33","first-page":"1877","volume":"33","author":"T Brown","year":"2020","unstructured":"Brown T, Mann B, Ryder N, Subbiah M, Kaplan JD, Dhariwal P, Neelakantan A, Shyam P, Sastry G, Askell A (2020) Language models are few-shot learners. Adv Neural Inf Process Syst 33:1877\u20131901","journal-title":"Adv Neural Inf Process Syst"},{"key":"2868_CR34","unstructured":"Achiam J, Adler S, Agarwal S, Ahmad L, Akkaya I, Aleman FL, Almeida D, Altenschmidt J, Altman S, Anadkat S, et al (2023) Gpt-4 technical report. arXiv preprint arXiv:2303.08774"},{"key":"2868_CR35","unstructured":"Touvron H, Lavril T, Izacard G, Martinet X, Lachaux M-A, Lacroix T, Rozi\u00e8re B, Goyal N, Hambro E, Azhar F, et al (2023) Llama: Open and efficient foundation language models. arXiv preprint arXiv:2302.13971"},{"key":"2868_CR36","doi-asserted-by":"crossref","unstructured":"Reynolds L, McDonell K (2021) Prompt programming for large language models: Beyond the few-shot paradigm. In: Extended abstracts of the 2021 CHI conference on human factors in computing systems, pp. 1\u20137","DOI":"10.1145\/3411763.3451760"},{"issue":"11","key":"2868_CR37","doi-asserted-by":"publisher","first-page":"2921","DOI":"10.14778\/3551793.3551841","volume":"15","author":"I Trummer","year":"2022","unstructured":"Trummer I (2022) Codexdb: Synthesizing code for query processing from natural language instructions using gpt-3 codex. Proceedings of the VLDB Endowment 15(11):2921\u20132928","journal-title":"Proceedings of the VLDB Endowment"},{"key":"2868_CR38","unstructured":"White J, Fu Q, Hays S, Sandborn M, Olea C, Gilbert H, Elnashar A, Spencer-Smith J, Schmidt DC (2023) A prompt pattern catalog to enhance prompt engineering with chatgpt. arXiv preprint arXiv:2302.11382"},{"issue":"1","key":"2868_CR39","doi-asserted-by":"publisher","first-page":"33","DOI":"10.1007\/s10916-023-01925-4","volume":"47","author":"M Cascella","year":"2023","unstructured":"Cascella M, Montomoli J, Bellini V, Bignami E (2023) Evaluating the feasibility of chatgpt in healthcare: an analysis of multiple clinical and research scenarios. J Med Syst 47(1):33","journal-title":"J Med Syst"},{"issue":"6","key":"2868_CR40","doi-asserted-by":"publisher","first-page":"100936","DOI":"10.1016\/j.ajogmf.2023.100936","volume":"5","author":"G Levin","year":"2023","unstructured":"Levin G, Meyer R, Kadoch E, Brezinov Y (2023) Identifying chatgpt-written obgyn abstracts using a simple tool. Am J Obstetr Gynecol MFM 5(6):100936","journal-title":"Am J Obstetr Gynecol MFM"},{"issue":"7","key":"2868_CR41","doi-asserted-by":"publisher","first-page":"587","DOI":"10.1093\/asj\/sjad042","volume":"43","author":"R Gupta","year":"2023","unstructured":"Gupta R, Pande P, Herzog I, Weisberger J, Chao J, Chaiyasate K, Lee ES (2023) Application of chatgpt in cosmetic plastic surgery: ally or antagonist? Aesthetic Surg J 43(7):587\u2013590","journal-title":"Aesthetic Surg J"},{"issue":"1","key":"2868_CR42","doi-asserted-by":"publisher","first-page":"4164","DOI":"10.1038\/s41598-023-31412-2","volume":"13","author":"A Lahat","year":"2023","unstructured":"Lahat A, Shachar E, Avidan B, Shatz Z, Glicksberg BS, Klang E (2023) Evaluating the use of large language model in identifying top research questions in gastroenterology. Sci Rep 13(1):4164","journal-title":"Sci Rep"},{"key":"2868_CR43","doi-asserted-by":"crossref","unstructured":"Lyu Q, Tan J, Zapadka ME, Ponnatapura J, Niu C, Myers KJ, Wang G, Whitlow CT (2023) Translating radiology reports into plain language using chatgpt and gpt-4 with prompt learning: Promising results, limitations, and potential. arXiv preprint arXiv:2303.09038","DOI":"10.1186\/s42492-023-00136-5"},{"key":"2868_CR44","doi-asserted-by":"crossref","unstructured":"Thorp HH (2023) ChatGPT is fun, but not an author. American Association for the Advancement of Science","DOI":"10.1126\/science.adg7879"},{"issue":"7947","key":"2868_CR45","doi-asserted-by":"publisher","first-page":"224","DOI":"10.1038\/d41586-023-00288-7","volume":"614","author":"EA Van Dis","year":"2023","unstructured":"Van Dis EA, Bollen J, Zuidema W, Van Rooij R, Bockting CL (2023) Chatgpt: five priorities for research. Nature 614(7947):224\u2013226","journal-title":"Nature"},{"issue":"8","key":"2868_CR46","first-page":"9","volume":"1","author":"A Radford","year":"2019","unstructured":"Radford A, Wu J, Child R, Luan D, Amodei D, Sutskever I (2019) Language models are unsupervised multitask learners. OpenAI blog 1(8):9","journal-title":"OpenAI blog"},{"key":"2868_CR47","doi-asserted-by":"crossref","unstructured":"Bender EM, Gebru T, McMillan-Major A, Shmitchell S (2021) On the dangers of stochastic parrots: Can language models be too big? In: Proceedings of the 2021 ACM Conference on Fairness, Accountability, and Transparency, pp. 610\u2013623","DOI":"10.1145\/3442188.3445922"},{"key":"2868_CR48","unstructured":"Bommasani R, Hudson DA, Adeli E, Altman R, Arora S, Arx S, Bernstein MS, Bohg J, Bosselut A, Brunskill E, et al (2021) On the opportunities and risks of foundation models. arXiv preprint arXiv:2108.07258"},{"key":"2868_CR49","doi-asserted-by":"crossref","unstructured":"Patel A, Raffel C, Callison-Burch C (2024) Datadreamer: A tool for synthetic data generation and reproducible llm workflows. arXiv preprint arXiv:2402.10379","DOI":"10.18653\/v1\/2024.acl-long.208"},{"issue":"4","key":"2868_CR50","doi-asserted-by":"publisher","first-page":"102744","DOI":"10.1016\/j.dsx.2023.102744","volume":"17","author":"R Vaishya","year":"2023","unstructured":"Vaishya R, Misra A, Vaish A (2023) Chatgpt: is this version good for healthcare and research? Diabetes Metab Syndrome Clin Res Rev 17(4):102744","journal-title":"Diabetes Metab Syndrome Clin Res Rev"},{"key":"2868_CR51","unstructured":"Yang Z (2019) Xlnet: Generalized autoregressive pretraining for language understanding. arXiv preprint arXiv:1906.08237"},{"key":"2868_CR52","doi-asserted-by":"crossref","unstructured":"Huang Z, Liang D, Xu P, Xiang B (2020) Improve transformer models with better relative position embeddings. arXiv preprint arXiv:2009.13658","DOI":"10.18653\/v1\/2020.findings-emnlp.298"},{"key":"2868_CR53","doi-asserted-by":"crossref","unstructured":"Liu F, Vuli\u0107 I, Korhonen A, Collier N (2021) Fast, effective, and self-supervised: Transforming masked language models into universal lexical and sentence encoders. arXiv preprint arXiv:2104.08027","DOI":"10.18653\/v1\/2021.emnlp-main.109"},{"key":"2868_CR54","unstructured":"Kaggle. https:\/\/www.kaggle.com\/competitions\/llm-detect-ai-generated-text\/data. Accessed: 2025-03-05"},{"key":"2868_CR55","doi-asserted-by":"publisher","first-page":"35239","DOI":"10.1007\/s11042-020-10082-6","volume":"80","author":"U Naseem","year":"2021","unstructured":"Naseem U, Razzak I, Eklund PW (2021) A survey of pre-processing techniques to improve short-text quality: a case study on hate speech detection on twitter. Multimedia Tools Appl 80:35239\u201335266","journal-title":"Multimedia Tools Appl"}],"container-title":["International Journal of Machine Learning and Cybernetics"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13042-025-02868-7.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s13042-025-02868-7","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13042-025-02868-7.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,3,16]],"date-time":"2026-03-16T09:54:30Z","timestamp":1773654870000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s13042-025-02868-7"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,1,29]]},"references-count":55,"journal-issue":{"issue":"2","published-print":{"date-parts":[[2026,2]]}},"alternative-id":["2868"],"URL":"https:\/\/doi.org\/10.1007\/s13042-025-02868-7","relation":{},"ISSN":["1868-8071","1868-808X"],"issn-type":[{"value":"1868-8071","type":"print"},{"value":"1868-808X","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,1,29]]},"assertion":[{"value":"4 February 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"11 November 2025","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"29 January 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no Conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"40"}}