{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,20]],"date-time":"2026-07-20T09:03:49Z","timestamp":1784538229937,"version":"3.55.0"},"reference-count":39,"publisher":"Springer Science and Business Media LLC","issue":"6","license":[{"start":{"date-parts":[[2026,3,5]],"date-time":"2026-03-05T00:00:00Z","timestamp":1772668800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by-nc-nd\/4.0"},{"start":{"date-parts":[[2026,7,20]],"date-time":"2026-07-20T00:00:00Z","timestamp":1784505600000},"content-version":"vor","delay-in-days":137,"URL":"https:\/\/creativecommons.org\/licenses\/by-nc-nd\/4.0"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62462010"],"award-info":[{"award-number":["62462010"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"name":"Guizhou Provincial Science and Technology Plan, China","award":["QianKe He Zhongda Zhuanxiang Zi[2024]003"],"award-info":[{"award-number":["QianKe He Zhongda Zhuanxiang Zi[2024]003"]}]},{"name":"Support for High-level Universities Outside the Province to Provide \u201cGroup-style\u201d Assistance to Guizhou Universities in Discipline Construction and Scientific Research Project","award":["QianKe He Rencai XKBF [2025]016"],"award-info":[{"award-number":["QianKe He Rencai XKBF [2025]016"]}]},{"name":"Big Data Security and Network Security Innovation Team of Guizhou Provincial High Education Institution, China","award":["[2023]052"],"award-info":[{"award-number":["[2023]052"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["J. King Saud Univ. Comput. Inf. Sci."],"published-print":{"date-parts":[[2026,8]]},"DOI":"10.1007\/s44443-026-00541-9","type":"journal-article","created":{"date-parts":[[2026,3,5]],"date-time":"2026-03-05T13:46:58Z","timestamp":1772718418000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Mitigating the impact of activation smoothing on weights: A channel permutation and smoothing-based LLMs quantization method"],"prefix":"10.1007","volume":"38","author":[{"given":"Yanyao","family":"Guan","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-0880-675X","authenticated-orcid":false,"given":"Yunhe","family":"Cui","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Langtao","family":"Zhao","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Guowei","family":"Shen","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yi","family":"Chen","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Chun","family":"Guo","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,3,5]]},"reference":[{"key":"541_CR1","doi-asserted-by":"crossref","unstructured":"Ashkboos S, Mohtashami A, Croci ML et al (2024) Quarot: Outlier-free 4-bit inference in rotated llms. Adv Neural Inf Process Syst 37:100213\u2013100240","DOI":"10.52202\/079017-3180"},{"key":"541_CR2","doi-asserted-by":"publisher","unstructured":"Bisk Y, Zellers R, Bras RL, et al (2020) PIQA: reasoning about physical commonsense in natural language. In: Proceedings of the thirty-fourth AAAI conference on artificial intelligence (AAAI 2020). AAAI Press, pp 7432\u20137439. https:\/\/doi.org\/10.1609\/AAAI.V34I05.6239","DOI":"10.1609\/AAAI.V34I05.6239"},{"key":"541_CR3","unstructured":"Chee J, Cai Y, Kuleshov V, et al (2023) Quip: 2-bit quantization of large language models with guarantees. In: Advances in neural information processing systems 36, NeurIPS 2023, New Orleans, LA, USA, December 10\u201316, 2023. http:\/\/papers.nips.cc\/paper_files\/paper\/2023\/hash\/0df38cd13520747e1e64e5b123a78ef8-Abstract-Conference.html"},{"key":"541_CR4","unstructured":"Clark P, Cowhey I, Etzioni O, et al (2018) Think you have solved question answering? try arc, the AI2 reasoning challenge. arxiv:1803.05457"},{"key":"541_CR5","unstructured":"Dettmers T, Lewis M, Belkada Y, et al (2022) Gpt3.int8(): 8-bit matrix multiplication for transformers at scale. In: Advances in neural information processing systems 35: Annual Conference on Neural Information Processing Systems 2022, NeurIPS 2022, New Orleans, LA, USA, November 28 - December 9, 2022, http:\/\/papers.nips.cc\/paper_files\/paper\/2022\/hash\/c3ba4962c05c49636d4c6206a97e9c8a-Abstract-Conference.html"},{"key":"541_CR6","unstructured":"Frantar E, Ashkboos S, Hoefler T, et al (2022) Gptq: Accurate post-training quantization for generative pre-trained transformers. arXiv:2210.17323"},{"key":"541_CR7","doi-asserted-by":"crossref","unstructured":"Frantar E, Castro RL, Chen J, et al (2024) Marlin: Mixed-precision auto-regressive parallel inference on large language models. arXiv:2408.11743","DOI":"10.1145\/3710848.3710871"},{"key":"541_CR8","unstructured":"Gao L, Biderman S, Black S, et al (2021) The pile: An 800gb dataset of diverse text for language modeling. arxiv:2101.00027"},{"key":"541_CR9","unstructured":"Gao L, Tow J, Abbasi B et al (2024) A framework for few-shot language model evaluation"},{"key":"541_CR10","doi-asserted-by":"publisher","unstructured":"Guo C, Tang J, Hu W, et al (2023) Olive: Accelerating large language models via hardware-friendly outlier-victim pair quantization. In: Proceedings of the 50th annual international symposium on computer architecture, ISCA 2023, Orlando, FL, USA, June 17\u201321, 2023. https:\/\/doi.org\/10.1145\/3579371.3589038","DOI":"10.1145\/3579371.3589038"},{"key":"541_CR11","unstructured":"Hu X, Cheng Y, Yang D, et al (2025) Ostquant: Refining large language model quantization with orthogonal and scaling transformations for better distribution fitting. arXiv:2501.13987"},{"key":"541_CR12","doi-asserted-by":"publisher","unstructured":"Jacob B, Kligys S, Chen B, et al (2018) Quantization and training of neural networks for efficient integer-arithmetic-only inference. In: 2018 IEEE conference on computer vision and pattern recognition (CVPR 2018), pp 2704\u20132713, https:\/\/doi.org\/10.1109\/CVPR.2018.00286, http:\/\/openaccess.thecvf.com\/content_cvpr_2018\/html\/Jacob_Quantization_and_Training_CVPR_2018_paper.html","DOI":"10.1109\/CVPR.2018.00286"},{"key":"541_CR13","doi-asserted-by":"crossref","unstructured":"Kim S, Choi Y, Oh J, et al (2025) Lightrot: A light-weighted rotation scheme and architecture for accurate low-bit large language model inference. IEEE J Emerging Sel Top Circ Syst","DOI":"10.1109\/JETCAS.2025.3558300"},{"key":"541_CR14","doi-asserted-by":"publisher","unstructured":"Kim YJ, Henry R, Fahim R, et al (2022) Who says elephants can\u2019t run: Bringing large scale moe models into cloud scale production. https:\/\/doi.org\/10.48550\/ARXIV.2211.10017, arxiv:2211.10017","DOI":"10.48550\/ARXIV.2211.10017"},{"key":"541_CR15","doi-asserted-by":"publisher","unstructured":"Lee C, Jin J, Kim T, et al (2024) OWQ: outlier-aware weight quantization for efficient fine-tuning and inference of large language models. In: Thirty-Eighth AAAI Conference on Artificial Intelligence, AAAI 2024, February 20\u201327, 2024, Vancouver, Canada. https:\/\/doi.org\/10.1609\/AAAI.V38I12.29237","DOI":"10.1609\/AAAI.V38I12.29237"},{"key":"541_CR16","unstructured":"Lee D, Han S, Cho T, et al (2023) SPQR: controlling q-ensemble independence with spiked random model for reinforcement learning. In: Advances in Neural Information Processing Systems 36, NeurIPS 2023, New Orleans, LA, USA, December 10\u201316, 2023. http:\/\/papers.nips.cc\/paper_files\/paper\/2023\/hash\/cdcaf772b4f8eda0385d0930517de64a-Abstract-Conference.html"},{"key":"541_CR17","doi-asserted-by":"publisher","first-page":"87766","DOI":"10.52202\/079017-2786","volume":"37","author":"H Lin","year":"2024","unstructured":"Lin H, Xu H, Wu Y et al (2024) Duquant: Distributing outliers via dual transformation makes stronger quantized llms. Adv Neural Inf Process Syst 37:87766\u201387800","journal-title":"Adv Neural Inf Process Syst"},{"key":"541_CR18","unstructured":"Lin J, Tang J, Tang H, et al (2024b) AWQ: activation-aware weight quantization for on-device LLM compression and acceleration. In: Proceedings of the seventh annual conference on machine learning and systems, MLSys 2024, Santa Clara, CA, USA, May 13\u201316, 2024. https:\/\/proceedings.mlsys.org\/paper_files\/paper\/2024\/hash\/42a452cbafa9dd64e9ba4aa95cc1ef21-Abstract-Conference.html"},{"key":"541_CR19","unstructured":"Liu J, Gong R, Wei X, et al (2024a) QLLM: accurate and efficient low-bitwidth quantization for large language models. In: The twelfth international conference on learning representations, ICLR 2024, Vienna, Austria, May 7\u201311, 2024. https:\/\/openreview.net\/forum?id=FIplmUWdm3"},{"key":"541_CR20","unstructured":"Liu Z, Zhao C, Fedorov I, et al (2024b) Spinquant: Llm quantization with learned rotations. arXiv:2405.16406"},{"key":"541_CR21","unstructured":"Merity S, Xiong C, Bradbury J, et al (2016) Pointer sentinel mixture models. arxiv:1609.07843"},{"key":"541_CR22","unstructured":"Raffel C, Shazeer N, Roberts A, et al (2020) Exploring the limits of transfer learning with a unified text-to-text transformer. J Mach Learn Res 21(140), 1\u201367. http:\/\/jmlr.org\/papers\/v21\/20-074.html"},{"key":"541_CR23","doi-asserted-by":"publisher","unstructured":"Rajpurkar P, Jia R, Liang P (2018) Know what you don\u2019t know: Unanswerable questions for SQuAD. In: Gurevych I, Miyao Y (eds) Proceedings of the 56th annual meeting of the association for computational linguistics (Volume 2: Short Papers). Association for Computational Linguistics, Melbourne, Australia, pp 784\u2013789. https:\/\/doi.org\/10.18653\/v1\/P18-2124, https:\/\/aclanthology.org\/P18-2124","DOI":"10.18653\/v1\/P18-2124"},{"key":"541_CR24","doi-asserted-by":"publisher","unstructured":"Sakaguchi K, Bras RL, Bhagavatula C, et al (2020) Winogrande: An adversarial winograd schema challenge at scale. In: Proceedings of the 34th AAAI conference on artificial intelligence (AAAI). AAAI Press, pp 8732\u20138740. https:\/\/doi.org\/10.1609\/AAAI.V34I05.6399","DOI":"10.1609\/AAAI.V34I05.6399"},{"key":"541_CR25","doi-asserted-by":"publisher","unstructured":"See A, Liu PJ, Manning CD (2017) Get to the point: Summarization with pointer-generator networks. In: Proceedings of the 55th annual meeting of the association for computational linguistics (Volume 1: Long Papers). Association for Computational Linguistics, Vancouver, Canada, pp 1073\u20131083. https:\/\/doi.org\/10.18653\/v1\/P17-1099, https:\/\/www.aclweb.org\/anthology\/P17-1099","DOI":"10.18653\/v1\/P17-1099"},{"key":"541_CR26","unstructured":"Shao W, Chen M, Zhang Z, et al (2024) Omniquant: Omnidirectionally calibrated quantization for large language models. In: The Twelfth international conference on learning representations, ICLR 2024, Vienna, Austria, May 7\u201311, 2024. https:\/\/openreview.net\/forum?id=8Wuvhh0LYW"},{"key":"541_CR27","doi-asserted-by":"publisher","unstructured":"Shen S, Dong Z, Ye J, et al (2020) Q-bert: Hessian based ultra low precision quantization of bert. In: Proceedings of the thirty-fourth AAAI conference on artificial intelligence (AAAI). AAAI Press, pp 8815\u20138821. https:\/\/doi.org\/10.1609\/AAAI.V34I05.6409","DOI":"10.1609\/AAAI.V34I05.6409"},{"key":"541_CR28","doi-asserted-by":"crossref","unstructured":"Socher R, Perelygin A, Wu J, et al (2013) Recursive deep models for semantic compositionality over a sentiment treebank. In: Proceedings of the 2013 conference on empirical methods in natural language processing. Association for Computational Linguistics, Seattle, Washington, USA, pp 1631\u20131642. https:\/\/www.aclweb.org\/anthology\/D13-1170","DOI":"10.18653\/v1\/D13-1170"},{"key":"541_CR29","doi-asserted-by":"publisher","unstructured":"Touvron H, Lavril T, Izacard G, et al (2023) LLaMA: open and efficient foundation language models. https:\/\/doi.org\/10.48550\/ARXIV.2302.13971, arxiv:2302.13971","DOI":"10.48550\/ARXIV.2302.13971"},{"key":"541_CR30","unstructured":"Tseng A, Chee J, Sun Q, et al (2024) Quip $$\\# $$: Even better llm quantization with hadamard incoherence and lattice codebooks. In: International conference on machine learning, PMLR, pp 48630\u201348656"},{"key":"541_CR31","unstructured":"Vaswani A, Shazeer N, Parmar N, et al (2017) Attention is all you need. In: Advances\u00a0in neural information processing systems (NeurIPS), pp 5998\u20136008. https:\/\/proceedings.neurips.cc\/paper\/2017\/hash\/3f5ee243547dee91fbd053c1c4a845aa-Abstract.html"},{"key":"541_CR32","doi-asserted-by":"crossref","unstructured":"Wang A, Singh A, Michael J, et al (2019) GLUE: A multi-task benchmark and analysis platform for natural language understanding. In the Proceedings of ICLR","DOI":"10.18653\/v1\/W18-5446"},{"key":"541_CR33","doi-asserted-by":"publisher","unstructured":"Wei X, Zhang Y, Li Y, et al (2023) Outlier suppression+: Accurate quantization of large language models by equivalent and effective shifting and scaling. In: Proceedings of the 2023 conference on empirical methods in natural language processing, EMNLP 2023, Singapore, December 6\u201310, 2023. https:\/\/doi.org\/10.18653\/V1\/2023.EMNLP-MAIN.102","DOI":"10.18653\/V1\/2023.EMNLP-MAIN.102"},{"key":"541_CR34","unstructured":"Xiao G, Lin J, Seznec M, et al (2023) Smoothquant: Accurate and efficient post-training quantization for large language models. In: International conference on machine learning, ICML 2023, 23\u201329 July 2023, Honolulu, Hawaii, USA. https:\/\/proceedings.mlr.press\/v202\/xiao23c.html"},{"key":"541_CR35","unstructured":"Yao Z, Aminabadi RY, Zhang M, et al (2022) Zeroquant: Efficient and affordable post-training quantization for large-scale transformers. In: Advances in neural information processing systems 35, NeurIPS 2022, New Orleans, LA, USA, November 28 - December 9, 2022. http:\/\/papers.nips.cc\/paper_files\/paper\/2022\/hash\/adf7fa39d65e2983d724ff7da57f00ac-Abstract-Conference.html"},{"key":"541_CR36","doi-asserted-by":"publisher","unstructured":"Yuan Z, Niu L, Liu J, et al (2023) RPTQ: reorder-based post-training quantization for large language models. https:\/\/doi.org\/10.48550\/ARXIV.2304.01089, arxiv:2304.01089","DOI":"10.48550\/ARXIV.2304.01089"},{"key":"541_CR37","doi-asserted-by":"publisher","unstructured":"Zellers R, Holtzman A, Bisk Y, et al (2019) Hellaswag: Can a machine really finish your sentence? In: Proceedings of the 57th conference of the association for computational linguistics (ACL). Association for Computational Linguistics, pp 4791\u20134800. https:\/\/doi.org\/10.18653\/V1\/P19-1472","DOI":"10.18653\/V1\/P19-1472"},{"key":"541_CR38","doi-asserted-by":"publisher","unstructured":"Zhang Y, Zhang P, Huang M, et al (2024) QQQ: quality quattuor-bit quantization for large language models. https:\/\/doi.org\/10.48550\/ARXIV.2406.09904, arxiv:2406.09904","DOI":"10.48550\/ARXIV.2406.09904"},{"key":"541_CR39","unstructured":"Zhao Y, Lin C, Zhu K, et al (2024) Atom: Low-bit quantization for efficient and accurate LLM serving. In: Proceedings of the seventh annual conference on machine learning and systems, MLSys 2024, Santa Clara, CA, USA, May 13\u201316, 2024. https:\/\/proceedings.mlsys.org\/paper_files\/paper\/2024\/hash\/5edb57c05c81d04beb716ef1d542fe9e-Abstract-Conference.html"}],"container-title":["Journal of King Saud University Computer and Information Sciences"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s44443-026-00541-9","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s44443-026-00541-9.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s44443-026-00541-9.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,7,20]],"date-time":"2026-07-20T08:16:13Z","timestamp":1784535373000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s44443-026-00541-9"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,3,5]]},"references-count":39,"journal-issue":{"issue":"6","published-print":{"date-parts":[[2026,8]]}},"alternative-id":["541"],"URL":"https:\/\/doi.org\/10.1007\/s44443-026-00541-9","relation":{},"ISSN":["1319-1578","2213-1248"],"issn-type":[{"value":"1319-1578","type":"print"},{"value":"2213-1248","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,3,5]]},"assertion":[{"value":"24 October 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"30 January 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"5 March 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":1,"name":"Ethics","label":"Competing Interests","group":{"name":"EthicsHeading","label":"Declarations"}}],"article-number":"341"}}