{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,8,14]],"date-time":"2026-08-14T10:27:55Z","timestamp":1786703275423,"version":"build-2736575974"},"reference-count":43,"publisher":"Springer Science and Business Media LLC","issue":"4","license":[{"start":{"date-parts":[[2026,3,30]],"date-time":"2026-03-30T00:00:00Z","timestamp":1774828800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,3,30]],"date-time":"2026-03-30T00:00:00Z","timestamp":1774828800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["CCF Trans. HPC"],"published-print":{"date-parts":[[2026,8]]},"DOI":"10.1007\/s42514-026-00274-1","type":"journal-article","created":{"date-parts":[[2026,3,30]],"date-time":"2026-03-30T14:55:15Z","timestamp":1774882515000},"page":"393-409","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Directive-based automatic function-level vectorization for simplified SIMD exploitation"],"prefix":"10.1007","volume":"8","author":[{"given":"Lili","family":"Liu","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Bo","family":"Zhao","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yingying","family":"Li","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Ping","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jinlong","family":"Xu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jinyang","family":"Yao","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,3,30]]},"reference":[{"key":"274_CR1","doi-asserted-by":"crossref","unstructured":"Allen, J.R., Kennedy, K., Porterfield, C., Warren, J.: Conversion of control dependence to data dependence. In: Proceedings of the 10th ACM SIGACT-SIGPLAN Symposium on Principles of Programming Languages, pp. 177\u2013189 (1983)","DOI":"10.1145\/567067.567085"},{"key":"274_CR2","doi-asserted-by":"crossref","unstructured":"Barredo, A., Cebrian, J.M., Moret\u00f3, M., Casas, M., Valero, M.: Improving predication efficiency through compaction\/restoration of simd instructions. In: 2020 IEEE International Symposium on High Performance Computer Architecture (HPCA), pp. 717\u2013728 (2020). IEEE","DOI":"10.1109\/HPCA47549.2020.00064"},{"key":"274_CR3","doi-asserted-by":"publisher","first-page":"371","DOI":"10.1016\/j.future.2020.10.036","volume":"116","author":"H Bian","year":"2021","unstructured":"Bian, H., Huang, J., Liu, L., Huang, D., Wang, X.: Albus: a method for efficiently processing spmv using simd and load balancing. Futur. Gener. Comput. Syst. 116, 371\u2013392 (2021)","journal-title":"Futur. Gener. Comput. Syst."},{"key":"274_CR4","unstructured":"Board, O.A.R.: OpenMP Application Program Interface Version 5.0 Reference Guide. Accessed 2024-09-07. https:\/\/www.openmp.org\/wp-content\/uploads\/OpenMPRef-5.0-0519-web.pdf (2018)"},{"key":"274_CR5","doi-asserted-by":"crossref","unstructured":"B\u00f6hm, C., Plant, C.: Massively parallel graph drawing and representation learning. In: 2020 IEEE International Conference on Big Data (Big Data), pp. 609\u2013616 (2020). IEEE","DOI":"10.1109\/BigData50022.2020.9377976"},{"issue":"2","key":"274_CR6","first-page":"241","volume":"44","author":"H Chen","year":"2016","unstructured":"Chen, H., Yang, C., Liu, S., Liu, Z.: An efficient simd parallel memory structure for radix-2 fft computation. Acta Electron. Sin. 44(2), 241\u2013246 (2016)","journal-title":"Acta Electron. Sin."},{"key":"274_CR7","unstructured":"Community, G.: GCC 12.3.0. Accessed 2024-09-07. https:\/\/ftp.gnu.org\/gnu\/gcc\/gcc-12.3.0 (2023)"},{"issue":"3","key":"274_CR8","first-page":"180","volume":"43","author":"J Feng","year":"2022","unstructured":"Feng, J., He, Y., Tao, Q.: Auto-vectorization: recent development and prospect. J. Commun. 43(3), 180\u2013195 (2022)","journal-title":"J. Commun."},{"issue":"1","key":"274_CR9","doi-asserted-by":"publisher","first-page":"1832522","DOI":"10.1155\/2022\/1832522","volume":"2022","author":"J Feng","year":"2022","unstructured":"Feng, J., He, Y., Tao, Q., Ma, H.: An slp vectorization method based on equivalent extended transformation. Wirel. Commun. Mob. Comput. 2022(1), 1832522 (2022)","journal-title":"Wirel. Commun. Mob. Comput."},{"issue":"12","key":"274_CR10","first-page":"2907","volume":"60","author":"J Feng","year":"2023","unstructured":"Feng, J., He, Y., Tao, Q., Ma, H.: Slp vectorization method based on multiple isomorphic transformations. J. Comput. Res. Dev. 60(12), 2907\u201329252927 (2023)","journal-title":"J. Comput. Res. Dev."},{"issue":"15","key":"274_CR11","doi-asserted-by":"publisher","first-page":"2209","DOI":"10.1093\/bioinformatics\/btaa963","volume":"37","author":"Y Gao","year":"2021","unstructured":"Gao, Y., Liu, Y., Ma, Y., Liu, B., Wang, Y., Xing, Y.: abpoa: an simd-based c library for fast partial order alignment using adaptive band. Bioinformatics 37(15), 2209\u20132211 (2021)","journal-title":"Bioinformatics"},{"key":"274_CR12","doi-asserted-by":"crossref","unstructured":"Hampton, M., Asanovic, K.: Compiling for vector-thread architectures. In: Proceedings of the 6th Annual IEEE\/ACM International Symposium on Code Generation and Optimization, pp. 205\u2013215 (2008)","DOI":"10.1145\/1356058.1356085"},{"key":"274_CR13","unstructured":"Intel: Intel\u00ae64 and IA-32 Architectures Software Developer Manuals. (2022). Accessed: 2024-09-07. https:\/\/www.intel.cn\/content\/www\/cn\/zh\/developer\/articles\/technical\/intel-sdm.html"},{"key":"274_CR14","unstructured":"Intel: ispc: Intel SPMD Program Compiler. Accessed 2024-09-07. https:\/\/github.com\/ispc\/ispc\/releases\/tag\/v1.18.0 (2022)"},{"key":"274_CR15","doi-asserted-by":"crossref","unstructured":"Kandiah, V., Lustig, D., Villa, O., Nellans, D., Hardavellas, N.: Parsimony: enabling simd\/vector programming in standard compiler flows. In: Proceedings of the 21st ACM\/IEEE International Symposium on Code Generation and Optimization, pp. 186\u2013198 (2023)","DOI":"10.1145\/3579990.3580019"},{"key":"274_CR16","doi-asserted-by":"crossref","unstructured":"Karrenberg, R., Hack, S.: Whole-function vectorization. Automatic SIMD vectorization of SSA-based control flow graphs, 85\u2013125 (2015)","DOI":"10.1007\/978-3-658-10113-8_6"},{"issue":"13","key":"274_CR17","doi-asserted-by":"publisher","first-page":"2209","DOI":"10.14778\/3275366.3284966","volume":"11","author":"T Kersten","year":"2018","unstructured":"Kersten, T., Leis, V., Kemper, A., Neumann, T., Pavlo, A., Boncz, P.: Everything you always wanted to know about compiled and vectorized queries but were afraid to ask. Proc. VLDB Endow. 11(13), 2209\u20132222 (2018)","journal-title":"Proc. VLDB Endow."},{"key":"274_CR18","doi-asserted-by":"crossref","unstructured":"Kong, M., Veras, R., Stock, K., Franchetti, F., Pouchet, L.-N., Sadayappan, P.: When polyhedral transformations meet simd code generation. In: Proceedings of the 34th ACM SIGPLAN Conference on Programming Language Design and Implementation, pp. 127\u2013138 (2013)","DOI":"10.1145\/2491956.2462187"},{"issue":"5","key":"274_CR19","doi-asserted-by":"publisher","first-page":"145","DOI":"10.1145\/358438.349320","volume":"35","author":"S Larsen","year":"2000","unstructured":"Larsen, S., Amarasinghe, S.: Exploiting superword level parallelism with multimedia instruction sets. Acm Sigplan Not. 35(5), 145\u2013156 (2000)","journal-title":"Acm Sigplan Not."},{"key":"274_CR20","doi-asserted-by":"crossref","unstructured":"Li, H., Han, J., Han, D.: Leveraging simd parallelism for accelerating network applications. In: Proceedings of the 4th Asia-Pacific Workshop on Networking, pp. 23\u201329 (2020)","DOI":"10.1145\/3411029.3411033"},{"key":"274_CR21","doi-asserted-by":"crossref","unstructured":"Liu, B., Laird, A., Tsang, W.H., Mahjour, B., Dehnavi, M.M.: Combining run-time checks and compile-time analysis to improve control flow auto-vectorization. In: Proceedings of the International Conference on Parallel Architectures and Compilation Techniques, pp. 439\u2013450 (2022)","DOI":"10.1145\/3559009.3569663"},{"key":"274_CR22","doi-asserted-by":"crossref","unstructured":"Maleki, S., Gao, Y., Garzar, M.J., Wong, T., Padua, D.A., et al.: An evaluation of vectorizing compilers. In: 2011 International Conference on Parallel Architectures and Compilation Techniques, pp. 372\u2013382 (2011). IEEE","DOI":"10.1109\/PACT.2011.68"},{"key":"274_CR23","doi-asserted-by":"crossref","unstructured":"Masten, M., Tyurin, E., Mitropoulou, K., Garcia, E., Saito, H.: Function\/kernel vectorization via loop vectorizer. In: 2018 IEEE\/ACM 5th Workshop on the LLVM Compiler Infrastructure in HPC (LLVM-HPC), pp. 39\u201348 (2018). IEEE","DOI":"10.1109\/LLVM-HPC.2018.8639483"},{"issue":"OOPSLA","key":"274_CR24","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1145\/3276480","volume":"2","author":"C Mendis","year":"2018","unstructured":"Mendis, C., Amarasinghe, S.: goslp: globally optimized superword level parallelism framework. Proc. ACM Program. Lang. 2(OOPSLA), 1\u201328 (2018)","journal-title":"Proc. ACM Program. Lang."},{"issue":"4","key":"274_CR25","doi-asserted-by":"publisher","first-page":"543","DOI":"10.1145\/3296979.3192413","volume":"53","author":"S Moll","year":"2018","unstructured":"Moll, S., Hack, S.: Partial control-flow linearization. ACM SIGPLAN Not. 53(4), 543\u2013556 (2018)","journal-title":"ACM SIGPLAN Not."},{"key":"274_CR26","unstructured":"Paktinatkeleshteri, R.: Efficient auto-vectorization for control-flow dependent loops through data permutation (2023)"},{"key":"274_CR27","doi-asserted-by":"crossref","unstructured":"Papaphilippou, P., HJ, K.P., Luk, W.: Simodense: a risc-v softcore optimised for exploring custom simd instructions. In: 2021 31st International Conference on Field-Programmable Logic and Applications (FPL), pp. 391\u2013397 (2021). IEEE","DOI":"10.1109\/FPL53798.2021.00082"},{"key":"274_CR28","doi-asserted-by":"crossref","unstructured":"Porpodas, V., Jones, T.M.: Throttling automatic vectorization: when less is more. In: 2015 International Conference on Parallel Architecture and Compilation (PACT), pp. 432\u2013444 (2015). IEEE","DOI":"10.1109\/PACT.2015.32"},{"key":"274_CR29","doi-asserted-by":"crossref","unstructured":"Porpodas, V., Ratnalikar, P.: Postslp: cross-region vectorization of fully or partially vectorized code. In: Languages and Compilers for Parallel Computing: 32nd International Workshop, LCPC 2019, Atlanta, GA, USA, October 22\u201324, 2019, Revised Selected Papers 32, pp. 15\u201331 (2021). Springer","DOI":"10.1007\/978-3-030-72789-5_2"},{"key":"274_CR30","doi-asserted-by":"crossref","unstructured":"Porpodas, V., Rocha, R.C., G\u00f3es, L.F.: Vw-slp: auto-vectorization with adaptive vector width. In: Proceedings of the 27th International Conference on Parallel Architectures and Compilation Techniques, pp. 1\u201315 (2018)","DOI":"10.1145\/3243176.3243189"},{"key":"274_CR31","doi-asserted-by":"crossref","unstructured":"Rayhan, Y., Aref, W.G.: Simd-ified r-tree query processing and optimization. In: Proceedings of the 31st ACM International Conference on Advances in Geographic Information Systems, pp. 1\u201310 (2023)","DOI":"10.1145\/3589132.3625610"},{"key":"274_CR32","doi-asserted-by":"crossref","unstructured":"Reiche, O., Kobylko, C., Hannig, F., Teich, J.: Auto-vectorization for image processing dsls. In: Proceedings of the 18th ACM SIGPLAN\/SIGBED Conference on Languages, Compilers, and Tools for Embedded Systems, pp. 21\u201330 (2017)","DOI":"10.1145\/3078633.3081039"},{"key":"274_CR33","unstructured":"RISC: RISC-V \u201cV\u201d Vector Extension. (2021). Accessed: 2024-09-07. https:\/\/github.com\/riscv\/riscv-v-spec\/releases\/tag\/v1.0-rc2"},{"issue":"4","key":"274_CR34","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1145\/3356842","volume":"16","author":"S Siso","year":"2019","unstructured":"Siso, S., Armour, W., Thiyagalingam, J.: Evaluating auto-vectorizing compilers through objective withdrawal of useful information. ACM Trans. Archit. Code Optim. (TACO) 16(4), 1\u201323 (2019)","journal-title":"ACM Trans. Archit. Code Optim. (TACO)"},{"issue":"2","key":"274_CR35","doi-asserted-by":"publisher","first-page":"26","DOI":"10.1109\/MM.2017.35","volume":"37","author":"N Stephens","year":"2017","unstructured":"Stephens, N., Biles, S., Boettcher, M., Eapen, J., Eyole, M., Gabrielli, G., Horsnell, M., Magklis, G., Martinez, A., Premillieu, N., et al.: The arm scalable vector extension. IEEE Micro 37(2), 26\u201339 (2017)","journal-title":"IEEE Micro"},{"key":"274_CR36","doi-asserted-by":"crossref","unstructured":"Tian, X., Saito, H., Girkar, M., Preis, S.V., Kozhukhov, S.S., Cherkasov, A.G., Nelson, C., Panchenko, N., Geva, R.: Compiling c\/c++ simd extensions for function and loop vectorizaion on multicore-simd processors. In: 2012 IEEE 26th International Parallel and Distributed Processing Symposium Workshops & PhD Forum, pp. 2349\u20132358 (2012). IEEE","DOI":"10.1109\/IPDPSW.2012.292"},{"issue":"6","key":"274_CR37","first-page":"1265","volume":"26","author":"G Wei","year":"2015","unstructured":"Wei, G., Rongcai, Z., Lin, H.: Simd automatic vectorization summary of compiler optimization. J. Softw. 26(6), 1265\u20131284 (2015)","journal-title":"J. Softw."},{"key":"274_CR38","unstructured":"Yamazaki, S.: Future possibilities and effectiveness of jit from elixir code of image processing and machine learning into native code with simd instructions. In: IPSJ Special Interest Group on Programming 136th Meeting (2021)"},{"key":"274_CR39","unstructured":"Yermalayeu, I., et al.: The Simd Library. Accessed 2024-09-07. https:\/\/github.com\/ermig1979\/Simd (2019)"},{"key":"274_CR40","doi-asserted-by":"crossref","unstructured":"Zhang, Z., Ou, Y., Liu, Y., Wang, C., Zhou, Y., Wang, X., Zhang, Y., Ouyang, Y., Shan, J., Wang, Y., et al.: Occamy: Elastically sharing a simd co-processor across multiple cpu cores. In: Proceedings of the 28th ACM International Conference on Architectural Support for Programming Languages and Operating Systems, Vol. 3, pp. 483\u2013497 (2023)","DOI":"10.1145\/3582016.3582046"},{"key":"274_CR41","doi-asserted-by":"crossref","unstructured":"Zhao, J., Li, B., Nie, W., Geng, Z., Zhang, R., Gao, X., Cheng, B., Wu, C., Cheng, Y., Li, Z., et al.: Akg: automatic kernel generation for neural processing units using polyhedral transformations. In: Proceedings of the 42nd ACM SIGPLAN International Conference on Programming Language Design and Implementation, pp. 1233\u20131248 (2021)","DOI":"10.1145\/3453483.3454106"},{"key":"274_CR42","doi-asserted-by":"crossref","unstructured":"Zheng, R., Pai, S.: Efficient execution of graph algorithms on cpu with simd extensions. In: 2021 IEEE\/ACM International Symposium on Code Generation and Optimization (CGO), pp. 262\u2013276 (2021). IEEE","DOI":"10.1109\/CGO51591.2021.9370326"},{"key":"274_CR43","unstructured":"Zhou, C., Hassman, Z., Xu, R., Shah, D., Richard, V., Li, Y.: Simd dataflow co-optimization for efficient neural networks inferences on cpus. arXiv preprint arXiv:2310.00574 (2023)"}],"container-title":["CCF Transactions on High Performance Computing"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s42514-026-00274-1.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s42514-026-00274-1","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s42514-026-00274-1.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,8,14]],"date-time":"2026-08-14T09:33:12Z","timestamp":1786699992000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s42514-026-00274-1"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,3,30]]},"references-count":43,"journal-issue":{"issue":"4","published-print":{"date-parts":[[2026,8]]}},"alternative-id":["274"],"URL":"https:\/\/doi.org\/10.1007\/s42514-026-00274-1","relation":{},"ISSN":["2524-4922","2524-4930"],"issn-type":[{"value":"2524-4922","type":"print"},{"value":"2524-4930","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,3,30]]},"assertion":[{"value":"19 July 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"31 January 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"30 March 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":1,"name":"Ethics","label":"Competing interest","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"Not applicable.","order":2,"name":"Ethics","label":"Ethical approval and consent to participate","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"All authors consent to the publication of this manuscript.","order":3,"name":"Ethics","label":"Consent for publication","group":{"name":"EthicsHeading","label":"Declarations"}}]}}