{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,4,11]],"date-time":"2026-04-11T03:32:34Z","timestamp":1775878354624,"version":"3.50.1"},"reference-count":57,"publisher":"Springer Science and Business Media LLC","issue":"22","license":[{"start":{"date-parts":[[2024,9,4]],"date-time":"2024-09-04T00:00:00Z","timestamp":1725408000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,9,4]],"date-time":"2024-09-04T00:00:00Z","timestamp":1725408000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimed Tools Appl"],"DOI":"10.1007\/s11042-024-20143-9","type":"journal-article","created":{"date-parts":[[2024,9,4]],"date-time":"2024-09-04T01:02:19Z","timestamp":1725411739000},"page":"25941-25958","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["Enhancing spoken dialect identification with stacked generalization of deep learning models"],"prefix":"10.1007","volume":"84","author":[{"ORCID":"https:\/\/orcid.org\/0000-0003-2649-4419","authenticated-orcid":false,"given":"Khaled","family":"Lounnas","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Mohamed","family":"Lichouri","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Mourad","family":"Abbas","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2024,9,4]]},"reference":[{"issue":"6","key":"20143_CR1","doi-asserted-by":"publisher","first-page":"691","DOI":"10.1017\/S1351324920000091","volume":"26","author":"A Hanani","year":"2020","unstructured":"Hanani A, Naser R (2020) Spoken Arabic dialect recognition using X-vectors. Nat Lang Eng 26(6):691\u2013700. https:\/\/doi.org\/10.1017\/S1351324920000091","journal-title":"Nat Lang Eng"},{"issue":"8","key":"20143_CR2","doi-asserted-by":"publisher","first-page":"23367","DOI":"10.1007\/s11042-023-16438-y","volume":"83","author":"AS Dhanjal","year":"2024","unstructured":"Dhanjal AS, Singh W (2024) A comprehensive survey on automatic speech recognition using neural networks. Multimed Tools Appli 83(8):23367\u201323412","journal-title":"Multimed Tools Appli"},{"issue":"12","key":"20143_CR3","doi-asserted-by":"publisher","first-page":"34499","DOI":"10.1007\/s11042-023-17094-y","volume":"83","author":"AA Alemu","year":"2024","unstructured":"Alemu AA, Melese MD, Salau AO (2024) Ethio-semitic language identification using convolutional neural networks with data augmentation. Multimed Tools Appl 83(12):34499\u201334514","journal-title":"Multimed Tools Appl"},{"key":"20143_CR4","doi-asserted-by":"crossref","unstructured":"Lonergan L, Qian M, Chiar\u00e1in NN, Gobl C, Chasaide AN (2023) Towards spoken dialect identification of Irish. arXiv:2307.07436","DOI":"10.21437\/SIGUL.2023-14"},{"key":"20143_CR5","doi-asserted-by":"publisher","first-page":"30205","DOI":"10.1007\/s11042-020-09321-7","volume":"79","author":"A Sharma","year":"2020","unstructured":"Sharma A, Kumar P, Maddukuri V, Madamshetti N, Kishore KG, Kavuru SSS, Roy PP (2020) Fast Griffin Lim based waveform generation strategy for text-to-speech synthesis. Multimed Tools Appli 79:30205\u201330233","journal-title":"Multimed Tools Appli"},{"key":"20143_CR6","doi-asserted-by":"crossref","unstructured":"Nazir O, Malik A, Singh S, Pathan ASK (2024) Multi speaker text-to-speech synthesis using generalized end-to-end loss function. Multimed Tools Appl 1-18","DOI":"10.1007\/s11042-024-18121-2"},{"issue":"2","key":"20143_CR7","doi-asserted-by":"publisher","first-page":"739","DOI":"10.11591\/ijai.v12.i2.pp739-746","volume":"12","author":"MA Humayun","year":"2023","unstructured":"Humayun MA, Yassin H, Abas PE (2023) Dialect classification using acoustic and linguistic features in Arabic speech. IAES Int J Artif Intell 12(2):739","journal-title":"IAES Int J Artif Intell"},{"key":"20143_CR8","doi-asserted-by":"crossref","unstructured":"Codru\u021b R, Ristea N, Ionescu R (2024, June) RoDia: a new dataset for Romanian dialect identification from speech. In: Findings of the association for computational linguistics: NAACL 2024 (pp 279-286)","DOI":"10.18653\/v1\/2024.findings-naacl.20"},{"key":"20143_CR9","doi-asserted-by":"crossref","unstructured":"Das HC, Bhattacharjee U (2024) Assamese dialect identification using static and dynamic features from vowel. J Adv Inf Technol 15(2)","DOI":"10.12720\/jait.15.2.306-321"},{"key":"20143_CR10","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2024.127664","volume":"587","author":"C Song","year":"2024","unstructured":"Song C, Ma Y, Xu Y, Chen H (2024) Multi-population evolutionary neural architecture search with stacked generalization. Neurocomputing 587:127664","journal-title":"Neurocomputing"},{"issue":"2","key":"20143_CR11","doi-asserted-by":"publisher","first-page":"248","DOI":"10.3390\/sym16020248","volume":"16","author":"S Aslam","year":"2024","unstructured":"Aslam S, Aslam H, Manzoor A, Chen H, Rasool A (2024) AntiPhishStack: LSTM-based stacked generalization model for optimized phishing URL detection. Symmetry 16(2):248","journal-title":"Symmetry"},{"key":"20143_CR12","doi-asserted-by":"publisher","first-page":"5131","DOI":"10.1007\/s11042-013-1587-5","volume":"74","author":"IJ Ding","year":"2015","unstructured":"Ding IJ, Yen CT (2015) Enhancing GMM speaker identification by incorporating SVM speaker verification for intelligent web-based speech applications. Multimed Tools Appl 74:5131\u20135140","journal-title":"Multimed Tools Appl"},{"key":"20143_CR13","doi-asserted-by":"crossref","unstructured":"Singh ST, Tiwari M (2024) A stacked generalization based meta-classifier for prediction of cloud workload. ICTACT J Soft Comput 14(4)","DOI":"10.21917\/ijsc.2024.0469"},{"issue":"7","key":"20143_CR14","doi-asserted-by":"publisher","first-page":"9565","DOI":"10.1007\/s11042-021-11439-1","volume":"82","author":"M Biswas","year":"2023","unstructured":"Biswas M, Rahaman S, Ahmadian A, Subari K, Singh PK (2023) Automatic spoken language identification using MFCC based time series features. Multimed Tools Appl 82(7):9565\u20139595","journal-title":"Multimed Tools Appl"},{"issue":"3","key":"20143_CR15","doi-asserted-by":"publisher","first-page":"3713","DOI":"10.1007\/s11042-022-13428-4","volume":"82","author":"D Khurana","year":"2023","unstructured":"Khurana D, Koli A, Khatter K, Singh S (2023) Natural language processing: state of the art, current trends and challenges. Multimed Tools Appl 82(3):3713\u20133744","journal-title":"Multimed Tools Appl"},{"key":"20143_CR16","doi-asserted-by":"publisher","first-page":"131","DOI":"10.1007\/s11042-011-0918-7","volume":"66","author":"MM Mustaquim","year":"2013","unstructured":"Mustaquim MM (2013) Automatic speech recognition-an approach for designing inclusive games. Multimed Tools Appl 66:131\u2013146","journal-title":"Multimed Tools Appl"},{"key":"20143_CR17","doi-asserted-by":"crossref","unstructured":"Xie Y (2019) A multimedia network independent learning aided translation system. Multimed Tools Appl 1-15","DOI":"10.1007\/s11042-019-7499-2"},{"key":"20143_CR18","doi-asserted-by":"publisher","first-page":"681","DOI":"10.1007\/s11042-012-1073-5","volume":"68","author":"T Athanaselis","year":"2014","unstructured":"Athanaselis T, Bakamidis S, Dologlou I, Argyriou EN, Symvonis A (2014) Making assistive reading tools user friendly: a new platform for Greek dyslexic students empowered by automatic speech recognition. Multimed Tools Appl 68:681\u2013699","journal-title":"Multimed Tools Appl"},{"key":"20143_CR19","doi-asserted-by":"publisher","first-page":"9","DOI":"10.1016\/j.procs.2018.03.002","volume":"128","author":"S Bougrine","year":"2018","unstructured":"Bougrine S, Cherroun H, Ziadi D (2018) Prosody-based spoken Algerian Arabic dialect identification. Procedia Comput Sci 128:9\u201317","journal-title":"Procedia Comput Sci"},{"key":"20143_CR20","unstructured":"Howard AG, Zhu M, Chen B, Kalenichenko D, Wang W, Weyand T, Adam H (2017) Mobilenets: efficient convolutional neural networks for mobile vision applications. arXiv:1704.04861"},{"key":"20143_CR21","doi-asserted-by":"crossref","unstructured":"Huang G, Liu Z, Van Der Maaten L, Weinberger KQ (2017) Densely connected convolutional networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp 4700\u20134708)","DOI":"10.1109\/CVPR.2017.243"},{"issue":"3","key":"20143_CR22","doi-asserted-by":"publisher","first-page":"593","DOI":"10.1007\/s10772-016-9351-7","volume":"19","author":"SS Agrawal","year":"2016","unstructured":"Agrawal SS, Jain A, Sinha S (2016) Analysis and modeling of acoustic information for automatic dialect classification. Int J Speech Technol 19(3):593\u2013609","journal-title":"Int J Speech Technol"},{"key":"20143_CR23","doi-asserted-by":"crossref","unstructured":"He K, Zhang X, Ren S, Sun J (2016) Deep residual learning for image recognition. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp 770\u2013778)","DOI":"10.1109\/CVPR.2016.90"},{"key":"20143_CR24","unstructured":"Simonyan K, Zisserman A (2015) Very deep convolutional networks for large-scale image recognition. In: ICLR"},{"issue":"4","key":"20143_CR25","doi-asserted-by":"publisher","first-page":"687","DOI":"10.1007\/s10772-016-9360-6","volume":"19","author":"M Hassine","year":"2016","unstructured":"Hassine M, Boussaid L, Messaoud H (2016) Maghrebian dialect recognition based on support vector machines and neural network classifiers. Int J Speech Technol 19(4):687\u2013695","journal-title":"Int J Speech Technol"},{"issue":"4","key":"20143_CR26","doi-asserted-by":"publisher","first-page":"1099","DOI":"10.1007\/s10772-019-09646-1","volume":"22","author":"NB Chittaragi","year":"2019","unstructured":"Chittaragi NB, Koolagudi SG (2019) Acoustic-phonetic feature based Kannada dialect identification from vowel sounds. Int J Speech Technol 22(4):1099\u20131113","journal-title":"Int J Speech Technol"},{"issue":"2","key":"20143_CR27","doi-asserted-by":"publisher","first-page":"251","DOI":"10.1007\/s10772-020-09678-y","volume":"23","author":"S Shivaprasad","year":"2020","unstructured":"Shivaprasad S, Sadanandam M (2020) Identification of regional dialects of Telugu language using text independent speech processing models. Int J Speech Technol 23(2):251\u2013258","journal-title":"Int J Speech Technol"},{"key":"20143_CR28","doi-asserted-by":"crossref","unstructured":"Bougrine S, Chorana A, Lakhdari A, Cherroun H (2017, April). Toward a web-based speech corpus for Algerian dialectal Arabic varieties. In: Proceedings of the Third Arabic Natural Language Processing Workshop (pp 138-146)","DOI":"10.18653\/v1\/W17-1317"},{"key":"20143_CR29","unstructured":"Lounnas K, Abbas M, Lichouri M (2019, September) Building a speech corpus based on Arabic podcasts for language and dialect identification. In: Proceedings of the 3rd International Conference on Natural Language and Speech Processing (pp 54-58)"},{"key":"20143_CR30","doi-asserted-by":"crossref","unstructured":"Biadsy F, Hirschberg JB (2009) Using prosody and phonotactics in Arabic dialect identification","DOI":"10.21437\/Interspeech.2009-77"},{"key":"20143_CR31","doi-asserted-by":"publisher","first-page":"45","DOI":"10.1016\/S1319-1578(08)80004-3","volume":"20","author":"M Alghamdi","year":"2008","unstructured":"Alghamdi M, Alhargan F, Alkanhal M, Alkhairy A, Eldesouki M, Alenazi A (2008) Saudi accented Arabic voice bank. J King Saud Univ-Comput Inf Sci 20:45\u201364","journal-title":"J King Saud Univ-Comput Inf Sci"},{"key":"20143_CR32","doi-asserted-by":"crossref","unstructured":"Lounnas K, Satori H, Hamidi M, Teffahi H, Abbas M, Lichouri M (2020, April) CLIASR: a combined automatic speech recognition and language identification system. In: 2020 1st International Conference on Innovative Research in Applied Science, Engineering and Technology (IRASET) (pp 1-5). IEEE","DOI":"10.1109\/IRASET48871.2020.9092020"},{"key":"20143_CR33","doi-asserted-by":"crossref","unstructured":"Lounnas K, Abbas M, Lichouri M, Hamidi M, Satori H, Teffahi H (2022) Enhancement of spoken digits recognition for under-resourced languages: case of Algerian and Moroccan dialects. Int J Speech Technol 1-13","DOI":"10.1007\/s10772-022-09971-y"},{"key":"20143_CR34","doi-asserted-by":"crossref","unstructured":"Barkat M, Ohala J, Pellegrino F (1999) Prosody as a distinctive feature for the discrimination of Arabic dialects. In: Sixth Eur Conf Speech Commun Technol","DOI":"10.21437\/Eurospeech.1999-102"},{"key":"20143_CR35","unstructured":"Komatsu M (2001, January) What constitutes acoustic evidence of prosody? The use of linear predictive coding residual signal in perceptual language identification. In: LACUS Forum (vol 28. pp 277\u2013287). Linguistic Association of Canada and the United States"},{"key":"20143_CR36","doi-asserted-by":"crossref","unstructured":"Sadanandam M (2021) HMM based language identification from speech utterances of popular indic languages using spectral and prosodic features HMM based language identification from speech utterances of popular indic languages using spectral and prosodic features","DOI":"10.18280\/ts.380232"},{"key":"20143_CR37","doi-asserted-by":"crossref","unstructured":"Biswas M, Rahaman S, Ahmadian A, Subari K, Singh PK (2022) Automatic spoken language identification using MFCC based time series features. Multimed Tools Appl 1-31","DOI":"10.1007\/s11042-021-11439-1"},{"key":"20143_CR38","doi-asserted-by":"crossref","unstructured":"Albadr MAA, Tiun S, Ayob M, Nazri MZA, AL-Dhief FT (2023) Grey wolf optimization-extreme learning machine for automatic spoken language identification. Multimed Tools Appl 82(18):27165\u201327191","DOI":"10.1007\/s11042-023-14473-3"},{"key":"20143_CR39","doi-asserted-by":"crossref","unstructured":"Godbole S, Jadhav V, Birajdar G (2020) Indian language identification using deep learning. In: ITM Web of Conferences (vol 32. p 01010). EDP Sciences","DOI":"10.1051\/itmconf\/20203201010"},{"key":"20143_CR40","unstructured":"Eldesouki M, Dalvi F, Sajjad H, Darwish K (2016, December) Qcri@ dsl 2016: Spoken arabic dialect identification using textual features. In: Proceedings of the Third Workshop on NLP for Similar Languages, Varieties and Dialects (VarDial3) (pp 221-226)"},{"issue":"19","key":"20143_CR41","doi-asserted-by":"publisher","first-page":"56883","DOI":"10.1007\/s11042-023-17782-9","volume":"83","author":"H Mukherjee","year":"2024","unstructured":"Mukherjee H, Dhar A, Obaidullah SM, Santosh KC, Phadikar S, Roy K, Pal U (2024) LIFA: language identification from audio with LPCC-G features. Multimed Tools Appl 83(19):56883\u201356907","journal-title":"Multimed Tools Appl"},{"key":"20143_CR42","doi-asserted-by":"crossref","unstructured":"Moftah M, Fakhr MW, El Ramly, S (2018, April) Arabic dialect identification based on motif discovery using GMM-UBM with different motif lengths. In: 2018 2nd International Conference on Natural Language and Speech Processing (ICNLSP) (pp 1-6). IEEE","DOI":"10.1109\/ICNLSP.2018.8374397"},{"key":"20143_CR43","doi-asserted-by":"crossref","unstructured":"Singh MK (2024) Multimedia application for forensic automatic speaker recognition from disguised voices using MFCC feature extraction and classification techniques. Multimed Tools Appl 1-19","DOI":"10.1007\/s11042-024-18602-4"},{"key":"20143_CR44","doi-asserted-by":"crossref","unstructured":"Sandler M, Howard A, Zhu M, Zhmoginov A, Chen LC (2018) Mobilenetv2: inverted residuals and linear bottlenecks. In: Proceedings of the IEEE conference on computer vision and pattern recognition (pp 4510-4520)","DOI":"10.1109\/CVPR.2018.00474"},{"key":"20143_CR45","doi-asserted-by":"crossref","unstructured":"Hanani A, Basha H, Sharaf Y, Taylor S (2015, October). Palestinian Arabic regional accent recognition. In: 2015 International Conference on Speech Technology and Human-Computer Dialogue (SpeD) (pp 1-6). IEEE","DOI":"10.1109\/SPED.2015.7343088"},{"key":"20143_CR46","unstructured":"Ziedan R, Micheal M, Alsammak A, Mursi M, Elmaghraby A (2016, September) A unified approach for arabic language dialect detection. In: Twenty ninth international conference on computers applications in industry and engineering (CAINE) (pp 165-170)"},{"key":"20143_CR47","doi-asserted-by":"crossref","unstructured":"Lounnas K, Demri L, Falek L, Teffahi H (2018, October) Automatic language identification for Berber and Arabic languages using prosodic features. In: 2018 International Conference on Electrical Sciences and Technologies in Maghreb (CISTEM) (pp 1-4). IEEE","DOI":"10.1109\/CISTEM.2018.8613414"},{"key":"20143_CR48","doi-asserted-by":"crossref","unstructured":"Lounnas K, Abbas M, Teffahi H, Lichouri M (2019, March) A language identification system based on voxforge speech corpus. In: International Conference on Advanced Machine Learning Technologies and Applications. Springer, Cham, (pp 529-534)","DOI":"10.1007\/978-3-030-14118-9_53"},{"key":"20143_CR49","doi-asserted-by":"crossref","unstructured":"Biadsy F, Hirschberg JB, Ellis DP (2011) Dialect and accent recognition using phonetic-segmentation supervectors","DOI":"10.21437\/Interspeech.2011-285"},{"key":"20143_CR50","doi-asserted-by":"crossref","unstructured":"Khurana S, Najafian M, Ali AM, Hanai TA, Belinkov Y, Glass JR (2017) QMDIS: QCRI-MIT advanced dialect identification system. INTERSPEECH","DOI":"10.21437\/Interspeech.2017-1391"},{"key":"20143_CR51","unstructured":"Bougrine S, Cherroun H, Ziadi D (2017) Hierarchical classification for spoken Arabic dialect identification using prosody: case of Algerian dialects. arXiv:1703.10065"},{"key":"20143_CR52","unstructured":"Michon E, Pham MQ, Crego JM, Senellart J (2018, August) Neural network architectures for Arabic dialect identification. In: Proceedings of the Fifth Workshop on NLP for Similar Languages, Varieties and Dialects (VarDial 2018) (pp 128-136)"},{"key":"20143_CR53","doi-asserted-by":"publisher","unstructured":"Bohra N, Bhatnagar V (2021) \"Language identification using stacked convolutional neural network (SCNN),\" 2021 11th International Conference on Cloud Computing, Data Science & Engineering (Confluence), pp 20-25. https:\/\/doi.org\/10.1109\/Confluence51648.2021.9377037","DOI":"10.1109\/Confluence51648.2021.9377037"},{"issue":"3","key":"20143_CR54","doi-asserted-by":"publisher","first-page":"235","DOI":"10.1007\/s10772-014-9223-y","volume":"17","author":"H Satori","year":"2014","unstructured":"Satori H, ElHaoussi F (2014) Investigation Amazigh speech recognition using CMU tools. Int J Speech Technol 17(3):235\u2013243","journal-title":"Int J Speech Technol"},{"key":"20143_CR55","first-page":"25","volume":"22","author":"H Satori","year":"2007","unstructured":"Satori H, Harti M, Chenfour N (2007) Arabic Speech Recognition System using CMU-Sphinx4. Corpus 22:25","journal-title":"Corpus"},{"key":"20143_CR56","doi-asserted-by":"publisher","first-page":"113160","DOI":"10.1016\/j.eswa.2019.113160","volume":"146","author":"S Agarwal","year":"2020","unstructured":"Agarwal S, Chowdary CR (2020) A-stacking and a-bagging: adaptive versions of ensemble learning algorithms for spoof fingerprint detection. Expert Syst Appl 146:113160","journal-title":"Expert Syst Appl"},{"key":"20143_CR57","doi-asserted-by":"publisher","DOI":"10.1016\/j.energy.2020.118874","volume":"214","author":"M Massaoudi","year":"2021","unstructured":"Massaoudi M, Refaat SS, Chihi I, Trabelsi M, Oueslati FS, Abu-Rub H (2021) A novel stacked generalization ensemble-based hybrid LGBM-XGB-MLP model for short-term load forecasting. Energy 214:118874","journal-title":"Energy"}],"container-title":["Multimedia Tools and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-024-20143-9.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11042-024-20143-9\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-024-20143-9.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,9,5]],"date-time":"2025-09-05T22:04:13Z","timestamp":1757109853000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11042-024-20143-9"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,9,4]]},"references-count":57,"journal-issue":{"issue":"22","published-online":{"date-parts":[[2025,7]]}},"alternative-id":["20143"],"URL":"https:\/\/doi.org\/10.1007\/s11042-024-20143-9","relation":{},"ISSN":["1573-7721"],"issn-type":[{"value":"1573-7721","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,9,4]]},"assertion":[{"value":"15 September 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"13 August 2024","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"18 August 2024","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"4 September 2024","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Competing Interests"}}]}}