{"doi":"10.5220/0005023002550260","title":"HMM-based Breath and Filled Pauses Elimination in ASR","abstract":null,"journal":"Proceedings of the 11th International Conference on Signal Processing and Multimedia Applications","year":2014,"id":687905,"datarank":0.24141568686511508,"base_score":1.6094379124341003,"endowment":1.6094379124341003,"self_citation_contribution":0.24141568686511508,"citation_network_contribution":0.0,"self_endowment_contribution":0.24141568686511508,"citer_contribution":0.0,"corpus_percentile":null,"corpus_rank":null,"citation_count":4,"citer_count":0,"citers_with_citation_signal":0,"citers_with_endowment":0,"datacite_reuse_total":0,"is_dataset":false,"is_dataset_confidence":null,"is_data_producer":false,"deposit_databanks":null,"is_oa":false,"file_count":0,"downloads":0,"has_version_chain":false,"published_date":null,"fair_score":null,"fair_percentile":null,"algorithm_id":"datarank_citation_only_1hop_v6","ranking_scope":"data_only","authors":[{"id":1797106,"name":"Tomasz Jadczyk","orcid":null,"position":1,"is_corresponding":false},{"id":1797107,"name":"Bartosz Ziółko","orcid":null,"position":2,"is_corresponding":false},{"id":1797105,"name":"Piotr Żelasko","orcid":null,"position":0,"is_corresponding":false}],"reference_count":0,"raw_metadata":{"has_enrichment":true,"resolved":true,"title":"HMM-based Breath and Filled Pauses Elimination in ASR","abstract":"The phenomena of filled pauses and breaths pose a challenge to Automatic Speech Recognition (ASR) systems dealing with spontaneous speech, including recognizer modules in Interactive Voice Reponse (IVR) systems. We suggest a method based on Hidden Markov Models (HMM), which is easily integrated into HMM-based ASR systems and allows detection of those disturbances without incorporating additional parameters. Our method involves training the models of disturbances and their insertion in the phrase Markov chain between word-final and word-initial phoneme models. Application of the method in our ASR shows improvement of recognition results in Polish telephonic speech corpus LUNA.","is_dataset_classified":null,"base_score":1.6094379124341003,"endowment":1.6094379124341003,"datacite_reuse_total":0,"file_count":0,"downloads":0,"views":0,"has_version_chain":false,"is_dataset":false,"is_oa":false,"pmid":"21097893","pmcid":null,"openalex_id":"https://openalex.org/W1987783679","authors":[],"funders":[],"total_grants":0,"fwci":0.4136,"citation_percentile":0.57041762,"influential_citations":0,"citation_trend":[{"year":2017,"count":2},{"year":2020,"count":1},{"year":2022,"count":1}],"oa_status":"gold","license":"cc-by-nc-nd","oa_locations":[{"url":"https://doi.org/10.5220/0005023002550260","host_type":""},{"url":"https://doi.org/10.5220/0005023002550260","host_type":""}],"fields_of_study":["Speech Recognition and Synthesis","Speech and dialogue systems","Phonetics and Phonology Research"],"mesh_terms":[],"keywords":["Hidden Markov model","Speech recognition","Computer science","Phrase","Word (group theory)","Artificial intelligence","Natural language processing","Pattern recognition (psychology)","Linguistics"],"sdg_mappings":[{"sdg_number":0,"sdg_label":"Peace, Justice and strong institutions"}],"linked_datasets":[],"clinical_trials":[],"software_tools":[],"database_accessions":[],"source":"live","citation_network_status":"fetched"},"created_at":"2026-08-19T09:26:07.419521Z","pmid":null,"pmcid":null,"fwci":null,"citation_percentile":null,"influential_citations":0,"oa_status":null,"license":null,"views":0,"total_file_size_bytes":0,"version_count":0,"fair_f":null,"fair_a":null,"fair_i":null,"fair_r":null,"fair_zscore":null,"fair_rationale":null,"fair_model":null,"fair_agent_version":null,"fair_fulltext_source":null,"fair_has_llm":null,"fair_computed_at":null,"clinical_trials":[],"software_tools":[],"db_accessions":[],"linked_datasets":[],"topics":[]}