{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/automatic-speech-recognition/papers/23","list_of":"/task/automatic-speech-recognition","task":"Automatic Speech Recognition (ASR)","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":23,"pages_in_order":31,"rows_per_page":100,"rows":[2201,2300],"of":3012,"counts":{"archive_papers_tagged":3012,"with_a_code_link":622,"where_syntology_ran_a_sample":77,"not_listed_spam_title":0,"listed":3012,"listed_where_code_ran":77,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":64,"every_run_a_failure_of_syntologys_instrument":13,"listed_with_a_run_with_no_instrument_failure":64,"listed_every_run_a_failure_of_syntologys_instrument":13,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/automatic-speech-recognition","prev":"/task/automatic-speech-recognition/papers/22","next":"/task/automatic-speech-recognition/papers/24","papers":[{"url":null,"slug":"multimodal-speech-recognition-with","title":"Multimodal Speech Recognition with Unstructured Audio Masking","date":"2020-10-16","arxiv_id":"2010.08642","repositories_listed":0,"syntology":null},{"url":null,"slug":"non-intrusive-speech-intelligibility","title":"Non-intrusive speech intelligibility prediction using automatic speech recognition derived measures","date":"2020-10-16","arxiv_id":"2010.08574","repositories_listed":0,"syntology":null},{"url":null,"slug":"lightweight-end-to-end-speech-recognition","title":"Lightweight End-to-End Speech Recognition from Raw Audio Data Using Sinc-Convolutions","date":"2020-10-15","arxiv_id":"2010.07597","repositories_listed":0,"syntology":null},{"url":null,"slug":"exploiting-spectral-augmentation-for-code","title":"Exploiting Spectral Augmentation for Code-Switched Spoken Language Identification","date":"2020-10-14","arxiv_id":"2010.07130","repositories_listed":0,"syntology":null},{"url":null,"slug":"improving-low-resource-code-switched-asr","title":"Improving Low Resource Code-switched ASR using Augmented Code-switched TTS","date":"2020-10-12","arxiv_id":"2010.05549","repositories_listed":0,"syntology":null},{"url":null,"slug":"universal-asr-unify-and-improve-streaming-asr-1","title":"Dual-mode ASR: Unify and Improve Streaming ASR with Full-context Modeling","date":"2020-10-12","arxiv_id":"2010.06030","repositories_listed":0,"syntology":null},{"url":null,"slug":"wer-we-are-and-wer-we-think-we-are","title":"WER we are and WER we think we are","date":"2020-10-07","arxiv_id":"2010.03432","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-study-on-lip-localization-techniques-used","title":"A Study on Lip Localization Techniques used for Lip reading from a Video","date":"2020-09-28","arxiv_id":"2009.13420","repositories_listed":0,"syntology":null},{"url":null,"slug":"fluentnet-end-to-end-detection-of-speech","title":"FluentNet: End-to-End Detection of Speech Disfluency with Deep Learning","date":"2020-09-23","arxiv_id":"2009.11394","repositories_listed":0,"syntology":null},{"url":null,"slug":"easyasr-a-distributed-machine-learning","title":"EasyASR: A Distributed Machine Learning Platform for End-to-end Automatic Speech Recognition","date":"2020-09-14","arxiv_id":"2009.06487","repositories_listed":0,"syntology":null},{"url":null,"slug":"multi-modal-embeddings-using-multi-task","title":"Multi-modal embeddings using multi-task learning for emotion recognition","date":"2020-09-10","arxiv_id":"2009.05019","repositories_listed":0,"syntology":null},{"url":null,"slug":"unmanned-aerial-vehicle-control-through","title":"Unmanned Aerial Vehicle Control Through Domain-based Automatic Speech Recognition","date":"2020-09-09","arxiv_id":"2009.04215","repositories_listed":0,"syntology":null},{"url":null,"slug":"robust-spoken-language-understanding-with-rl","title":"Robust Spoken Language Understanding with RL-based Value Error Recovery","date":"2020-09-07","arxiv_id":"2009.03095","repositories_listed":0,"syntology":null},{"url":null,"slug":"silent-speech-interfaces-for-speech","title":"Silent Speech Interfaces for Speech Restoration: A Review","date":"2020-09-04","arxiv_id":"2009.02110","repositories_listed":0,"syntology":null},{"url":null,"slug":"voice-conversion-by-cascading-automatic","title":"Voice Conversion by Cascading Automatic Speech Recognition and Text-to-Speech Synthesis with Prosody Transfer","date":"2020-09-03","arxiv_id":"2009.01475","repositories_listed":0,"syntology":null},{"url":null,"slug":"convolutional-speech-recognition-with-pitch","title":"Convolutional Speech Recognition with Pitch and Voice Quality Features","date":"2020-09-02","arxiv_id":"2009.01309","repositories_listed":0,"syntology":null},{"url":null,"slug":"multi-view-attention-based-speech-enhancement","title":"Multi-view Attention-based Speech Enhancement Model for Noise-robust Automatic Speech Recognition","date":"2020-09-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"aphasic-speech-recognition-using-a-mixture-of","title":"Aphasic Speech Recognition using a Mixture of Speech Intelligibility Experts","date":"2020-08-25","arxiv_id":"2008.10788","repositories_listed":0,"syntology":null},{"url":null,"slug":"learned-transferable-architectures-can","title":"Learned Transferable Architectures Can Surpass Hand-Designed Architectures for Large Scale Speech Recognition","date":"2020-08-25","arxiv_id":"2008.11589","repositories_listed":0,"syntology":null},{"url":null,"slug":"improving-tail-performance-of-a-deliberation","title":"Improving Tail Performance of a Deliberation E2E ASR Model Using a Large Text Corpus","date":"2020-08-24","arxiv_id":"2008.10491","repositories_listed":0,"syntology":null},{"url":null,"slug":"cross-utterance-language-models-with-acoustic","title":"Cross-Utterance Language Models with Acoustic Error Sampling","date":"2020-08-19","arxiv_id":"2009.01008","repositories_listed":0,"syntology":null},{"url":null,"slug":"speech-to-semantics-improve-asr-and-nlu","title":"Speech To Semantics: Improve ASR and NLU Jointly via All-Neural Interfaces","date":"2020-08-14","arxiv_id":"2008.06173","repositories_listed":0,"syntology":null},{"url":null,"slug":"conv-transformer-transducer-low-latency-low","title":"Conv-Transformer Transducer: Low Latency, Low Frame Rate, Streamable End-to-End Speech Recognition","date":"2020-08-13","arxiv_id":"2008.05750","repositories_listed":0,"syntology":null},{"url":null,"slug":"large-scale-transfer-learning-for-low","title":"Large-scale Transfer Learning for Low-resource Spoken Language Understanding","date":"2020-08-13","arxiv_id":"2008.05671","repositories_listed":0,"syntology":null},{"url":"/paper/masri-headset-a-maltese-corpus-for-speech-1","slug":"masri-headset-a-maltese-corpus-for-speech-1","title":"MASRI-HEADSET: A Maltese Corpus for Speech Recognition","date":"2020-08-13","arxiv_id":"2008.05760","repositories_listed":0,"syntology":null},{"url":null,"slug":"online-automatic-speech-recognition-with","title":"Online Automatic Speech Recognition with Listen, Attend and Spell Model","date":"2020-08-12","arxiv_id":"2008.05514","repositories_listed":0,"syntology":null},{"url":null,"slug":"transfer-learning-approaches-for-streaming","title":"Transfer Learning Approaches for Streaming End-to-End Speech Recognition System","date":"2020-08-12","arxiv_id":"2008.05086","repositories_listed":0,"syntology":null},{"url":null,"slug":"transformer-with-bidirectional-decoder-for","title":"Transformer with Bidirectional Decoder for Speech Recognition","date":"2020-08-11","arxiv_id":"2008.04481","repositories_listed":0,"syntology":null},{"url":null,"slug":"subword-regularization-an-analysis-of","title":"Subword Regularization: An Analysis of Scalability and Generalization for End-to-End Automatic Speech Recognition","date":"2020-08-10","arxiv_id":"2008.04034","repositories_listed":0,"syntology":null},{"url":null,"slug":"lrspeech-extremely-low-resource-speech","title":"LRSpeech: Extremely Low-Resource Speech Synthesis and Recognition","date":"2020-08-09","arxiv_id":"2008.03687","repositories_listed":0,"syntology":null},{"url":null,"slug":"deep-learning-based-dereverberation-of","title":"Deep Learning Based Dereverberation of Temporal Envelopesfor Robust Speech Recognition","date":"2020-08-07","arxiv_id":"2008.03339","repositories_listed":0,"syntology":null},{"url":null,"slug":"investigation-of-speaker-adaptation-methods","title":"Investigation of Speaker-adaptation methods in Transformer based ASR","date":"2020-08-07","arxiv_id":"2008.03247","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-transfer-learning-method-for-speech-emotion","title":"A Transfer Learning Method for Speech Emotion Recognition from Automatic Speech Recognition","date":"2020-08-06","arxiv_id":"2008.02863","repositories_listed":0,"syntology":null},{"url":null,"slug":"iterative-compression-of-end-to-end-asr-model","title":"Iterative Compression of End-to-End ASR Model using AutoML","date":"2020-08-06","arxiv_id":"2008.02897","repositories_listed":0,"syntology":null},{"url":null,"slug":"shouted-speech-compensation-for-speaker","title":"Shouted Speech Compensation for Speaker Verification Robust to Vocal Effort Conditions","date":"2020-08-06","arxiv_id":"2008.02487","repositories_listed":0,"syntology":null},{"url":null,"slug":"unsupervised-cross-domain-singing-voice","title":"Unsupervised Cross-Domain Singing Voice Conversion","date":"2020-08-06","arxiv_id":"2008.02830","repositories_listed":0,"syntology":null},{"url":null,"slug":"this-is-houston-say-again-please-the-behavox","title":"\"This is Houston. Say again, please\". The Behavox system for the Apollo-11 Fearless Steps Challenge (phase II)","date":"2020-08-04","arxiv_id":"2008.01504","repositories_listed":0,"syntology":null},{"url":null,"slug":"weakly-supervised-construction-of-asr-systems","title":"Weakly Supervised Construction of ASR Systems with Massive Video Data","date":"2020-08-04","arxiv_id":"2008.01300","repositories_listed":0,"syntology":null},{"url":null,"slug":"modular-end-to-end-automatic-speech","title":"Modular End-to-end Automatic Speech Recognition Framework for Acoustic-to-word Model","date":"2020-07-31","arxiv_id":"2008.00953","repositories_listed":0,"syntology":null},{"url":null,"slug":"utterance-wise-meeting-transcription-system","title":"Utterance-Wise Meeting Transcription System Using Asynchronous Distributed Microphones","date":"2020-07-31","arxiv_id":"2007.15868","repositories_listed":0,"syntology":null},{"url":null,"slug":"developing-rnn-t-models-surpassing-high","title":"Developing RNN-T Models Surpassing High-Performance Hybrid Models with Customization Capability","date":"2020-07-30","arxiv_id":"2007.15188","repositories_listed":0,"syntology":null},{"url":null,"slug":"exploiting-cross-lingual-knowledge-in","title":"Exploiting Cross-Lingual Knowledge in Unsupervised Acoustic Modeling for Low-Resource Languages","date":"2020-07-29","arxiv_id":"2007.15074","repositories_listed":0,"syntology":null},{"url":null,"slug":"neural-kalman-filtering-for-speech","title":"Neural Kalman Filtering for Speech Enhancement","date":"2020-07-28","arxiv_id":"2007.13962","repositories_listed":0,"syntology":null},{"url":null,"slug":"effects-of-language-relatedness-for-cross-1","title":"Effects of Language Relatedness for Cross-lingual Transfer Learning in Character-Based Language Models","date":"2020-07-22","arxiv_id":"2007.11648","repositories_listed":0,"syntology":null},{"url":null,"slug":"audio-adversarial-examples-for-robust-hybrid","title":"Audio Adversarial Examples for Robust Hybrid CTC/Attention Speech Recognition","date":"2020-07-21","arxiv_id":"2007.10723","repositories_listed":0,"syntology":null},{"url":null,"slug":"neural-architecture-search-for-speech","title":"Neural Architecture Search For LF-MMI Trained Time Delay Neural Networks","date":"2020-07-17","arxiv_id":"2007.08818","repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-an-automated-soap-note-classifying","title":"Towards an Automated SOAP Note: Classifying Utterances from Medical Conversations","date":"2020-07-17","arxiv_id":"2007.08749","repositories_listed":0,"syntology":null},{"url":null,"slug":"the-faults-in-our-asrs-an-overview-of-attacks","title":"SoK: The Faults in our ASRs: An Overview of Attacks against Automatic Speech Recognition and Speaker Identification Systems","date":"2020-07-13","arxiv_id":"2007.06622","repositories_listed":0,"syntology":null},{"url":null,"slug":"massively-multilingual-asr-50-languages-1","title":"Massively Multilingual ASR: 50 Languages, 1 Model, 1 Billion Parameters","date":"2020-07-06","arxiv_id":"2007.03001","repositories_listed":0,"syntology":null},{"url":null,"slug":"deep-graph-random-process-for-relational","title":"Deep Graph Random Process for Relational-Thinking-Based Speech Recognition","date":"2020-07-04","arxiv_id":"2007.02126","repositories_listed":0,"syntology":null},{"url":null,"slug":"robust-prediction-of-punctuation-and-1","title":"Robust Prediction of Punctuation and Truecasing for Medical ASR","date":"2020-07-04","arxiv_id":"2007.02025","repositories_listed":0,"syntology":null},{"url":null,"slug":"cuni-neural-asr-with-phoneme-level","title":"CUNI Neural ASR with Phoneme-Level Intermediate Step for\\textasciitildeNon-Native\\textasciitildeSLT at IWSLT 2020","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":"/paper/end-to-end-offline-speech-translation-system","slug":"end-to-end-offline-speech-translation-system","title":"End-to-End Offline Speech Translation System for IWSLT 2020 using Modality Agnostic Meta-Learning","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"how-accents-confound-probing-for-accent","title":"How Accents Confound: Probing for Accent Information in End-to-End Speech Recognition Systems","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"investigating-the-effect-of-auxiliary","title":"Investigating the effect of auxiliary objectives for the automated grading of learner English speech transcriptions","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"large-vocabulary-read-speech-corpora-for-four-1","title":"Large Vocabulary Read Speech Corpora for Four Ethiopian Languages: Amharic, Tigrigna, Oromo, and Wolaytta","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-and-multiresolution-speech","title":"Multimodal and Multiresolution Speech Recognition with Transformers","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"robust-neural-machine-translation-with-asr","title":"Robust Neural Machine Translation with ASR Errors","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"simulspeech-end-to-end-simultaneous-speech-to","title":"SimulSpeech: End-to-End Simultaneous Speech to Text Translation","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"start-before-end-and-end-to-end-neural-speech","title":"Start-Before-End and End-to-End: Neural Speech Translation by AppTek and RWTH Aachen University","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"the-afrl-iwslt-2020-systems-work-from-home","title":"The AFRL IWSLT 2020 Systems: Work-From-Home Edition","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"tigrinya-automatic-speech-recognition-with","title":"Tigrinya Automatic Speech recognition with Morpheme based recognition units","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-end-2-end-learning-for-predicting","title":"Towards end-2-end learning for predicting behavior codes from spoken utterances in psychotherapy conversations","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-stream-translation-adaptive","title":"Towards Stream Translation: Adaptive Computation Time for Simultaneous Machine Translation","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-understanding-asr-error-correction","title":"Towards Understanding ASR Error Correction for Medical Conversations","date":"2020-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"multi-view-frequency-lstm-an-efficient","title":"Multi-view Frequency LSTM: An Efficient Frontend for Automatic Speech Recognition","date":"2020-06-30","arxiv_id":"2007.00131","repositories_listed":0,"syntology":null},{"url":null,"slug":"neural-machine-translation-for-multilingual","title":"Neural Machine Translation for Multilingual Grapheme-to-Phoneme Conversion","date":"2020-06-25","arxiv_id":"2006.14194","repositories_listed":0,"syntology":null},{"url":null,"slug":"streaming-transformer-asr-with-blockwise","title":"Streaming Transformer ASR with Blockwise Synchronous Inference","date":"2020-06-25","arxiv_id":"2006.14941","repositories_listed":0,"syntology":null},{"url":null,"slug":"efficient-active-learning-for-automatic","title":"Boosting Active Learning for Speech Recognition with Noisy Pseudo-labeled Samples","date":"2020-06-19","arxiv_id":"2006.11021","repositories_listed":0,"syntology":null},{"url":null,"slug":"joint-speaker-counting-speech-recognition-and","title":"Joint Speaker Counting, Speech Recognition, and Speaker Identification for Overlapped Speech of Any Number of Speakers","date":"2020-06-19","arxiv_id":"2006.10930","repositories_listed":0,"syntology":null},{"url":null,"slug":"end-to-end-code-switching-language-models-for","title":"End-to-End Code Switching Language Models for Automatic Speech Recognition","date":"2020-06-16","arxiv_id":"2006.08870","repositories_listed":0,"syntology":null},{"url":null,"slug":"quantization-of-acoustic-model-parameters-in","title":"Quantization of Acoustic Model Parameters in Automatic Speech Recognition Framework","date":"2020-06-16","arxiv_id":"2006.09054","repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-automated-assessment-of-stuttering","title":"Towards Automated Assessment of Stuttering and Stuttering Therapy","date":"2020-06-16","arxiv_id":"2006.09222","repositories_listed":0,"syntology":null},{"url":null,"slug":"exploration-of-end-to-end-asr-for-openstt","title":"Exploration of End-to-End ASR for OpenSTT -- Russian Open Speech-to-Text Dataset","date":"2020-06-15","arxiv_id":"2006.08274","repositories_listed":0,"syntology":null},{"url":null,"slug":"the-jhu-multi-microphone-multi-speaker-asr","title":"The JHU Multi-Microphone Multi-Speaker ASR System for the CHiME-6 Challenge","date":"2020-06-14","arxiv_id":"2006.07898","repositories_listed":0,"syntology":null},{"url":null,"slug":"data-augmentation-for-training-dialog-models","title":"Data Augmentation for Training Dialog Models Robust to Speech Recognition Errors","date":"2020-06-10","arxiv_id":"2006.05635","repositories_listed":0,"syntology":null},{"url":null,"slug":"improving-cross-lingual-transfer-learning-for","title":"Improving Cross-Lingual Transfer Learning for End-to-End Speech Recognition with Speech Translation","date":"2020-06-09","arxiv_id":"2006.05474","repositories_listed":0,"syntology":null},{"url":null,"slug":"learning-not-to-discriminate-task-agnostic","title":"Learning not to Discriminate: Task Agnostic Learning for Improving Monolingual and Code-switched Speech Recognition","date":"2020-06-09","arxiv_id":"2006.05257","repositories_listed":0,"syntology":null},{"url":null,"slug":"on-the-effectiveness-of-neural-text","title":"On the Effectiveness of Neural Text Generation based Data Augmentation for Recognition of Morphologically Rich Speech","date":"2020-06-09","arxiv_id":"2006.05129","repositories_listed":0,"syntology":null},{"url":null,"slug":"contextual-rnn-t-for-open-domain-asr","title":"Contextual RNN-T For Open Domain ASR","date":"2020-06-04","arxiv_id":"2006.03411","repositories_listed":0,"syntology":null},{"url":null,"slug":"multi-talker-asr-for-an-unknown-number-of","title":"Multi-talker ASR for an unknown number of sources: Joint training of source counting, separation and ASR","date":"2020-06-04","arxiv_id":"2006.02786","repositories_listed":0,"syntology":null},{"url":null,"slug":"transfer-learning-for-british-sign-language-1","title":"Transfer Learning for British Sign Language Modelling","date":"2020-06-03","arxiv_id":"2006.02144","repositories_listed":0,"syntology":null},{"url":null,"slug":"analyzing-the-quality-and-stability-of-a","title":"Analyzing the Quality and Stability of a Streaming End-to-End On-Device Speech Recognizer","date":"2020-06-02","arxiv_id":"2006.01416","repositories_listed":0,"syntology":null},{"url":null,"slug":"detecting-audio-attacks-on-asr-systems-with","title":"Detecting Audio Attacks on ASR Systems with Dropout Uncertainty","date":"2020-06-02","arxiv_id":"2006.01906","repositories_listed":0,"syntology":null},{"url":null,"slug":"dilated-u-net-based-approach-for-multichannel","title":"Dilated U-net based approach for multichannel speech enhancement from First-Order Ambisonics recordings","date":"2020-06-02","arxiv_id":"2006.01708","repositories_listed":0,"syntology":null},{"url":null,"slug":"an-effective-contextual-language-modeling","title":"An Effective Contextual Language Modeling Framework for Speech Summarization with Augmented Features","date":"2020-06-01","arxiv_id":"2006.01189","repositories_listed":0,"syntology":null},{"url":null,"slug":"analyse-de-l-effet-de-la-r-everb-eration-sur","title":"Analyse de l'effet de la r\\'everb\\'eration sur la reconnaissance automatique de la parole (Analyzing how reverberation affects Automatic Speech Recognition)","date":"2020-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"constrained-variational-autoencoder-for","title":"Constrained Variational Autoencoder for improving EEG based Speech Recognition Systems","date":"2020-06-01","arxiv_id":"2006.02902","repositories_listed":0,"syntology":null},{"url":null,"slug":"introduction-d-informations-s-emantiques-dans","title":"Introduction d'informations s\\'emantiques dans un syst\\`eme de reconnaissance de la parole (Despite spectacular advances in recent years, the Automatic Speech Recognition (ASR) systems still make mistakes, especially in noisy environments)","date":"2020-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"learning-to-recognize-code-switched-speech","title":"Learning to Recognize Code-switched Speech Without Forgetting Monolingual Speech Recognition","date":"2020-06-01","arxiv_id":"2006.00782","repositories_listed":0,"syntology":null},{"url":null,"slug":"reconnaissance-automatique-de-la-parole-g-en","title":"Reconnaissance automatique de la parole : g\\'en\\'eration des prononciations non natives pour l'enrichissement du lexique (In this study we propose a method for lexicon adaptation in order to improve the automatic speech recognition (ASR) of non-native speakers)","date":"2020-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"sur-l-utilisation-de-la-reconnaissance","title":"Sur l'utilisation de la reconnaissance automatique de la parole pour l'aide au diagnostic diff\\'erentiel entre la maladie de Parkinson et l'AMS (On using automatic speech recognition for the differential diagnosis of Parkinson's Disease and MSA This article presents a study regarding the contribution of automatic speech processing in the differential diagnosis between Parkinson's disease and MSA (Multi-System Atrophies))","date":"2020-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"dynamic-masking-for-improved-stability-in","title":"Dynamic Masking for Improved Stability in Spoken Language Translation","date":"2020-05-30","arxiv_id":"2006.00249","repositories_listed":0,"syntology":null},{"url":null,"slug":"an-audio-enriched-bert-based-framework-for","title":"An Audio-enriched BERT-based Framework for Spoken Multiple-choice Question Answering","date":"2020-05-25","arxiv_id":"2005.12142","repositories_listed":0,"syntology":null},{"url":null,"slug":"an-end-to-end-mispronunciation-detection","title":"An End-to-End Mispronunciation Detection System for L2 English Speech Leveraging Novel Anti-Phone Modeling","date":"2020-05-25","arxiv_id":"2005.11950","repositories_listed":0,"syntology":null},{"url":"/paper/ft-speech-danish-parliament-speech-corpus","slug":"ft-speech-danish-parliament-speech-corpus","title":"FT Speech: Danish Parliament Speech Corpus","date":"2020-05-25","arxiv_id":"2005.12368","repositories_listed":0,"syntology":null},{"url":null,"slug":"large-scale-evaluation-of-importance-maps-in","title":"Large scale evaluation of importance maps in automatic speech recognition","date":"2020-05-21","arxiv_id":"2005.10929","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-comparison-of-label-synchronous-and-frame","title":"A Comparison of Label-Synchronous and Frame-Synchronous End-to-End Models for Speech Recognition","date":"2020-05-20","arxiv_id":"2005.10113","repositories_listed":0,"syntology":null},{"url":null,"slug":"early-stage-lm-integration-using-local-and","title":"Early Stage LM Integration Using Local and Global Log-Linear Combination","date":"2020-05-20","arxiv_id":"2005.10049","repositories_listed":0,"syntology":null},{"url":null,"slug":"investigation-of-large-margin-softmax-in","title":"Investigation of Large-Margin Softmax in Neural Language Modeling","date":"2020-05-20","arxiv_id":"2005.10089","repositories_listed":0,"syntology":null}],"record_sha256":"c791f3f3cb1fe1573492faebc8aaa86ea418b2654aa83feb419027dfda00a760","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}