{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/speech-recognition-1/papers/31","list_of":"/task/speech-recognition-1","task":"speech-recognition","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":31,"pages_in_order":58,"rows_per_page":100,"rows":[3001,3100],"of":5715,"counts":{"archive_papers_tagged":5715,"with_a_code_link":1277,"where_syntology_ran_a_sample":162,"not_listed_spam_title":0,"listed":5715,"listed_where_code_ran":162,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":134,"every_run_a_failure_of_syntologys_instrument":28,"listed_with_a_run_with_no_instrument_failure":134,"listed_every_run_a_failure_of_syntologys_instrument":28,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/speech-recognition-1","prev":"/task/speech-recognition-1/papers/30","next":"/task/speech-recognition-1/papers/32","papers":[{"url":null,"slug":"decoupled-federated-learning-for-asr-with-non","title":"Decoupled Federated Learning for ASR with Non-IID Data","date":"2022-06-18","arxiv_id":"2206.09102","repositories_listed":0,"syntology":null},{"url":null,"slug":"developing-a-speech-recognition-system-for","title":"Developing a Speech Recognition System for Recognizing Tonal Speech Signals Using a Convolutional Neural Network","date":"2022-06-17","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"a-ctc-triggered-siamese-network-with-spatial","title":"A CTC Triggered Siamese Network with Spatial-Temporal Dropout for Speech Recognition","date":"2022-06-16","arxiv_id":"2206.08031","repositories_listed":0,"syntology":null},{"url":null,"slug":"draft-a-novel-framework-to-reduce-domain","title":"DRAFT: A Novel Framework to Reduce Domain Shifting in Self-supervised Learning and Its Application to Children's ASR","date":"2022-06-16","arxiv_id":"2206.07931","repositories_listed":0,"syntology":null},{"url":null,"slug":"exploiting-cross-domain-and-cross-lingual","title":"Exploiting Cross-domain And Cross-Lingual Ultrasound Tongue Imaging Features For Elderly And Dysarthric Speech Recognition","date":"2022-06-15","arxiv_id":"2206.07327","repositories_listed":0,"syntology":null},{"url":null,"slug":"exploring-capabilities-of-monolingual-audio","title":"Exploring Capabilities of Monolingual Audio Transformers using Large Datasets in Automatic Speech Recognition of Czech","date":"2022-06-15","arxiv_id":"2206.07627","repositories_listed":0,"syntology":null},{"url":null,"slug":"residual-language-model-for-end-to-end-speech","title":"Residual Language Model for End-to-end Speech Recognition","date":"2022-06-15","arxiv_id":"2206.07430","repositories_listed":0,"syntology":null},{"url":null,"slug":"the-zevomos-entry-to-voicemos-challenge-2022","title":"The ZevoMOS entry to VoiceMOS Challenge 2022","date":"2022-06-15","arxiv_id":"2206.07448","repositories_listed":0,"syntology":null},{"url":null,"slug":"transformer-based-automatic-speech","title":"Transformer-based Automatic Speech Recognition of Formal and Colloquial Czech in MALACH Project","date":"2022-06-15","arxiv_id":"2206.07666","repositories_listed":0,"syntology":null},{"url":null,"slug":"toward-zero-oracle-word-error-rate-on-the","title":"Toward Zero Oracle Word Error Rate on the Switchboard Benchmark","date":"2022-06-13","arxiv_id":"2206.06192","repositories_listed":0,"syntology":null},{"url":null,"slug":"learning-based-data-storage-vision-technical","title":"Learning-Based Data Storage [Vision] (Technical Report)","date":"2022-06-12","arxiv_id":"2206.05778","repositories_listed":0,"syntology":null},{"url":null,"slug":"investigation-of-ensemble-features-of-self","title":"Investigation of Ensemble features of Self-Supervised Pretrained Models for Automatic Speech Recognition","date":"2022-06-11","arxiv_id":"2206.05518","repositories_listed":0,"syntology":null},{"url":null,"slug":"ahd-convnet-for-speech-emotion-classification","title":"AHD ConvNet for Speech Emotion Classification","date":"2022-06-10","arxiv_id":"2206.05286","repositories_listed":0,"syntology":null},{"url":null,"slug":"training-neural-networks-using-sat-solvers","title":"Training Neural Networks using SAT solvers","date":"2022-06-10","arxiv_id":"2206.04833","repositories_listed":0,"syntology":null},{"url":null,"slug":"context-based-out-of-vocabulary-word-recovery","title":"Context-based out-of-vocabulary word recovery for ASR systems in Indian languages","date":"2022-06-09","arxiv_id":"2206.04305","repositories_listed":0,"syntology":null},{"url":null,"slug":"face-dubbing-lip-synchronous-voice-preserving","title":"Face-Dubbing++: Lip-Synchronous, Voice Preserving Translation of Videos","date":"2022-06-09","arxiv_id":"2206.04523","repositories_listed":0,"syntology":null},{"url":null,"slug":"joint-encoder-decoder-self-supervised-pre","title":"Joint Encoder-Decoder Self-Supervised Pre-training for ASR","date":"2022-06-09","arxiv_id":"2206.04465","repositories_listed":0,"syntology":null},{"url":null,"slug":"legonn-building-modular-encoder-decoder","title":"LegoNN: Building Modular Encoder-Decoder Models","date":"2022-06-07","arxiv_id":"2206.03318","repositories_listed":0,"syntology":null},{"url":null,"slug":"fednst-federated-noisy-student-training-for","title":"FedNST: Federated Noisy Student Training for Automatic Speech Recognition","date":"2022-06-06","arxiv_id":"2206.02797","repositories_listed":0,"syntology":null},{"url":null,"slug":"lip-listening-mixing-senses-to-understand","title":"Lip-Listening: Mixing Senses to Understand Lips using Cross Modality Knowledge Distillation for Word-Based Models","date":"2022-06-05","arxiv_id":"2207.05692","repositories_listed":0,"syntology":null},{"url":null,"slug":"pronunciation-dictionary-free-multilingual","title":"Pronunciation Dictionary-Free Multilingual Speech Synthesis by Combining Unsupervised and Supervised Phonetic Representations","date":"2022-06-02","arxiv_id":"2206.00951","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-semi-automated-live-interlingual","title":"A Semi-Automated Live Interlingual Communication Workflow Featuring Intralingual Respeaking: Evaluation and Benchmarking","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"a-systematic-approach-to-derive-a-refined","title":"A Systematic Approach to Derive a Refined Speech Corpus for Sinhala","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"adversarial-speech-generation-and-natural","title":"Adversarial Speech Generation and Natural Speech Recovery for Speech Content Protection","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"automatic-speech-recognition-for-irish-the","title":"Automatic Speech Recognition for Irish: the ABAIR-ÉIST System","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"bea-base-a-benchmark-for-asr-of-spontaneous-1","title":"BEA-Base: A Benchmark for ASR of Spontaneous Hungarian","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"building-open-source-speech-technology-for","title":"Building Open-source Speech Technology for Low-resource Minority Languages with SáMi as an Example – Tools, Methods and Experiments","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"conversational-speech-recognition-needs-data","title":"Conversational Speech Recognition Needs Data? Experiments with Austrian German","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"developing-automatic-speech-recognition-for","title":"Developing Automatic Speech Recognition for Scottish Gaelic","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"development-and-evaluation-of-speech-1","title":"Development and Evaluation of Speech Recognition for the Welsh Language","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"development-of-automatic-speech-recognition","title":"Development of Automatic Speech Recognition for the Documentation of Cook Islands Māori","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"diabiz-an-annotated-corpus-of-polish-call","title":"DiaBiz – an Annotated Corpus of Polish Call Center Dialogs","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"evaluation-of-off-the-shelf-speech-1","title":"Evaluation of Off-the-shelf Speech Recognizers on Different Accents in a Dialogue Domain","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"extracting-linguistic-knowledge-from-speech-a","title":"Extracting Linguistic Knowledge from Speech: A Study of Stop Realization in 5 Romance Languages","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"generating-synthetic-clinical-speech-data","title":"Generating Synthetic Clinical Speech Data through Simulated ASR Deletion Error","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"handwriting-recognition-for-scottish-gaelic","title":"Handwriting recognition for Scottish Gaelic","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"huqariq-a-multilingual-speech-corpus-of-1","title":"Huqariq: A Multilingual Speech Corpus of Native Languages of Peru forSpeech Recognition","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"listra-automatic-speech-translation-english-1","title":"LiSTra Automatic Speech Translation: English to Lingala Case Study","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"mesures-linguistiques-automatiques-pour","title":"Mesures linguistiques automatiques pour l’évaluation des systèmes de Reconnaissance Automatique de la Parole (Automated linguistic measures for automatic speech recognition systems’ evaluation)","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"multilingual-transfer-learning-for-children","title":"Multilingual Transfer Learning for Children Automatic Speech Recognition","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"multiword-expressions-and-the-low-resource","title":"Multiword Expressions and the Low-Resource Scenario from the Perspective of a Local Oral Culture","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"parlamentparla-a-speech-corpus-of-catalan","title":"ParlamentParla: A Speech Corpus of Catalan Parliamentary Sessions","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"parlaspeech-hr-a-freely-available-asr-dataset","title":"ParlaSpeech-HR - a Freely Available ASR Dataset for Croatian Bootstrapped from the ParlaMint Corpus","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"post-stroke-speech-transcription-challenge","title":"Post-Stroke Speech Transcription Challenge (Task B): Correctness Detection in Anomia Diagnosis with Imperfect Transcripts","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"progress-in-multilingual-speech-recognition","title":"Progress in Multilingual Speech Recognition for Low Resource Languages Kurmanji Kurdish, Cree and Inuktut","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"rusavic-corpus-russian-audio-visual-speech-in","title":"RUSAVIC Corpus: Russian Audio-Visual Speech in Cars","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"samromur-children-an-icelandic-speech-corpus","title":"Samrómur Children: An Icelandic Speech Corpus","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"samromur-crowd-sourcing-large-amounts-of-data","title":"Samrómur: Crowd-sourcing large amounts of data","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"snow-mountain-dataset-of-audio-recordings-of","title":"Snow Mountain: Dataset of Audio Recordings of The Bible in Low Resource Languages","date":"2022-06-01","arxiv_id":"2206.01205","repositories_listed":0,"syntology":null},{"url":null,"slug":"standard-german-subtitling-of-swiss-german-tv","title":"Standard German Subtitling of Swiss German TV content: the PASSAGE Project","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-a-unified-asr-system-for-the-armenian","title":"Towards a Unified ASR System for the Armenian Standards","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-an-open-source-dutch-speech","title":"Towards an Open-Source Dutch Speech Recognition System for the Healthcare Domain","date":"2022-06-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"adversarial-synthesis-based-data-augmentation","title":"Adversarial synthesis based data-augmentation for code-switched spoken language identification","date":"2022-05-30","arxiv_id":"2205.15747","repositories_listed":0,"syntology":null},{"url":null,"slug":"speaker-identification-using-speech","title":"Speaker Identification using Speech Recognition","date":"2022-05-29","arxiv_id":"2205.14649","repositories_listed":0,"syntology":null},{"url":null,"slug":"adaptive-activation-network-for-low-resource","title":"Adaptive Activation Network For Low Resource Multilingual Speech Recognition","date":"2022-05-28","arxiv_id":"2205.14326","repositories_listed":0,"syntology":null},{"url":null,"slug":"is-lip-region-of-interest-sufficient-for","title":"Is Lip Region-of-Interest Sufficient for Lipreading?","date":"2022-05-28","arxiv_id":"2205.14295","repositories_listed":0,"syntology":null},{"url":null,"slug":"acoustic-to-articulatory-speech-inversion-1","title":"Acoustic-to-articulatory Speech Inversion with Multi-task Learning","date":"2022-05-27","arxiv_id":"2205.13755","repositories_listed":0,"syntology":null},{"url":null,"slug":"contrastive-siamese-network-for-semi","title":"Contrastive Siamese Network for Semi-supervised Speech Recognition","date":"2022-05-27","arxiv_id":"2205.14054","repositories_listed":0,"syntology":null},{"url":null,"slug":"punctuation-restoration-in-spanish-customer","title":"Punctuation Restoration in Spanish Customer Support Transcripts using Transfer Learning","date":"2022-05-27","arxiv_id":"2205.13961","repositories_listed":0,"syntology":null},{"url":null,"slug":"clinical-dialogue-transcription-error","title":"Clinical Dialogue Transcription Error Correction using Seq2Seq Models","date":"2022-05-26","arxiv_id":"2205.13572","repositories_listed":0,"syntology":null},{"url":null,"slug":"contextual-adapters-for-personalized-speech","title":"Contextual Adapters for Personalized Speech Recognition in Neural Transducers","date":"2022-05-26","arxiv_id":"2205.13660","repositories_listed":0,"syntology":null},{"url":null,"slug":"joint-training-of-speech-enhancement-and-self","title":"Joint Training of Speech Enhancement and Self-supervised Model for Noise-robust ASR","date":"2022-05-26","arxiv_id":"2205.13293","repositories_listed":0,"syntology":null},{"url":null,"slug":"an-investigation-on-applying-acoustic-feature","title":"An Investigation on Applying Acoustic Feature Conversion to ASR of Adult and Child Speech","date":"2022-05-25","arxiv_id":"2205.12477","repositories_listed":0,"syntology":null},{"url":null,"slug":"heterogeneous-reservoir-computing-models-for","title":"Heterogeneous Reservoir Computing Models for Persian Speech Recognition","date":"2022-05-25","arxiv_id":"2205.12594","repositories_listed":0,"syntology":null},{"url":null,"slug":"improving-ctc-based-asr-models-with-gated","title":"Improving CTC-based ASR Models with Gated Interlayer Collaboration","date":"2022-05-25","arxiv_id":"2205.12462","repositories_listed":0,"syntology":null},{"url":null,"slug":"investigating-lexical-replacements-for-arabic","title":"Investigating Lexical Replacements for Arabic-English Code-Switched Data Augmentation","date":"2022-05-25","arxiv_id":"2205.12649","repositories_listed":0,"syntology":null},{"url":null,"slug":"on-building-spoken-language-understanding","title":"On Building Spoken Language Understanding Systems for Low Resourced Languages","date":"2022-05-25","arxiv_id":"2205.12818","repositories_listed":0,"syntology":null},{"url":null,"slug":"semantic-preserved-communication-system-for","title":"Semantic-preserved Communication System for Highly Efficient Speech Transmission","date":"2022-05-25","arxiv_id":"2205.12727","repositories_listed":0,"syntology":null},{"url":null,"slug":"adaptive-multilingual-speech-recognition-with","title":"Adaptive multilingual speech recognition with pretrained models","date":"2022-05-24","arxiv_id":"2205.12304","repositories_listed":0,"syntology":null},{"url":null,"slug":"dpsnn-a-differentially-private-spiking-neural","title":"DPSNN: A Differentially Private Spiking Neural Network with Temporal Enhanced Pooling","date":"2022-05-24","arxiv_id":"2205.12718","repositories_listed":0,"syntology":null},{"url":null,"slug":"multi-level-modeling-units-for-end-to-end","title":"Multi-Level Modeling Units for End-to-End Mandarin Speech Recognition","date":"2022-05-24","arxiv_id":"2205.11998","repositories_listed":0,"syntology":null},{"url":null,"slug":"calibrate-and-refine-a-novel-and-agile","title":"Calibrate and Refine! A Novel and Agile Framework for ASR-error Robust Intent Detection","date":"2022-05-23","arxiv_id":"2205.11008","repositories_listed":0,"syntology":null},{"url":null,"slug":"deep-learning-for-visual-speech-analysis-a","title":"Deep Learning for Visual Speech Analysis: A Survey","date":"2022-05-22","arxiv_id":"2205.10839","repositories_listed":0,"syntology":null},{"url":null,"slug":"self-supervised-speech-representation","title":"Self-Supervised Speech Representation Learning: A Review","date":"2022-05-21","arxiv_id":"2205.10643","repositories_listed":0,"syntology":null},{"url":null,"slug":"neuralecho-a-self-attentive-recurrent-neural","title":"NeuralEcho: A Self-Attentive Recurrent Neural Network For Unified Acoustic Echo Suppression And Speech Enhancement","date":"2022-05-20","arxiv_id":"2205.10401","repositories_listed":0,"syntology":null},{"url":null,"slug":"predicting-electrode-array-impedance-after","title":"Predicting electrode array impedance after one month from cochlear implantation surgery","date":"2022-05-20","arxiv_id":"2205.10021","repositories_listed":0,"syntology":null},{"url":null,"slug":"set-based-meta-interpolation-for-few-task","title":"Set-based Meta-Interpolation for Few-Task Meta-Learning","date":"2022-05-20","arxiv_id":"2205.09990","repositories_listed":0,"syntology":null},{"url":null,"slug":"automatic-spoken-language-identification-1","title":"Automatic Spoken Language Identification using a Time-Delay Neural Network","date":"2022-05-19","arxiv_id":"2205.09564","repositories_listed":0,"syntology":null},{"url":null,"slug":"content-context-factorized-representations","title":"Content-Context Factorized Representations for Automated Speech Recognition","date":"2022-05-19","arxiv_id":"2205.09872","repositories_listed":0,"syntology":null},{"url":null,"slug":"insights-on-neural-representations-for-end-to","title":"Insights on Neural Representations for End-to-End Speech Recognition","date":"2022-05-19","arxiv_id":"2205.09456","repositories_listed":0,"syntology":null},{"url":null,"slug":"minimising-biasing-word-errors-for-contextual","title":"Minimising Biasing Word Errors for Contextual ASR with the Tree-Constrained Pointer Generator","date":"2022-05-18","arxiv_id":"2205.09058","repositories_listed":0,"syntology":null},{"url":null,"slug":"deploying-self-supervised-learning-in-the","title":"Deploying self-supervised learning in the wild for hybrid automatic speech recognition","date":"2022-05-17","arxiv_id":"2205.08598","repositories_listed":0,"syntology":null},{"url":null,"slug":"streaming-noise-context-aware-enhancement-for","title":"Streaming Noise Context Aware Enhancement For Automatic Speech Recognition in Multi-Talker Environments","date":"2022-05-17","arxiv_id":"2205.08555","repositories_listed":0,"syntology":null},{"url":null,"slug":"accented-speech-recognition-benchmarking-pre","title":"Accented Speech Recognition: Benchmarking, Pre-training, and Diverse Data","date":"2022-05-16","arxiv_id":"2205.08014","repositories_listed":0,"syntology":null},{"url":null,"slug":"improved-consistency-training-for-semi","title":"Improved Consistency Training for Semi-Supervised Sequence-to-Sequence ASR via Speech Chain Reconstruction and Self-Transcribing","date":"2022-05-14","arxiv_id":"2205.06963","repositories_listed":0,"syntology":null},{"url":null,"slug":"pretraining-approaches-for-spoken-language","title":"Pretraining Approaches for Spoken Language Recognition: TalTech Submission to the OLR 2021 Challenge","date":"2022-05-14","arxiv_id":"2205.07083","repositories_listed":0,"syntology":null},{"url":null,"slug":"personalized-adversarial-data-augmentation","title":"Personalized Adversarial Data Augmentation for Dysarthric and Elderly Speech Recognition","date":"2022-05-13","arxiv_id":"2205.06445","repositories_listed":0,"syntology":null},{"url":null,"slug":"unified-modeling-of-multi-domain-multi-device","title":"Unified Modeling of Multi-Domain Multi-Device ASR Systems","date":"2022-05-13","arxiv_id":"2205.06655","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-closer-look-at-audio-visual-multi-person","title":"A Closer Look at Audio-Visual Multi-Person Speech Recognition and Active Speaker Selection","date":"2022-05-11","arxiv_id":"2205.05684","repositories_listed":0,"syntology":null},{"url":null,"slug":"end-to-end-multi-person-audio-visual","title":"End-to-End Multi-Person Audio/Visual Automatic Speech Recognition","date":"2022-05-11","arxiv_id":"2205.05586","repositories_listed":0,"syntology":null},{"url":null,"slug":"improved-meta-learning-for-low-resource","title":"Improved Meta Learning for Low Resource Speech Recognition","date":"2022-05-11","arxiv_id":"2205.06182","repositories_listed":0,"syntology":null},{"url":null,"slug":"best-of-both-worlds-multi-task-audio-visual","title":"Best of Both Worlds: Multi-task Audio-Visual Automatic Speech Recognition and Active Speaker Detection","date":"2022-05-10","arxiv_id":"2205.05206","repositories_listed":0,"syntology":null},{"url":null,"slug":"separator-transducer-segmenter-streaming","title":"Separator-Transducer-Segmenter: Streaming Recognition and Segmentation of Multi-party Speech","date":"2022-05-10","arxiv_id":"2205.05199","repositories_listed":0,"syntology":null},{"url":null,"slug":"speaker-reinforcement-using-target-source","title":"Speaker Reinforcement Using Target Source Extraction for Robust Automatic Speech Recognition","date":"2022-05-09","arxiv_id":"2205.04433","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-conformer-based-waveform-domain-neural","title":"A Conformer-based Waveform-domain Neural Acoustic Echo Canceller Optimized for ASR Accuracy","date":"2022-05-06","arxiv_id":"2205.03481","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-highly-adaptive-acoustic-model-for-accurate","title":"A Highly Adaptive Acoustic Model for Accurate Multi-Dialect Speech Recognition","date":"2022-05-06","arxiv_id":"2205.03027","repositories_listed":0,"syntology":null},{"url":null,"slug":"hearing-voices-at-the-national-library-a","title":"Hearing voices at the National Library -- a speech corpus and acoustic model for the Swedish language","date":"2022-05-06","arxiv_id":"2205.03026","repositories_listed":0,"syntology":null},{"url":null,"slug":"online-model-compression-for-federated","title":"Online Model Compression for Federated Learning with Large Models","date":"2022-05-06","arxiv_id":"2205.03494","repositories_listed":0,"syntology":null},{"url":null,"slug":"design-of-a-novel-korean-learning-application","title":"Design of a novel Korean learning application for efficient pronunciation correction","date":"2022-05-04","arxiv_id":"2205.02001","repositories_listed":0,"syntology":null},{"url":null,"slug":"on-trac-consortium-systems-for-the-iwslt-2022","title":"ON-TRAC Consortium Systems for the IWSLT 2022 Dialect and Low-resource Speech Translation Tasks","date":"2022-05-04","arxiv_id":"2205.01987","repositories_listed":0,"syntology":null}],"record_sha256":"fa621061ad0167f4ac5fd06f9dc5c236e5a8ae9938a6a82eb381350112e0c8c0","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}