{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/speech-recognition/papers/53","list_of":"/task/speech-recognition","task":"Speech Recognition","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":53,"pages_in_order":65,"rows_per_page":100,"rows":[5201,5300],"of":6433,"counts":{"archive_papers_tagged":6433,"with_a_code_link":1373,"where_syntology_ran_a_sample":196,"not_listed_spam_title":0,"listed":6433,"listed_where_code_ran":196,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":162,"every_run_a_failure_of_syntologys_instrument":34,"listed_with_a_run_with_no_instrument_failure":162,"listed_every_run_a_failure_of_syntologys_instrument":34,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/speech-recognition","prev":"/task/speech-recognition/papers/52","next":"/task/speech-recognition/papers/54","papers":[{"url":null,"slug":"automatic-context-window-composition-for","title":"Automatic context window composition for distant speech recognition","date":"2018-05-26","arxiv_id":"1805.10498","repositories_listed":0,"syntology":null},{"url":null,"slug":"geometric-understanding-of-deep-learning","title":"Geometric Understanding of Deep Learning","date":"2018-05-26","arxiv_id":"1805.10451","repositories_listed":0,"syntology":null},{"url":null,"slug":"task-dependent-modulation-of-the-visual","title":"Task-dependent modulation of the visual sensory thalamus assists visual-speech recognition","date":"2018-05-24","arxiv_id":"1805.05682","repositories_listed":0,"syntology":null},{"url":null,"slug":"asr-based-features-for-emotion-recognition-a","title":"ASR-based Features for Emotion Recognition: A Transfer Learning Approach","date":"2018-05-23","arxiv_id":"1805.09197","repositories_listed":0,"syntology":null},{"url":null,"slug":"unsupervised-cross-modal-alignment-of-speech","title":"Unsupervised Cross-Modal Alignment of Speech and Text Embedding Spaces","date":"2018-05-18","arxiv_id":"1805.07467","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-comparison-of-modeling-units-in-sequence-to","title":"A Comparison of Modeling Units in Sequence-to-Sequence Speech Recognition with the Transformer on Mandarin Chinese","date":"2018-05-16","arxiv_id":"1805.06239","repositories_listed":0,"syntology":null},{"url":null,"slug":"composing-finite-state-transducers-on-gpus","title":"Composing Finite State Transducers on GPUs","date":"2018-05-16","arxiv_id":"1805.06383","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-purely-end-to-end-system-for-multi-speaker","title":"A Purely End-to-end System for Multi-speaker Speech Recognition","date":"2018-05-15","arxiv_id":"1805.05826","repositories_listed":0,"syntology":null},{"url":null,"slug":"improved-asr-for-under-resourced-languages","title":"Improved ASR for Under-Resourced Languages Through Multi-Task Learning with Acoustic Landmarks","date":"2018-05-15","arxiv_id":"1805.05574","repositories_listed":0,"syntology":null},{"url":null,"slug":"understanding-and-controlling-user","title":"Gradient-Leaks: Understanding and Controlling Deanonymization in Federated Learning","date":"2018-05-15","arxiv_id":"1805.05838","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-comparable-study-of-modeling-units-for-end","title":"A comparable study of modeling units for end-to-end Mandarin speech recognition","date":"2018-05-10","arxiv_id":"1805.03832","repositories_listed":0,"syntology":null},{"url":null,"slug":"transfer-learning-from-adult-to-children-for","title":"Transfer Learning from Adult to Children for Speech Recognition: Evaluation, Analysis and Recommendations","date":"2018-05-08","arxiv_id":"1805.03322","repositories_listed":0,"syntology":null},{"url":null,"slug":"boosting-noise-robustness-of-acoustic-model","title":"Boosting Noise Robustness of Acoustic Model via Deep Adversarial Training","date":"2018-05-02","arxiv_id":"1805.01357","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-comparative-study-of-extremely-low-resource","title":"A Comparative Study of Extremely Low-Resource Transliteration of the World's Languages","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"a-leveled-reading-corpus-of-modern-standard","title":"A Leveled Reading Corpus of Modern Standard Arabic","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"a-real-life-french-accented-corpus-of-air","title":"A Real-life, French-accented Corpus of Air Traffic Control Communications","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"a-semi-autonomous-system-for-creating-a-human","title":"A Semi-autonomous System for Creating a Human-Machine Interaction Corpus in Virtual Reality: Application to the ACORFORMed System for Training Doctors to Break Bad News","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"a-vietnamese-dialog-act-corpus-based-on-iso","title":"A Vietnamese Dialog Act Corpus Based on ISO 24617-2 standard","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"a-web-service-for-pre-segmenting-very-long","title":"A Web Service for Pre-segmenting Very Long Transcribed Speech Recordings","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"an-application-for-building-a-polish","title":"An Application for Building a Polish Telephone Speech Corpus","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"asr-for-documenting-acutely-under-resourced","title":"ASR for Documenting Acutely Under-Resourced Indigenous Languages","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"building-open-javanese-and-sundanese-corpora","title":"Building Open Javanese and Sundanese Corpora for Multilingual Text-to-Speech","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"classification-of-closely-related-sub","title":"Classification of Closely Related Sub-dialects of Arabic Using Support-Vector Machines","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"collecting-code-switched-data-from-social","title":"Collecting Code-Switched Data from Social Media","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"collection-and-analysis-of-code-switch","title":"Collection and Analysis of Code-switch Egyptian Arabic-English Speech Corpus","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"construction-of-english-french-multimodal","title":"Construction of English-French Multimodal Affective Conversational Corpus from TV Dramas","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"contextual-dependencies-in-time-continuous","title":"Contextual Dependencies in Time-Continuous Multidimensional Affect Recognition","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"cpjd-corpus-crowdsourced-parallel-speech","title":"CPJD Corpus: Crowdsourced Parallel Speech Corpus of Japanese Dialects","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"creating-lithuanian-and-latvian-speech","title":"Creating Lithuanian and Latvian Speech Corpora from Inaccurately Annotated Web Data","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"dart-a-large-dataset-of-dialectal-arabic","title":"DART: A Large Dataset of Dialectal Arabic Tweets","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"data-driven-pronunciation-modeling-of-swiss","title":"Data-Driven Pronunciation Modeling of Swiss German Dialectal Speech for Automatic Speech Recognition","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"design-and-development-of-speech-corpora-for","title":"Design and Development of Speech Corpora for Air Traffic Control Training","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"discovering-canonical-indian-english-accents","title":"Discovering Canonical Indian English Accents: A Crowdsourcing-based Approach","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"evaluation-of-feature-space-speaker","title":"Evaluation of Feature-Space Speaker Adaptation for End-to-End Acoustic Models","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"farmi-a-framework-for-recording-multi-modal","title":"FARMI: A FrAmework for Recording Multi-Modal Interactions","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"from-asolved-problemsa-to-new-challenges-a","title":"From `Solved Problems' to New Challenges: A Report on LDC Activities","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"improved-transcription-and-indexing-of-oral","title":"Improved Transcription and Indexing of Oral History Interviews for Digital Humanities Research","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"increasing-the-accessibility-of-time-aligned","title":"Increasing the Accessibility of Time-Aligned Speech Corpora with Spokes Mix","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"matics-software-suite-new-tools-for","title":"Matics Software Suite: New Tools for Evaluation and Data Exploration","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"mirasvoice-a-bilingual-english-persian-speech","title":"MirasVoice: A bilingual (English-Persian) speech corpus","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"mocca-measure-of-confidence-for-corpus","title":"MOCCA: Measure of Confidence for Corpus Analysis - Automatic Reliability Check of Transcript and Automatic Segmentation","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"multilingual-parallel-corpus-for-global","title":"Multilingual Parallel Corpus for Global Communication Plan","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"neural-caption-generation-for-news-images","title":"Neural Caption Generation for News Images","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"open-asr-for-icelandic-resources-and-a","title":"Open ASR for Icelandic: Resources and a Baseline System","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"parallel-corpora-in-mboshi-bantu-c25-congo","title":"Parallel Corpora in Mboshi (Bantu C25, Congo-Brazzaville)","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"phonetically-balanced-code-mixed-speech","title":"Phonetically Balanced Code-Mixed Speech Corpus for Hindi-English Automatic Speech Recognition","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"pronunciation-variants-and-asr-of-colloquial","title":"Pronunciation Variants and ASR of Colloquial Speech: A Case Study on Czech","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"simulating-asr-errors-for-training-slu","title":"Simulating ASR errors for training SLU systems","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"speech-rate-calculations-with-short","title":"Speech Rate Calculations with Short Utterances: A Study from a Speech-to-Speech, Machine Translation Mediated Map Task","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"text-normalization-infrastructure-that-scales","title":"Text Normalization Infrastructure that Scales to Hundreds of Language Varieties","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"the-waw-corpus-the-first-corpus-of","title":"The WAW Corpus: The First Corpus of Interpreted Speeches and their Translations for English and Arabic","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-an-automatic-assessment-of","title":"Towards an Automatic Assessment of Crowdsourced Data for NLU","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-processing-of-the-oral-history","title":"Towards Processing of the Oral History Interviews and Related Printed Documents","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"unified-guidelines-and-resources-for-arabic","title":"Unified Guidelines and Resources for Arabic Dialect Orthography","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"using-discourse-information-for-education","title":"Using Discourse Information for Education with a Spanish-Chinese Parallel Corpus","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"vast-a-corpus-of-video-annotation-for-speech","title":"VAST: A Corpus of Video Annotation for Speech Technologies","date":"2018-05-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"automatic-documentation-of-icd-codes-with-far","title":"Automatic Documentation of ICD Codes with Far-Field Speech Recognition","date":"2018-04-30","arxiv_id":"1804.11046","repositories_listed":0,"syntology":null},{"url":null,"slug":"investigations-on-end-to-end-audiovisual","title":"Investigations on End-to-End Audiovisual Fusion","date":"2018-04-30","arxiv_id":"1804.11127","repositories_listed":0,"syntology":null},{"url":null,"slug":"sparse-persistent-rnns-squeezing-large","title":"Sparse Persistent RNNs: Squeezing Large Recurrent Networks On-Chip","date":"2018-04-26","arxiv_id":"1804.10223","repositories_listed":0,"syntology":null},{"url":null,"slug":"end-to-end-multimodal-speech-recognition","title":"End-to-End Multimodal Speech Recognition","date":"2018-04-25","arxiv_id":"1804.09713","repositories_listed":0,"syntology":null},{"url":null,"slug":"recent-progresses-in-deep-learning-based","title":"Recent Progresses in Deep Learning based Acoustic Models (Updated)","date":"2018-04-25","arxiv_id":"1804.09298","repositories_listed":0,"syntology":null},{"url":null,"slug":"an-information-theoretic-view-for-deep","title":"An Information-Theoretic View for Deep Learning","date":"2018-04-24","arxiv_id":"1804.09060","repositories_listed":0,"syntology":null},{"url":null,"slug":"automatic-speech-recognition-for-launch","title":"Automatic speech recognition for launch control center communication using recurrent neural networks with data augmentation and custom language model","date":"2018-04-24","arxiv_id":"1804.09552","repositories_listed":0,"syntology":null},{"url":null,"slug":"multi-head-decoder-for-end-to-end-speech","title":"Multi-Head Decoder for End-to-End Speech Recognition","date":"2018-04-22","arxiv_id":"1804.08050","repositories_listed":0,"syntology":null},{"url":null,"slug":"neural-compatibility-modeling-with-attentive","title":"Neural Compatibility Modeling with Attentive Knowledge Distillation","date":"2018-04-17","arxiv_id":"1805.00313","repositories_listed":0,"syntology":null},{"url":null,"slug":"precise-detection-of-speech-endpoints","title":"Precise Detection of Speech Endpoints Dynamically: A Wavelet Convolution based approach","date":"2018-04-17","arxiv_id":"1804.06159","repositories_listed":0,"syntology":null},{"url":"/paper/neural-network-language-modeling-with-letter","slug":"neural-network-language-modeling-with-letter","title":"Neural Network Language Modeling with Letter-based Features and Importance Sampling","date":"2018-04-15","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"language-recognition-using-time-delay-deep","title":"Language Recognition using Time Delay Deep Neural Network","date":"2018-04-13","arxiv_id":"1804.05000","repositories_listed":0,"syntology":null},{"url":null,"slug":"global-snr-estimation-of-speech-signals-using","title":"Global SNR Estimation of Speech Signals using Entropy and Uncertainty Estimates from Dropout Networks","date":"2018-04-12","arxiv_id":"1804.04353","repositories_listed":0,"syntology":null},{"url":null,"slug":"vision-as-an-interlingua-learning","title":"Vision as an Interlingua: Learning Multilingual Semantic Embeddings of Untranscribed Speech","date":"2018-04-09","arxiv_id":"1804.03052","repositories_listed":0,"syntology":null},{"url":null,"slug":"joint-learning-of-interactive-spoken-content","title":"Joint Learning of Interactive Spoken Content Retrieval and Trainable User Simulator","date":"2018-04-01","arxiv_id":"1804.00318","repositories_listed":0,"syntology":null},{"url":null,"slug":"espnet-end-to-end-speech-processing-toolkit","title":"ESPnet: End-to-End Speech Processing Toolkit","date":"2018-03-30","arxiv_id":"1804.00015","repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-unsupervised-automatic-speech","title":"Towards Unsupervised Automatic Speech Recognition Trained by Unaligned Speech and Text only","date":"2018-03-29","arxiv_id":"1803.10952","repositories_listed":0,"syntology":null},{"url":null,"slug":"machine-speech-chain-with-one-shot-speaker","title":"Machine Speech Chain with One-shot Speaker Adaptation","date":"2018-03-28","arxiv_id":"1803.10525","repositories_listed":0,"syntology":null},{"url":"/paper/the-fifth-chime-speech-separation-and","slug":"the-fifth-chime-speech-separation-and","title":"The fifth 'CHiME' Speech Separation and Recognition Challenge: Dataset, task and baselines","date":"2018-03-28","arxiv_id":"1803.10609","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-multi-discriminator-cyclegan-for","title":"A Multi-Discriminator CycleGAN for Unsupervised Non-Parallel Speech Domain Adaptation","date":"2018-03-27","arxiv_id":"1804.00522","repositories_listed":0,"syntology":null},{"url":"/paper/building-state-of-the-art-distant-speech","slug":"building-state-of-the-art-distant-speech","title":"Building state-of-the-art distant speech recognition using the CHiME-4 challenge with a setup of speech enhancement baseline","date":"2018-03-27","arxiv_id":"1803.10109","repositories_listed":0,"syntology":null},{"url":null,"slug":"comprehending-real-numbers-development-of","title":"Comprehending Real Numbers: Development of Bengali Real Number Speech Corpus","date":"2018-03-27","arxiv_id":"1803.10136","repositories_listed":0,"syntology":null},{"url":null,"slug":"multi-modal-data-augmentation-for-end-to-end","title":"Multi-Modal Data Augmentation for End-to-End ASR","date":"2018-03-27","arxiv_id":"1803.10299","repositories_listed":0,"syntology":null},{"url":null,"slug":"student-teacher-learning-for-blstm-mask-based","title":"Student-Teacher Learning for BLSTM Mask-based Speech Enhancement","date":"2018-03-27","arxiv_id":"1803.10013","repositories_listed":0,"syntology":null},{"url":null,"slug":"clipping-free-attacks-against-artificial","title":"Clipping free attacks against artificial neural networks","date":"2018-03-26","arxiv_id":"1803.09468","repositories_listed":0,"syntology":null},{"url":null,"slug":"spectral-feature-mapping-with-mimic-loss-for","title":"Spectral feature mapping with mimic loss for robust speech recognition","date":"2018-03-26","arxiv_id":"1803.09816","repositories_listed":0,"syntology":null},{"url":null,"slug":"low-resource-speech-to-text-translation","title":"Low-Resource Speech-to-Text Translation","date":"2018-03-24","arxiv_id":"1803.09164","repositories_listed":0,"syntology":null},{"url":null,"slug":"exploring-the-robustness-of-features-and","title":"Exploring the robustness of features and enhancement on speech recognition systems in highly-reverberant real environments","date":"2018-03-23","arxiv_id":"1803.09013","repositories_listed":0,"syntology":null},{"url":null,"slug":"highly-reverberant-real-environment-database","title":"Highly-Reverberant Real Environment database: HRRE","date":"2018-03-23","arxiv_id":"1801.09651","repositories_listed":0,"syntology":null},{"url":null,"slug":"acoustic-feature-learning-using-cross-domain","title":"Acoustic feature learning using cross-domain articulatory measurements","date":"2018-03-19","arxiv_id":"1803.06805","repositories_listed":0,"syntology":null},{"url":null,"slug":"speech-to-text-and-text-to-speech-recognition","title":"Speech to text and text to speech recognition systems-Areview","date":"2018-03-17","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"tbd-benchmarking-and-analyzing-deep-neural","title":"TBD: Benchmarking and Analyzing Deep Neural Network Training","date":"2018-03-16","arxiv_id":"1803.06905","repositories_listed":0,"syntology":null},{"url":null,"slug":"advancing-connectionist-temporal","title":"Advancing Connectionist Temporal Classification With Attention Modeling","date":"2018-03-15","arxiv_id":"1803.05563","repositories_listed":0,"syntology":null},{"url":null,"slug":"resource-aware-design-of-a-deep-convolutional","title":"Resource aware design of a deep convolutional-recurrent neural network for speech recognition through audio-visual sensor fusion","date":"2018-03-13","arxiv_id":"1803.04840","repositories_listed":0,"syntology":null},{"url":null,"slug":"from-nodes-to-networks-evolving-recurrent","title":"From Nodes to Networks: Evolving Recurrent Neural Networks","date":"2018-03-12","arxiv_id":"1803.04439","repositories_listed":0,"syntology":null},{"url":null,"slug":"speech-recognition-keyword-spotting-through","title":"Speech Recognition: Keyword Spotting Through Image Recognition","date":"2018-03-10","arxiv_id":"1803.03759","repositories_listed":0,"syntology":null},{"url":null,"slug":"extracting-domain-invariant-features-by","title":"Extracting Domain Invariant Features by Unsupervised Learning for Robust Automatic Speech Recognition","date":"2018-03-07","arxiv_id":"1803.02551","repositories_listed":0,"syntology":null},{"url":null,"slug":"on-modular-training-of-neural-acoustics-to","title":"On Modular Training of Neural Acoustics-to-Word Model for LVCSR","date":"2018-03-03","arxiv_id":"1803.01090","repositories_listed":0,"syntology":null},{"url":null,"slug":"the-history-began-from-alexnet-a","title":"The History Began from AlexNet: A Comprehensive Survey on Deep Learning Approaches","date":"2018-03-03","arxiv_id":"1803.01164","repositories_listed":0,"syntology":null},{"url":null,"slug":"challenges-in-speech-recognition-and","title":"Challenges in Speech Recognition and Translation of High-Value Low-Density Polysynthetic Languages","date":"2018-03-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"evaluating-automatic-speech-recognition-in","title":"Evaluating Automatic Speech Recognition in Translation","date":"2018-03-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"neural-network-methods-for-natural-language","title":"Neural Network Methods for Natural Language Processing by Yoav Goldberg","date":"2018-03-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"on-the-derivational-entropy-of-left-to-right","title":"On the Derivational Entropy of Left-to-Right Probabilistic Finite-State Automata and Hidden Markov Models","date":"2018-03-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"portable-speech-to-speech-translation-on-an","title":"Portable Speech-to-Speech Translation on an Android Smartphone: The MFLTS System","date":"2018-03-01","arxiv_id":null,"repositories_listed":0,"syntology":null}],"record_sha256":"a98d7c00927f3fedaf586b3cc3c0a9211f35e5e01d92f0d956f588519dd88e78","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}