{"url":"/method/e-branchformer","slug":"e-branchformer","name":"E-Branchformer","full_name":"E-Branchformer","full_name_withheld":false,"description_markdown":"E-BRANCHFORMER: BRANCHFORMER WITH ENHANCED MERGING FOR SPEECH RECOGNITION","description_state":"present","introduced_year":null,"introduced_by":{"title":"E-Branchformer: Branchformer with Enhanced merging for speech recognition","paper":"/paper/e-branchformer-branchformer-with-enhanced","first_author":"Kwangyoun Kim","n_authors":7,"url_abs":null,"archive_paper_url":"https://paperswithcode.com/paper/e-branchformer-branchformer-with-enhanced"},"source":{"url":"https://arxiv.org/abs/2210.00077v2","title":"E-Branchformer: Branchformer with Enhanced merging for speech recognition","url_on_a_paper_host":true},"code_snippet_url":null,"code_snippet_url_on_a_code_host":false,"categories":[{"area":"Natural Language Processing","area_id":"natural-language-processing","collection":"Transformers","url":"/methods/category/transformers","pwc_aliases":[]}],"n_papers_tagged":13,"archive_num_papers":13,"papers_newest_first":[{"paper":null,"title":"Whale: Large-Scale multilingual ASR model with w2v-BERT and E-Branchformer with large speech data","date":"2025-06-02","arxiv_id":"2506.01439","n_code_links":0,"syntology":null},{"paper":null,"title":"Late fusion ensembles for speech recognition on diverse input audio representations","date":"2024-12-01","arxiv_id":"2412.01861","n_code_links":0,"syntology":null},{"paper":null,"title":"Improving Automatic Speech Recognition with Decoder-Centric Regularisation in Encoder-Decoder Models","date":"2024-10-22","arxiv_id":"2410.17437","n_code_links":0,"syntology":null},{"paper":"/paper/2408-02369","title":"The NPU-ASLP System Description for Visual Speech Recognition in CNVSRC 2024","date":"2024-08-05","arxiv_id":"2408.02369","n_code_links":1,"syntology":null},{"paper":"/paper/multi-convformer-extending-conformer-with","title":"Multi-Convformer: Extending Conformer with Multiple Convolution Kernels","date":"2024-07-04","arxiv_id":"2407.03718","n_code_links":1,"syntology":null},{"paper":null,"title":"Low-resource speech recognition and dialect identification of Irish in a multi-task framework","date":"2024-05-02","arxiv_id":"2405.01293","n_code_links":0,"syntology":null},{"paper":null,"title":"Enhancing Lip Reading with Multi-Scale Video and Multi-Encoder","date":"2024-04-08","arxiv_id":"2404.05466","n_code_links":0,"syntology":null},{"paper":"/paper/owsm-v3-1-better-and-faster-open-whisper","title":"OWSM v3.1: Better and Faster Open Whisper-Style Speech Models based on E-Branchformer","date":"2024-01-30","arxiv_id":"2401.16658","n_code_links":1,"syntology":null},{"paper":"/paper/the-npu-aslp-liauto-system-description-for","title":"The NPU-ASLP-LiAuto System Description for Visual Speech Recognition in CNVSRC 2023","date":"2024-01-07","arxiv_id":"2401.06788","n_code_links":2,"syntology":null},{"paper":null,"title":"Tensor decomposition for minimization of E2E SLU model toward on-device processing","date":"2023-06-02","arxiv_id":"2306.01247","n_code_links":0,"syntology":null},{"paper":"/paper/a-comparative-study-on-e-branchformer-vs","title":"A Comparative Study on E-Branchformer vs Conformer in Speech Recognition, Translation, and Understanding Tasks","date":"2023-05-18","arxiv_id":"2305.11073","n_code_links":2,"syntology":null},{"paper":null,"title":"The NPU-ASLP System for Audio-Visual Speech Recognition in MISP 2022 Challenge","date":"2023-03-11","arxiv_id":"2303.06341","n_code_links":0,"syntology":null},{"paper":"/paper/e-branchformer-branchformer-with-enhanced","title":"E-Branchformer: Branchformer with Enhanced merging for speech recognition","date":"2022-09-30","arxiv_id":"2210.00077","n_code_links":1,"syntology":null}],"papers_shown":13,"tasks":[{"task":"/task/speech-recognition","name":"Speech Recognition","papers":12},{"task":"/task/speech-recognition-1","name":"speech-recognition","papers":12},{"task":"/task/automatic-speech-recognition-2","name":"Automatic Speech Recognition","papers":5},{"task":"/task/automatic-speech-recognition","name":"Automatic Speech Recognition (ASR)","papers":5},{"task":"/task/decoder","name":"Decoder","papers":5},{"task":"/task/visual-speech-recognition","name":"Visual Speech Recognition","papers":3},{"task":"/task/spoken-language-understanding","name":"Spoken Language Understanding","papers":2},{"task":"/task/audio-visual-speech-recognition","name":"Audio-Visual Speech Recognition","papers":1},{"task":"/task/dialect-identification","name":"Dialect Identification","papers":1},{"task":"/task/language-modeling","name":"Language Modeling","papers":1},{"task":"/task/language-modelling","name":"Language Modelling","papers":1},{"task":"/task/lip-reading","name":"Lip Reading","papers":1},{"task":"/task/lipreading","name":"Lipreading","papers":1},{"task":"/task/task-2","name":"Task 2","papers":1},{"task":"/task/tensor-decomposition","name":"Tensor Decomposition","papers":1}],"tasks_shown":15,"n_tasks":15,"usage_by_year":[{"year":"2022","papers":1},{"year":"2023","papers":3},{"year":"2024","papers":8},{"year":"2025","papers":1}],"row_source":"methods_table","archive":{"source":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","archive_url":"https://paperswithcode.com/method/e-branchformer"},"syntology_read_at":"2026-09-24T18:15:14+00:00"}