{"url":"/task/automatic-phoneme-recognition","name":"Automatic Phoneme Recognition","slug":"automatic-phoneme-recognition","description_markdown":"**Automatic Phoneme Recognition** (APR) involves converting spoken language into a sequence of phonemes, which are the distinct units of sound that distinguish one word from another in a given language. It is designed to transcribe spoken words into their textual phonetic representations in real-time, enabling detailed analysis of speech patterns, pronunciation, and linguistic nuances. The goal of Automatic Phoneme Recognition is to accurately identify and transcribe phonemes, considering variations in accent, pronunciation, and speaking style, as well as background noise and other factors that can affect speech quality. This technology is crucial for linguistic research, speech therapy, language learning applications, and enhancing the performance of speech recognition systems.","categories":[{"name":"Speech","url":"/area/speech"}],"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","slug_source":"archive_url"},"counts":{"papers_tagged":5,"papers_with_code":1,"benchmarks":6,"benchmark_tables_in_archive":6,"benchmark_tables_shown":6,"benchmark_tables_withheld_as_spam":0,"benchmark_definition":"a leaderboard table with at least one row; benchmark_tables_shown also counts the zero-row tables; benchmark_tables_in_archive adds the tables withheld as spam","datasets":6,"subtasks":0,"parent_tasks":1},"benchmarks":[{"leaderboard":"/sota/automatic-phoneme-recognition-on-vibravox","slug":"automatic-phoneme-recognition-on-vibravox","dataset":"VibraVox (headset microphone)","dataset_url":"/dataset/vibravox-headset-microphone","rows_in_archive":1,"metrics":["Test PER"],"first_row_in_archive_order":{"model":"medium wav2vec2.0","paper_title":"Vibravox: A Dataset of French Speech Captured with Body-conduction Audio Sensors","paper_url":"/paper/vibravox-a-dataset-of-french-speech-captured","paper_date":"2024-07-16","arxiv_id":"2407.11828","code_links":[{"title":"jhauret/vibravox","url":"https://github.com/jhauret/vibravox"}],"syntology":null}},{"leaderboard":"/sota/automatic-phoneme-recognition-on-vibravox-1","slug":"automatic-phoneme-recognition-on-vibravox-1","dataset":"VibraVox (forehead accelerometer)","dataset_url":"/dataset/vibravox-forehead-accelerometer","rows_in_archive":1,"metrics":["Test PER"],"first_row_in_archive_order":{"model":"medium wav2vec2.0","paper_title":"Vibravox: A Dataset of French Speech Captured with Body-conduction Audio Sensors","paper_url":"/paper/vibravox-a-dataset-of-french-speech-captured","paper_date":"2024-07-16","arxiv_id":"2407.11828","code_links":[{"title":"jhauret/vibravox","url":"https://github.com/jhauret/vibravox"}],"syntology":null}},{"leaderboard":"/sota/automatic-phoneme-recognition-on-vibravox-2","slug":"automatic-phoneme-recognition-on-vibravox-2","dataset":"VibraVox (soft in-ear microphone)","dataset_url":"/dataset/vibravox-soft-in-ear-microphone","rows_in_archive":1,"metrics":["Test PER"],"first_row_in_archive_order":{"model":"medium wav2vec2.0","paper_title":"Vibravox: A Dataset of French Speech Captured with Body-conduction Audio Sensors","paper_url":"/paper/vibravox-a-dataset-of-french-speech-captured","paper_date":"2024-07-16","arxiv_id":"2407.11828","code_links":[{"title":"jhauret/vibravox","url":"https://github.com/jhauret/vibravox"}],"syntology":null}},{"leaderboard":"/sota/automatic-phoneme-recognition-on-vibravox-3","slug":"automatic-phoneme-recognition-on-vibravox-3","dataset":"VibraVox (rigid in-ear microphone)","dataset_url":"/dataset/vibravox","rows_in_archive":1,"metrics":["Test PER"],"first_row_in_archive_order":{"model":"medium wav2vec2.0","paper_title":"Vibravox: A Dataset of French Speech Captured with Body-conduction Audio Sensors","paper_url":"/paper/vibravox-a-dataset-of-french-speech-captured","paper_date":"2024-07-16","arxiv_id":"2407.11828","code_links":[{"title":"jhauret/vibravox","url":"https://github.com/jhauret/vibravox"}],"syntology":null}},{"leaderboard":"/sota/automatic-phoneme-recognition-on-vibravox-4","slug":"automatic-phoneme-recognition-on-vibravox-4","dataset":"VibraVox (throat microphone)","dataset_url":"/dataset/vibravox-throat-microphone","rows_in_archive":1,"metrics":["Test PER"],"first_row_in_archive_order":{"model":"medium wav2vec2.0","paper_title":"Vibravox: A Dataset of French Speech Captured with Body-conduction Audio Sensors","paper_url":"/paper/vibravox-a-dataset-of-french-speech-captured","paper_date":"2024-07-16","arxiv_id":"2407.11828","code_links":[{"title":"jhauret/vibravox","url":"https://github.com/jhauret/vibravox"}],"syntology":null}},{"leaderboard":"/sota/automatic-phoneme-recognition-on-vibravox-5","slug":"automatic-phoneme-recognition-on-vibravox-5","dataset":"VibraVox (temple vibration pickup)","dataset_url":"/dataset/vibravox-temple-vibration-pickup","rows_in_archive":1,"metrics":["Test PER"],"first_row_in_archive_order":{"model":"medium wav2vec2.0","paper_title":"Vibravox: A Dataset of French Speech Captured with Body-conduction Audio Sensors","paper_url":"/paper/vibravox-a-dataset-of-french-speech-captured","paper_date":"2024-07-16","arxiv_id":"2407.11828","code_links":[{"title":"jhauret/vibravox","url":"https://github.com/jhauret/vibravox"}],"syntology":null}}],"datasets":[{"url":"/dataset/vibravox-forehead-accelerometer","name":"VibraVox (forehead accelerometer)","full_name":"","num_papers_in_archive":1},{"url":"/dataset/vibravox-headset-microphone","name":"VibraVox (headset microphone)","full_name":"","num_papers_in_archive":1},{"url":"/dataset/vibravox","name":"VibraVox (rigid in-ear microphone)","full_name":"","num_papers_in_archive":1},{"url":"/dataset/vibravox-soft-in-ear-microphone","name":"VibraVox (soft in-ear microphone)","full_name":"","num_papers_in_archive":1},{"url":"/dataset/vibravox-temple-vibration-pickup","name":"VibraVox (temple vibration pickup)","full_name":"","num_papers_in_archive":1},{"url":"/dataset/vibravox-throat-microphone","name":"VibraVox (throat microphone)","full_name":"","num_papers_in_archive":1}],"subtasks":[],"parent_tasks":[{"url":"/task/automatic-speech-recognition","name":"Automatic Speech Recognition (ASR)"}],"papers":{"order":"repositories listed in the archive (desc), then date (desc); the archive holds no stars","population":"papers tagged with this task that list at least one repository in the archive","shown":1,"of":1,"tagged_in_all":5,"items":[{"url":"/paper/vibravox-a-dataset-of-french-speech-captured","title":"Vibravox: A Dataset of French Speech Captured with Body-conduction Audio Sensors","date":"2024-07-16","arxiv_id":"2407.11828","repositories_listed":1,"syntology":null}],"syntology_records":0,"syntology_note":"a paper without a record is not a recorded non-run: it may lack an arXiv id or simply be absent from the graph layer"},"description_links":{"kept":0,"unwrapped_to_text":0,"bare_urls_linked":0,"relative_images_dropped":0,"rule":"internal links are kept only when the target slug exists in the catalog"},"syntology":{"read_at":"2026-09-25T09:33:49+00:00","claim":"Per-sample execution status on synthesized fixtures ('ran N of M samples'); not a correctness claim and not a ranking signal.","status_vocabulary":{"ran_honours":"ran, honoured the contract we drafted","ran_violates":"ran, violated the contract we drafted","ran_draft_wrong":"ran; our contract draft was wrong, not the code","ran_fixture":"ran; our fixture could not drive it","ran":"ran on a synthesized input","unverified":"unverified (harvested, no recorded run)"}},"not_shown":{"libraries":"the archive has no per-task library table","trend_sparklines":"the Trend column of the benchmarks table was a rendered image; it is not in the archive","social_and_latest_sorts":"stars and social signals are not in the archive"}}