{"about":{"site":"https://codewithpapers.app","non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page"},"url":"/code/feature-loss","entry":"feature_loss","source":"Syntology graph, per-sample; not an archive number","read_at":"2026-09-24T18:15:14+00:00","claim":"Names are grouped by exact entry-name string. Same-named routines are NOT asserted to be equivalent; 'ran' means executed on a synthesized fixture, not correctness. n_samples_ran = sum of by_status over every status except 'unverified' (ran_draft_wrong and ran_fixture are failures of Syntology's instrument, not of the code); n_papers_ran = papers with at least one such sample.","status_vocabulary":{"ran_honours":"ran, honoured the contract we drafted","ran_violates":"ran, violated the contract we drafted","ran_draft_wrong":"ran; our contract draft was wrong, not the code","ran_fixture":"ran; our fixture could not drive it","ran":"ran on a synthesized input","unverified":"unverified (harvested, no recorded run)"},"n_papers":27,"n_papers_ran":20,"units":"n_samples, n_samples_ran, n_samples_fingerprinted and by_status count distinct code bodies (code_sha256); n_places and n_places_pointer_only count places, one per (paper, code body) pair, which is also the unit of the samples list","n_samples":8,"n_samples_ran":5,"n_samples_fingerprinted":0,"n_places":28,"n_places_pointer_only":5,"by_status":{"ran_honours":1,"ran_violates":0,"ran_draft_wrong":0,"ran_fixture":0,"ran":4,"unverified":3},"syntology":{"atlas_url":null,"mcp":null,"mcp_per_sample":{"tool":"get_code","arguments_in":"samples[].mcp_get_code"},"developers":"https://syntology.ai/developers"},"samples":[{"arxiv_id":"2504.09839","paper":"/paper/safespeech-robust-and-universal-voice","title":"SafeSpeech: Robust and Universal Voice Protection Against Malicious Speech Synthesis","date":"2025-04-14","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"wxzyd123/safespeech","path":"bert_vits2/losses.py","file_url":"https://github.com/wxzyd123/safespeech/blob/HEAD/bert_vits2/losses.py","status":"unverified","verification_level":0,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"24e8b4731dc3e070","mcp_get_code":{"code_sha256":"24e8b4731dc3e070"}},{"arxiv_id":"2501.04926","paper":"/paper/flowhigh-towards-efficient-and-high-quality","title":"FLowHigh: Towards Efficient and High-Quality Audio Super-Resolution with Single-Step Flow Matching","date":"2025-01-09","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"jjunak-yun/FLowHigh_code","path":"vocoder/BIGVGAN/bigvgan/models.py","file_url":"https://github.com/jjunak-yun/FLowHigh_code/blob/HEAD/vocoder/BIGVGAN/bigvgan/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2410.20742","paper":"/paper/mitigating-unauthorized-speech-synthesis-for","title":"Mitigating Unauthorized Speech Synthesis for Voice Protection","date":"2024-10-28","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"wxzyd123/pivotal_objective_perturbation","path":"vits/losses.py","file_url":"https://github.com/wxzyd123/pivotal_objective_perturbation/blob/HEAD/vits/losses.py","status":"unverified","verification_level":0,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"NONE","inline_ok":false,"code_sha256_prefix":"24e8b4731dc3e070","mcp_get_code":{"code_sha256":"24e8b4731dc3e070"}},{"arxiv_id":"2409.12121","paper":"/paper/wmcodec-end-to-end-neural-speech-codec-with","title":"WMCodec: End-to-End Neural Speech Codec with Deep Watermarking for Authenticity Verification","date":null,"month_inferred_from_arxiv_id":"2024-09","title_source":"archive","repo":"zjzser/wmcodec","path":"models.py","file_url":"https://github.com/zjzser/wmcodec/blob/HEAD/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"NONE","inline_ok":false,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2406.10581","paper":"/paper/crossfuse-a-novel-cross-attention-mechanism","title":"CrossFuse: A Novel Cross Attention Mechanism based Infrared and Visible Image Fusion Approach","date":"2024-06-15","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"hli1221/crossfuse","path":"network/loss.py","file_url":"https://github.com/hli1221/crossfuse/blob/HEAD/network/loss.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"NONE","inline_ok":false,"code_sha256_prefix":"849efbd3c9fe8c54","mcp_get_code":{"code_sha256":"849efbd3c9fe8c54"}},{"arxiv_id":"2406.02250","paper":"/paper/multi-stage-speech-bandwidth-extension-with","title":"Multi-Stage Speech Bandwidth Extension with Flexible Sampling Rate Control","date":"2024-06-04","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"yxlu-0102/AP-BWE","path":"models/model.py","file_url":"https://github.com/yxlu-0102/AP-BWE/blob/HEAD/models/model.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"4518cce23c4f900e","mcp_get_code":{"code_sha256":"4518cce23c4f900e"}},{"arxiv_id":"2406.00320","paper":"/paper/frieren-efficient-video-to-audio-generation","title":"Frieren: Efficient Video-to-Audio Generation Network with Rectified Flow Matching","date":"2024-06-01","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"cyanbx/Frieren-V2A","path":"Frieren/vocoder/bigvgan/models.py","file_url":"https://github.com/cyanbx/Frieren-V2A/blob/HEAD/Frieren/vocoder/bigvgan/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"Apache-2.0","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2311.14957","paper":"/paper/multi-scale-sub-band-constant-q-transform","title":"Multi-Scale Sub-Band Constant-Q Transform Discriminator for High-Fidelity Vocoder","date":null,"month_inferred_from_arxiv_id":"2023-11","title_source":"archive","repo":"nvidia/bigvgan","path":"loss.py","file_url":"https://github.com/nvidia/bigvgan/blob/HEAD/loss.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"95f8dc8b58c53187","mcp_get_code":{"code_sha256":"95f8dc8b58c53187"}},{"arxiv_id":"2310.01889","paper":"/paper/ring-attention-with-blockwise-transformers","title":"Ring Attention with Blockwise Transformers for Near-Infinite Context","date":"2023-10-03","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"seongho608/ringformer","path":"losses.py","file_url":"https://github.com/seongho608/ringformer/blob/HEAD/losses.py","status":"unverified","verification_level":0,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"24e8b4731dc3e070","mcp_get_code":{"code_sha256":"24e8b4731dc3e070"}},{"arxiv_id":"2310.00014","paper":"/paper/fewer-token-neural-speech-codec-with-time","title":"Fewer-token Neural Speech Codec with Time-invariant Codes","date":null,"month_inferred_from_arxiv_id":"2023-10","title_source":"archive","repo":"y-ren16/ticodec","path":"academicodec/models/ticodec/models.py","file_url":"https://github.com/y-ren16/ticodec/blob/HEAD/academicodec/models/ticodec/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"NONE","inline_ok":false,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2306.00814","paper":"/paper/vocos-closing-the-gap-between-time-domain-and","title":"Vocos: Closing the gap between time-domain and Fourier-based neural vocoders for high-quality audio synthesis","date":"2023-06-01","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"IAHispano/Applio","path":"rvc/train/losses.py","file_url":"https://github.com/IAHispano/Applio/blob/HEAD/rvc/train/losses.py","status":"unverified","verification_level":0,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"0f7c6286a8d000f3","mcp_get_code":{"code_sha256":"0f7c6286a8d000f3"}},{"arxiv_id":"2305.19709","paper":"/paper/xphonebert-a-pre-trained-multilingual-model","title":"XPhoneBERT: A Pre-trained Multilingual Model for Phoneme Representations for Text-to-Speech","date":"2023-05-31","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"vinairesearch/xphonebert","path":"VITS_with_XPhoneBERT/losses.py","file_url":"https://github.com/vinairesearch/xphonebert/blob/HEAD/VITS_with_XPhoneBERT/losses.py","status":"unverified","verification_level":0,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"24e8b4731dc3e070","mcp_get_code":{"code_sha256":"24e8b4731dc3e070"}},{"arxiv_id":"2305.13516","paper":"/paper/scaling-speech-technology-to-1000-languages-1","title":"Scaling Speech Technology to 1,000+ Languages","date":"2023-05-22","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"ylacombe/finetune-hf-vits","path":"run_vits_finetuning.py","file_url":"https://github.com/ylacombe/finetune-hf-vits/blob/HEAD/run_vits_finetuning.py","status":"ran_honours","verification_level":1,"contract_check":"HONOURS","metamorphic_tier":"well_formed","behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"8feddae84063cbbc","mcp_get_code":{"code_sha256":"8feddae84063cbbc"}},{"arxiv_id":"2305.06908","paper":"/paper/comospeech-one-step-speech-and-singing-voice","title":"CoMoSpeech: One-Step Speech and Singing Voice Synthesis via Consistency Model","date":"2023-05-11","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"zhenye234/CoMoSpeech","path":"hifi-gan/models.py","file_url":"https://github.com/zhenye234/CoMoSpeech/blob/HEAD/hifi-gan/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2303.05309","paper":"/paper/mixspeech-cross-modality-self-learning-with","title":"MixSpeech: Cross-Modality Self-Learning with Audio-Visual Stream Mixup for Visual Speech Translation and Recognition","date":"2023-03-09","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"rongjiehuang/transpeech","path":"research/TranSpeech/hifigan/models.py","file_url":"https://github.com/rongjiehuang/transpeech/blob/HEAD/research/TranSpeech/hifigan/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2212.14227","paper":"/paper/styletts-vc-one-shot-voice-conversion-by","title":"StyleTTS-VC: One-Shot Voice Conversion by Knowledge Transfer from Style-Based TTS Models","date":"2022-12-29","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"yl4579/StyleTTS-VC","path":"Demo/hifi-gan/vocoder.py","file_url":"https://github.com/yl4579/StyleTTS-VC/blob/HEAD/Demo/hifi-gan/vocoder.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2212.09730","paper":"/paper/speaking-style-conversion-with-discrete-self","title":"Speaking Style Conversion in the Waveform Domain Using Discrete Self-Supervised Units","date":"2022-12-19","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"gallilmaimon/DISSC","path":"sr/models.py","file_url":"https://github.com/gallilmaimon/DISSC/blob/HEAD/sr/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2206.04658","paper":"/paper/bigvgan-a-universal-neural-vocoder-with-large","title":"BigVGAN: A Universal Neural Vocoder with Large-Scale Training","date":"2022-06-09","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"sh-lee-prml/BigVGAN","path":"losses.py","file_url":"https://github.com/sh-lee-prml/BigVGAN/blob/HEAD/losses.py","status":"unverified","verification_level":0,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"24e8b4731dc3e070","mcp_get_code":{"code_sha256":"24e8b4731dc3e070"}},{"arxiv_id":"2205.15439","paper":"/paper/styletts-a-style-based-generative-model-for","title":"StyleTTS: A Style-Based Generative Model for Natural and Diverse Text-to-Speech Synthesis","date":"2022-05-30","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"yl4579/StyleTTS","path":"Demo/hifi-gan/vocoder.py","file_url":"https://github.com/yl4579/StyleTTS/blob/HEAD/Demo/hifi-gan/vocoder.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2205.04421","paper":"/paper/naturalspeech-end-to-end-text-to-speech","title":"NaturalSpeech: End-to-End Text to Speech Synthesis with Human-Level Quality","date":"2022-05-09","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"daniilrobnikov/vits2","path":"losses.py","file_url":"https://github.com/daniilrobnikov/vits2/blob/HEAD/losses.py","status":"unverified","verification_level":0,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"f9a2eedb621a3dda","mcp_get_code":{"code_sha256":"f9a2eedb621a3dda"}},{"arxiv_id":"2203.13086","paper":"/paper/hifi-a-unified-framework-for-neural-vocoding","title":"HiFi++: a Unified Framework for Bandwidth Extension and Speech Enhancement","date":"2022-03-24","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"rishikksh20/HiFiplusplus-pytorch","path":"models.py","file_url":"https://github.com/rishikksh20/HiFiplusplus-pytorch/blob/HEAD/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2203.02395","paper":"/paper/istftnet-fast-and-lightweight-mel-spectrogram","title":"iSTFTNet: Fast and Lightweight Mel-Spectrogram Vocoder Incorporating Inverse Short-Time Fourier Transform","date":"2022-03-04","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"hcy71o/autovocoder","path":"models.py","file_url":"https://github.com/hcy71o/autovocoder/blob/HEAD/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"Apache-2.0","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2202.13277","paper":"/paper/learning-the-beauty-in-songs-neural-singing","title":"Learning the Beauty in Songs: Neural Singing Voice Beautifier","date":"2022-02-27","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"MoonInTheRiver/DiffSinger","path":"modules/hifigan/hifigan.py","file_url":"https://github.com/MoonInTheRiver/DiffSinger/blob/HEAD/modules/hifigan/hifigan.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2109.13821","paper":"/paper/diffusion-based-voice-conversion-with-fast","title":"Diffusion-Based Voice Conversion with Fast Maximum Likelihood Sampling Scheme","date":"2021-09-28","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"trinhtuanvubk/diff-vc","path":"hifi-gan/models.py","file_url":"https://github.com/trinhtuanvubk/diff-vc/blob/HEAD/hifi-gan/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":false,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2106.07889","paper":"/paper/univnet-a-neural-vocoder-with-multi","title":"UnivNet: A Neural Vocoder with Multi-Resolution Spectrogram Discriminators for High-Fidelity Waveform Generation","date":"2021-06-15","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"rishikksh20/UnivNet-pytorch","path":"loss.py","file_url":"https://github.com/rishikksh20/UnivNet-pytorch/blob/HEAD/loss.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2106.06103","paper":"/paper/conditional-variational-autoencoder-with","title":"Conditional Variational Autoencoder with Adversarial Learning for End-to-End Text-to-Speech","date":null,"month_inferred_from_arxiv_id":"2021-06","title_source":"archive","repo":"maum-ai/phaseaug","path":"models.py","file_url":"https://github.com/maum-ai/phaseaug/blob/HEAD/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"BSD-3-Clause","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}},{"arxiv_id":"2106.06103","paper":"/paper/conditional-variational-autoencoder-with","title":"Conditional Variational Autoencoder with Adversarial Learning for End-to-End Text-to-Speech","date":null,"month_inferred_from_arxiv_id":"2021-06","title_source":"archive","repo":"jaywalnut310/vits","path":"losses.py","file_url":"https://github.com/jaywalnut310/vits/blob/HEAD/losses.py","status":"unverified","verification_level":0,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"24e8b4731dc3e070","mcp_get_code":{"code_sha256":"24e8b4731dc3e070"}},{"arxiv_id":"2006.04558","paper":"/paper/fastspeech-2-fast-and-high-quality-end-to-end","title":"FastSpeech 2: Fast and High-Quality End-to-End Text to Speech","date":"2020-06-08","month_inferred_from_arxiv_id":null,"title_source":"archive","repo":"shivammehta25/BetterFastSpeech2","path":"fs2/hifigan/models.py","file_url":"https://github.com/shivammehta25/BetterFastSpeech2/blob/HEAD/fs2/hifigan/models.py","status":"ran","verification_level":1,"contract_check":null,"metamorphic_tier":null,"behaviour_fingerprint":false,"licence":"MIT","inline_ok":true,"code_sha256_prefix":"e453b51f0ed5fb28","mcp_get_code":{"code_sha256":"e453b51f0ed5fb28"}}]}