{"about":{"site":"https://codewithpapers.app","non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page"},"url":"/paper/kapre-on-gpu-audio-preprocessing-layers-for-a","title":"Kapre: On-GPU Audio Preprocessing Layers for a Quick Implementation of Deep Neural Network Models with Keras","arxiv_id":"1706.05781","date":"2017-06-19","proceeding":null,"authors":["Keunwoo Choi","Deokjin Joo","Ju-ho Kim"],"abstract":"We introduce Kapre, Keras layers for audio and music signal preprocessing.\nMusic research using deep neural networks requires a heavy and tedious\npreprocessing stage, for which audio processing parameters are often ignored in\nparameter optimisation. To solve this problem, Kapre implements time-frequency\nconversions, normalisation, and data augmentation as Keras layers. We report\nsimple benchmark results, showing real-time on-GPU preprocessing adds a\nreasonable amount of computation.","url_abs":"http://arxiv.org/abs/1706.05781v1","url_pdf":"http://arxiv.org/pdf/1706.05781v1.pdf","source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","row_kind":"abstracts"},"code_links":[{"paper_slug":"kapre-on-gpu-audio-preprocessing-layers-for-a","repo_url":"https://github.com/keunwoochoi/kapre","is_official":1,"mentioned_in_paper":1,"mentioned_in_github":1,"framework":"tf","reach":{"status":"unanswered"}},{"paper_slug":"kapre-on-gpu-audio-preprocessing-layers-for-a","repo_url":"https://github.com/Otochess/Audio","is_official":0,"mentioned_in_paper":0,"mentioned_in_github":1,"framework":"tf","reach":{"status":"unanswered"}},{"paper_slug":"kapre-on-gpu-audio-preprocessing-layers-for-a","repo_url":"https://github.com/RishitJainn/Music-Genre-Classification-ChatBot","is_official":0,"mentioned_in_paper":0,"mentioned_in_github":1,"framework":"tf","reach":{"status":"unanswered"}},{"paper_slug":"kapre-on-gpu-audio-preprocessing-layers-for-a","repo_url":"https://github.com/godisloveforme/instrumentClassifer","is_official":0,"mentioned_in_paper":0,"mentioned_in_github":1,"framework":"tf","reach":{"status":"unanswered"}},{"paper_slug":"kapre-on-gpu-audio-preprocessing-layers-for-a","repo_url":"https://github.com/morningkaya/Audio-Classification2","is_official":0,"mentioned_in_paper":0,"mentioned_in_github":1,"framework":"tf","reach":{"status":"unanswered"}},{"paper_slug":"kapre-on-gpu-audio-preprocessing-layers-for-a","repo_url":"https://github.com/ritiksharma373/Music_genre_classification_chatbot","is_official":0,"mentioned_in_paper":0,"mentioned_in_github":1,"framework":"tf","reach":{"status":"unanswered"}},{"paper_slug":"kapre-on-gpu-audio-preprocessing-layers-for-a","repo_url":"https://github.com/seth814/Audio-Classification","is_official":0,"mentioned_in_paper":0,"mentioned_in_github":1,"framework":"tf","reach":{"status":"unanswered"}}],"tasks":[{"task_slug":"data-augmentation","task_name":"Data Augmentation"},{"task_slug":null,"task_name":"GPU"}],"methods":[],"datasets_introduced":[],"methods_introduced":[],"results":[],"syntology":{"atlas_url":"https://app.syntology.ai/?focus=1706.05781","mcp":null,"developers":"https://syntology.ai/developers"},"arxiv_metadata":null,"syntology_extracted_results":null}