{"about":{"site":"https://codewithpapers.app","non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page"},"url":"/paper/sense-vocabulary-compression-through-the","title":"Sense Vocabulary Compression through the Semantic Knowledge of WordNet for Neural Word Sense Disambiguation","arxiv_id":"1905.05677","date":"2019-05-14","proceeding":"GWC 2019 7","authors":["Loïc Vial","Benjamin Lecouteux","Didier Schwab"],"abstract":"In this article, we tackle the issue of the limited quantity of manually sense annotated corpora for the task of word sense disambiguation, by exploiting the semantic relationships between senses such as synonymy, hypernymy and hyponymy, in order to compress the sense vocabulary of Princeton WordNet, and thus reduce the number of different sense tags that must be observed to disambiguate all words of the lexical database. We propose two different methods that greatly reduces the size of neural WSD models, with the benefit of improving their coverage without additional training data, and without impacting their precision. In addition to our method, we present a WSD system which relies on pre-trained BERT word vectors in order to achieve results that significantly outperform the state of the art on all WSD evaluation tasks.","url_abs":"https://arxiv.org/abs/1905.05677v3","url_pdf":"https://arxiv.org/pdf/1905.05677v3.pdf","source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","row_kind":"abstracts"},"code_links":[{"paper_slug":"sense-vocabulary-compression-through-the","repo_url":"https://github.com/getalp/disambiguate","is_official":1,"mentioned_in_paper":1,"mentioned_in_github":1,"framework":"pytorch","reach":{"status":"ok","spdx":"MIT"}},{"paper_slug":"sense-vocabulary-compression-through-the","repo_url":"https://github.com/Gozzo18/WSD-Final-Homework---NLP","is_official":0,"mentioned_in_paper":0,"mentioned_in_github":1,"framework":"tf","reach":null}],"tasks":[{"task_slug":"word-sense-disambiguation","task_name":"Word Sense Disambiguation"}],"methods":[{"method_slug":"adam","method_name":"Adam"},{"method_slug":"attention","method_name":"Attention"},{"method_slug":"attention-dropout","method_name":"Attention Dropout"},{"method_slug":"bert","method_name":"BERT"},{"method_slug":"dense-connections","method_name":"Dense Connections"},{"method_slug":"dropout","method_name":"Dropout"},{"method_slug":"layer-normalization","method_name":"Layer Normalization"},{"method_slug":"linear-layer","method_name":"Linear Layer"},{"method_slug":"linear-warmup-with-linear-decay","method_name":"Linear Warmup With Linear Decay"},{"method_slug":"multi-head-attention","method_name":"Multi-Head Attention"},{"method_slug":"residual-connection","method_name":"Residual Connection"},{"method_slug":"softmax","method_name":"Softmax"},{"method_slug":"weight-decay","method_name":"Weight Decay"},{"method_slug":"wordpiece","method_name":"WordPiece"}],"datasets_introduced":[],"methods_introduced":[],"results":[{"leaderboard":"/sota/word-sense-disambiguation-on-semeval-2007","task":"Word Sense Disambiguation","dataset":"SemEval 2007 Task 17","model":"SemCor+WNGC, hypernyms","rank_in_archive_order":1,"of":9,"metrics":{"F1":"73.4"},"uses_additional_data":false},{"leaderboard":"/sota/word-sense-disambiguation-on-semeval-2007-1","task":"Word Sense Disambiguation","dataset":"SemEval 2007 Task 7","model":"SemCor+WNGC, hypernyms","rank_in_archive_order":1,"of":10,"metrics":{"F1":"90.4"},"uses_additional_data":false},{"leaderboard":"/sota/word-sense-disambiguation-on-semeval-2013","task":"Word Sense Disambiguation","dataset":"SemEval 2013 Task 12","model":"SemCor+WNGC, hypernyms","rank_in_archive_order":1,"of":12,"metrics":{"F1":"78.7"},"uses_additional_data":false},{"leaderboard":"/sota/word-sense-disambiguation-on-semeval-2015","task":"Word Sense Disambiguation","dataset":"SemEval 2015 Task 13","model":"SemCor+WNGC, hypernyms","rank_in_archive_order":1,"of":6,"metrics":{"F1":"82.6"},"uses_additional_data":false},{"leaderboard":"/sota/word-sense-disambiguation-on-senseval-2","task":"Word Sense Disambiguation","dataset":"SensEval 2","model":"SemCor+WNGC, hypernyms","rank_in_archive_order":1,"of":11,"metrics":{"F1":"79.7"},"uses_additional_data":false},{"leaderboard":"/sota/word-sense-disambiguation-on-senseval-3-task","task":"Word Sense Disambiguation","dataset":"SensEval 3 Task 1","model":"SemCor+WNGC, hypernyms","rank_in_archive_order":1,"of":11,"metrics":{"F1":"77.8"},"uses_additional_data":false},{"leaderboard":"/sota/word-sense-disambiguation-on-supervised","task":"Word Sense Disambiguation","dataset":"Supervised:","model":"SemCor+WNGC, hypernyms","rank_in_archive_order":9,"of":27,"metrics":{"SemEval 2007":"73.4","SemEval 2013":"78.7","SemEval 2015":"82.6","Senseval 2":"79.7","Senseval 3":"77.8"},"uses_additional_data":false}],"syntology":{"atlas_url":"https://app.syntology.ai/?focus=1905.05677","mcp":null,"developers":"https://syntology.ai/developers"},"arxiv_metadata":null,"syntology_extracted_results":null}