{"url":"/method/characterbert","slug":"characterbert","name":"CharacterBERT","full_name":"CharacterBERT","full_name_withheld":false,"description_markdown":"CharacterBERT is a variant of [BERT](https://paperswithcode.com/method/bert) that **drops the wordpiece system** and **replaces it with a CharacterCNN module** just like the one [ELMo](https://paperswithcode.com/method/elmo) uses to produce its first layer representation. This allows CharacterBERT to represent any input token without splitting it into wordpieces. Moreover, this frees BERT from the burden of a domain-specific wordpiece vocabulary which may not be suited to your domain of interest (e.g. medical domain). Finally, it allows the model to be more robust to noisy inputs.","description_state":"present","introduced_year":null,"introduced_by":{"title":null,"paper":null,"first_author":null,"n_authors":0,"url_abs":null,"archive_paper_url":null},"source":{"url":"https://arxiv.org/abs/2010.10392v3","title":"CharacterBERT: Reconciling ELMo and BERT for Word-Level Open-Vocabulary Representations From Characters","url_on_a_paper_host":true},"code_snippet_url":"https://github.com/helboukkouri/character-bert#using-characterbert-in-practice","code_snippet_url_on_a_code_host":true,"categories":[{"area":"Natural Language Processing","area_id":"natural-language-processing","collection":"Language Models","url":"/methods/category/language-models","pwc_aliases":[]}],"n_papers_tagged":8,"archive_num_papers":null,"papers_newest_first":[{"paper":"/paper/continuous-prompt-tuning-based-textual","title":"Continuous Prompt Tuning Based Textual Entailment Model for E-commerce Entity Typing","date":"2022-11-04","arxiv_id":"2211.02483","n_code_links":1,"syntology":null},{"paper":"/paper/characterbert-and-self-teaching-for-improving","title":"CharacterBERT and Self-Teaching for Improving the Robustness of Dense Retrievers on Queries with Typos","date":"2022-04-01","arxiv_id":"2204.00716","n_code_links":1,"syntology":null},{"paper":"/paper/signal-in-noise-exploring-meaning-encoded-in","title":"Signal in Noise: Exploring Meaning Encoded in Random Character Sequences with Character-Aware Language Models","date":"2022-03-15","arxiv_id":"2203.07911","n_code_links":1,"syntology":null},{"paper":null,"title":"Exploring Meaning Encoded in Random Character Sequences with Character-Aware Language Models","date":"2021-11-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/imposing-relation-structure-in-language-model","title":"Imposing Relation Structure in Language-Model Embeddings Using Contrastive Learning","date":"2021-09-02","arxiv_id":"2109.00840","n_code_links":1,"syntology":null},{"paper":"/paper/uniparma-semeval-2021-task-5-toxic-spans","title":"UniParma at SemEval-2021 Task 5: Toxic Spans Detection Using CharacterBERT and Bag-of-Words Model","date":"2021-03-17","arxiv_id":"2103.09645","n_code_links":1,"syntology":null},{"paper":"/paper/characterbert-reconciling-elmo-and-bert-for","title":"CharacterBERT: Reconciling ELMo and BERT for Word-Level Open-Vocabulary Representations From Characters","date":"2020-10-20","arxiv_id":"2010.10392","n_code_links":2,"syntology":{"ran":5,"of":8,"unverified":3,"pointer_only":0}},{"paper":"/paper/aggressive-language-identification-using-word","title":"Aggressive Language Identification Using Word Embeddings and Sentiment Features","date":"2018-08-01","arxiv_id":null,"n_code_links":1,"syntology":null}],"papers_shown":8,"tasks":[{"task":"/task/language-modeling","name":"Language Modeling","papers":2},{"task":"/task/language-modelling","name":"Language Modelling","papers":2},{"task":"/task/natural-language-inference","name":"Natural Language Inference","papers":2},{"task":"/task/relation-extraction","name":"Relation Extraction","papers":2},{"task":"/task/aggression-identification","name":"Aggression Identification","papers":1},{"task":"/task/machine-learning","name":"BIG-bench Machine Learning","papers":1},{"task":"/task/clinical-concept-extraction","name":"Clinical Concept Extraction","papers":1},{"task":"/task/contrastive-learning","name":"Contrastive Learning","papers":1},{"task":"/task/drug-drug-interaction-extraction","name":"Drug–drug Interaction Extraction","papers":1},{"task":"/task/entity-typing","name":"Entity Typing","papers":1},{"task":"/task/language-identification","name":"Language Identification","papers":1},{"task":"/task/named-entity-recognition-1","name":"Named Entity Recognition","papers":1},{"task":"/task/named-entity-recognition-ner","name":"Named Entity Recognition (NER)","papers":1},{"task":"/task/passage-retrieval","name":"Passage Retrieval","papers":1},{"task":null,"name":"Relation","papers":1},{"task":"/task/retrieval","name":"Retrieval","papers":1},{"task":"/task/semantic-similarity","name":"Semantic Similarity","papers":1},{"task":"/task/sentence","name":"Sentence","papers":1},{"task":"/task/sentence-embeddings","name":"Sentence Embeddings","papers":1},{"task":"/task/toxic-spans-detection","name":"Toxic Spans Detection","papers":1}],"tasks_shown":20,"n_tasks":22,"usage_by_year":[{"year":"2018","papers":1},{"year":"2020","papers":1},{"year":"2021","papers":3},{"year":"2022","papers":3}],"row_source":"embedded","archive":{"source":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","archive_url":"https://paperswithcode.com/method/characterbert"},"syntology_read_at":"2026-09-24T18:15:14+00:00"}