{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/bert/papers/39","list_of":"/method/bert","method":"BERT","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":39,"pages_in_order":70,"rows_per_page":100,"rows":[3801,3900],"of":6938,"counts":{"archive_papers_tagged":6938,"with_a_code_link":2862,"where_syntology_ran_a_sample":640,"not_listed_spam_title":0,"listed":6938,"listed_where_code_ran":640,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":520,"every_run_a_failure_of_syntologys_instrument":120,"listed_with_a_run_with_no_instrument_failure":520,"listed_every_run_a_failure_of_syntologys_instrument":120,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/bert","prev":"/method/bert/papers/38","next":"/method/bert/papers/40","papers":[{"paper":null,"slug":"that-is-a-good-looking-car-visual-aspect","title":"That is a good looking car !: Visual Aspect based Sentiment Controlled Personalized Response Generation","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"tree-knowledge-distillation-for-compressing","title":"Tree Knowledge Distillation for Compressing Transformer-Based Language Models","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"uncovering-surprising-event-boundaries-in","title":"Uncovering Surprising Event Boundaries in Narratives","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"understand-before-answer-improve-temporal","title":"Understand before Answer: Improve Temporal Reading Comprehension via Precise Question Understanding","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"vee-bert-accelerating-bert-inference-for","title":"VEE-BERT: Accelerating BERT Inference for Named Entity Recognition via Vote Early Exiting","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/what-do-tokens-know-about-their-characters","slug":"what-do-tokens-know-about-their-characters","title":"What do tokens know about their characters and how do they know it?","date":"2022-01-16","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"what-role-does-bert-play-in-the-neural","title":"What Role Does BERT Play in the Neural Machine Translation Encoder?","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"automatic-correction-of-syntactic-dependency","title":"Automatic Correction of Syntactic Dependency Annotation Differences","date":"2022-01-15","arxiv_id":"2201.05891","n_code_links":0,"syntology":null},{"paper":null,"slug":"automatic-lexical-simplification-for-turkish","title":"Automatic Lexical Simplification for Turkish","date":"2022-01-15","arxiv_id":"2201.05878","n_code_links":0,"syntology":null},{"paper":null,"slug":"machine-learning-for-food-review-and","title":"Machine Learning for Food Review and Recommendation","date":"2022-01-15","arxiv_id":"2201.10978","n_code_links":0,"syntology":null},{"paper":null,"slug":"polarity-and-subjectivity-detection-with","title":"Polarity and Subjectivity Detection with Multitask Learning and BERT Embedding","date":"2022-01-14","arxiv_id":"2201.05363","n_code_links":0,"syntology":null},{"paper":"/paper/knowledge-graph-augmented-network-towards","slug":"knowledge-graph-augmented-network-towards","title":"Knowledge Graph Augmented Network Towards Multiview Representation Learning for Aspect-based Sentiment Analysis","date":"2022-01-13","arxiv_id":"2201.04831","n_code_links":1,"syntology":null},{"paper":"/paper/lp-bert-multi-task-pre-training-knowledge","slug":"lp-bert-multi-task-pre-training-knowledge","title":"Multi-task Pre-training Language Model for Semantic Network Completion","date":"2022-01-13","arxiv_id":"2201.04843","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-automated-error-analysis-learning-to","title":"Towards Automated Error Analysis: Learning to Characterize Errors","date":"2022-01-13","arxiv_id":"2201.05017","n_code_links":0,"syntology":null},{"paper":"/paper/diagnosing-bert-with-retrieval-heuristics","slug":"diagnosing-bert-with-retrieval-heuristics","title":"Diagnosing BERT with Retrieval Heuristics","date":"2022-01-12","arxiv_id":"2201.04458","n_code_links":1,"syntology":null},{"paper":"/paper/generative-adversarial-network-for-text-to","slug":"generative-adversarial-network-for-text-to","title":"Generative Adversarial Network for Text-to-Face Synthesis and Manipulation with Pretrained BERT Model","date":"2022-01-12","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/promptbert-improving-bert-sentence-embeddings-1","slug":"promptbert-improving-bert-sentence-embeddings-1","title":"PromptBERT: Improving BERT Sentence Embeddings with Prompts","date":"2022-01-12","arxiv_id":"2201.04337","n_code_links":1,"syntology":{"ran":1,"of":5,"n_ran_checked":1,"n_instrument":0,"unverified":4,"pointer_only":5,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["kongds/prompt-bert"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"a-feature-extraction-based-model-for-hate","title":"A Feature Extraction based Model for Hate Speech Identification","date":"2022-01-11","arxiv_id":"2201.04227","n_code_links":0,"syntology":null},{"paper":null,"slug":"explaining-prediction-uncertainty-of-pre","title":"Explaining Predictive Uncertainty by Looking Back at Model Explanations","date":"2022-01-11","arxiv_id":"2201.03742","n_code_links":0,"syntology":null},{"paper":null,"slug":"quantifying-robustness-to-adversarial-word","title":"Quantifying Robustness to Adversarial Word Substitutions","date":"2022-01-11","arxiv_id":"2201.03829","n_code_links":0,"syntology":null},{"paper":"/paper/bert-for-sentiment-analysis-pre-trained-and","slug":"bert-for-sentiment-analysis-pre-trained-and","title":"BERT for Sentiment Analysis: Pre-trained and Fine-Tuned Alternatives","date":"2022-01-10","arxiv_id":"2201.03382","n_code_links":2,"syntology":null},{"paper":"/paper/black-box-tuning-for-language-model-as-a","slug":"black-box-tuning-for-language-model-as-a","title":"Black-Box Tuning for Language-Model-as-a-Service","date":"2022-01-10","arxiv_id":"2201.03514","n_code_links":2,"syntology":{"ran":3,"of":4,"n_ran_checked":3,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["txsun1997/black-box-tuning"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"fully-automatic-scoring-of-handwritten","title":"Handwriting recognition and automatic scoring for descriptive answers in Japanese language tests","date":"2022-01-10","arxiv_id":"2201.03215","n_code_links":0,"syntology":null},{"paper":null,"slug":"tiltedbert-resource-adjustable-version-of","title":"Latency Adjustable Transformer Encoder for Language Understanding","date":"2022-01-10","arxiv_id":"2201.03327","n_code_links":0,"syntology":null},{"paper":null,"slug":"self-training-vision-language-berts-with-a","title":"Self-Training Vision Language BERTs with a Unified Conditional Model","date":"2022-01-06","arxiv_id":"2201.02010","n_code_links":0,"syntology":null},{"paper":null,"slug":"formal-analysis-of-art-proxy-learning-of","title":"Formal Analysis of Art: Proxy Learning of Visual Concepts from Style Through Language Models","date":"2022-01-05","arxiv_id":"2201.01819","n_code_links":0,"syntology":null},{"paper":"/paper/learning-audio-visual-speech-representation-1","slug":"learning-audio-visual-speech-representation-1","title":"Learning Audio-Visual Speech Representation by Masked Multimodal Cluster Prediction","date":"2022-01-05","arxiv_id":"2201.02184","n_code_links":2,"syntology":null},{"paper":"/paper/relationship-extraction-for-knowledge-graph","slug":"relationship-extraction-for-knowledge-graph","title":"Comparison of biomedical relationship extraction methods and models for knowledge graph creation","date":"2022-01-05","arxiv_id":"2201.01647","n_code_links":0,"syntology":null},{"paper":"/paper/an-adversarial-benchmark-for-fake-news","slug":"an-adversarial-benchmark-for-fake-news","title":"An Adversarial Benchmark for Fake News Detection Models","date":"2022-01-03","arxiv_id":"2201.00912","n_code_links":1,"syntology":null},{"paper":null,"slug":"which-student-is-best-a-comprehensive","title":"Which Student is Best? A Comprehensive Knowledge Distillation Exam for Task-Specific BERT Models","date":"2022-01-03","arxiv_id":"2201.00558","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-sensitivity-of-deep-learning-based-text","title":"On Sensitivity of Deep Learning Based Text Classification Algorithms to Practical Input Perturbations","date":"2022-01-02","arxiv_id":"2201.00318","n_code_links":0,"syntology":null},{"paper":"/paper/continual-stereo-matching-of-continuous","slug":"continual-stereo-matching-of-continuous","title":"Continual Stereo Matching of Continuous Driving Scenes With Growing Architecture","date":"2022-01-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"expanding-large-pre-trained-unimodal-models","title":"Expanding Large Pre-Trained Unimodal Models With Multimodal Information Injection for Image-Text Multimodal Classification","date":"2022-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"spaceedit-learning-a-unified-editing-space-1","title":"SpaceEdit: Learning a Unified Editing Space for Open-Domain Image Color Editing","date":"2022-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/clustering-vietnamese-conversations-from","slug":"clustering-vietnamese-conversations-from","title":"Clustering Vietnamese Conversations From Facebook Page To Build Training Dataset For Chatbot","date":"2021-12-31","arxiv_id":"2112.15338","n_code_links":1,"syntology":null},{"paper":null,"slug":"automatic-mixed-precision-quantization-search","title":"Automatic Mixed-Precision Quantization Search of BERT","date":"2021-12-30","arxiv_id":"2112.14938","n_code_links":0,"syntology":null},{"paper":"/paper/dense-to-sparse-gate-for-mixture-of-experts-1","slug":"dense-to-sparse-gate-for-mixture-of-experts-1","title":"EvoMoE: An Evolutional Mixture-of-Experts Training Framework via Dense-To-Sparse Gate","date":"2021-12-29","arxiv_id":"2112.14397","n_code_links":2,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["codecaution/evomoe"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"the-university-of-texas-at-dallas-hltri-s","title":"The University of Texas at Dallas HLTRI's Participation in EPIC-QA: Searching for Entailed Questions Revealing Novel Answer Nuggets","date":"2021-12-28","arxiv_id":"2112.13946","n_code_links":0,"syntology":null},{"paper":null,"slug":"contextual-sentence-analysis-for-the","title":"Contextual Sentence Analysis for the Sentiment Prediction on Financial Data","date":"2021-12-27","arxiv_id":"2112.13790","n_code_links":0,"syntology":null},{"paper":"/paper/event-based-clinical-findings-extraction-from","slug":"event-based-clinical-findings-extraction-from","title":"Event-based clinical findings extraction from radiology reports with pre-trained language model","date":"2021-12-27","arxiv_id":"2112.13512","n_code_links":1,"syntology":null},{"paper":null,"slug":"mind-the-gap-cross-lingual-information","title":"Mind the Gap: Cross-Lingual Information Retrieval with Hierarchical Knowledge Enhancement","date":"2021-12-27","arxiv_id":"2112.13510","n_code_links":0,"syntology":null},{"paper":"/paper/multi-image-visual-question-answering","slug":"multi-image-visual-question-answering","title":"Multi-Image Visual Question Answering","date":"2021-12-27","arxiv_id":"2112.13706","n_code_links":1,"syntology":null},{"paper":null,"slug":"secondary-use-of-clinical-problem-list","title":"Secondary Use of Clinical Problem List Entries for Neural Network-Based Disease Code Assignment","date":"2021-12-27","arxiv_id":"2112.13756","n_code_links":0,"syntology":null},{"paper":null,"slug":"understanding-roberta-s-mood-the-role-of","title":"Evaluating Contextual Embeddings and their Extraction Layers for Depression Assessment","date":"2021-12-27","arxiv_id":"2112.13795","n_code_links":0,"syntology":null},{"paper":"/paper/an-ensemble-of-pre-trained-transformer-models","slug":"an-ensemble-of-pre-trained-transformer-models","title":"An Ensemble of Pre-trained Transformer Models For Imbalanced Multiclass Malware Classification","date":"2021-12-25","arxiv_id":"2112.13236","n_code_links":1,"syntology":null},{"paper":"/paper/cabace-injecting-character-sequence","slug":"cabace-injecting-character-sequence","title":"CABACE: Injecting Character Sequence Information and Domain Knowledge for Enhanced Acronym and Long-Form Extraction","date":"2021-12-25","arxiv_id":"2112.13237","n_code_links":1,"syntology":null},{"paper":"/paper/deeper-clinical-document-understanding-using","slug":"deeper-clinical-document-understanding-using","title":"Deeper Clinical Document Understanding Using Relation Extraction","date":"2021-12-25","arxiv_id":"2112.13259","n_code_links":1,"syntology":null},{"paper":"/paper/distilling-the-knowledge-of-romanian-berts","slug":"distilling-the-knowledge-of-romanian-berts","title":"Distilling the Knowledge of Romanian BERTs Using Multiple Teachers","date":"2021-12-23","arxiv_id":"2112.12650","n_code_links":1,"syntology":null},{"paper":null,"slug":"adaptive-beam-search-to-enhance-on-device","title":"Adaptive Beam Search to Enhance On-device Abstractive Summarization","date":"2021-12-22","arxiv_id":"2201.02739","n_code_links":0,"syntology":null},{"paper":null,"slug":"consistency-and-coherence-from-points-of","title":"Consistency and Coherence from Points of Contextual Similarity","date":"2021-12-22","arxiv_id":"2112.11638","n_code_links":0,"syntology":null},{"paper":null,"slug":"db-bert-a-database-tuning-tool-that-reads-the","title":"DB-BERT: a Database Tuning Tool that \"Reads the Manual\"","date":"2021-12-21","arxiv_id":"2112.10925","n_code_links":0,"syntology":null},{"paper":"/paper/predicting-job-titles-from-job-descriptions","slug":"predicting-job-titles-from-job-descriptions","title":"Predicting Job Titles from Job Descriptions with Multi-label Text Classification","date":"2021-12-21","arxiv_id":"2112.11052","n_code_links":1,"syntology":null},{"paper":null,"slug":"training-dataset-and-dictionary-sizes-matter","title":"Training dataset and dictionary sizes matter in BERT models: the case of Baltic languages","date":"2021-12-20","arxiv_id":"2112.10553","n_code_links":0,"syntology":null},{"paper":null,"slug":"data-augmentation-for-mental-health","title":"Data Augmentation for Mental Health Classification on Social Media","date":"2021-12-19","arxiv_id":"2112.10064","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-transformers-for-hate-speech","title":"Leveraging Transformers for Hate Speech Detection in Conversational Code-Mixed Tweets","date":"2021-12-18","arxiv_id":"2112.09986","n_code_links":0,"syntology":null},{"paper":null,"slug":"low-resource-learning-with-knowledge-graphs-a","title":"Zero-shot and Few-shot Learning with Knowledge Graphs: A Comprehensive Survey","date":"2021-12-18","arxiv_id":"2112.10006","n_code_links":0,"syntology":null},{"paper":null,"slug":"syntactic-gcn-bert-based-chinese-event","title":"Syntactic-GCN Bert based Chinese Event Extraction","date":"2021-12-18","arxiv_id":"2112.09939","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-high-precision-health-relatedness-score-for","title":"A High-Precision Health-relatedness Score for Phrases to Mine Cause-Effect Statements from the Web","date":"2021-12-17","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"can-machine-learning-tools-support-the","title":"Can Machine Learning Tools Support the Identification of Sustainable Design Leads From Product Reviews? Opportunities and Challenges","date":"2021-12-17","arxiv_id":"2112.09391","n_code_links":0,"syntology":null},{"paper":null,"slug":"challenging-america-modeling-language-in","title":"Challenging America: Modeling language in longer time scales","date":"2021-12-17","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/explain-edit-and-understand-rethinking-user","slug":"explain-edit-and-understand-rethinking-user","title":"Explain, Edit, and Understand: Rethinking User Study Design for Evaluating Model Explanations","date":"2021-12-17","arxiv_id":"2112.09669","n_code_links":1,"syntology":null},{"paper":null,"slug":"joint-chinese-word-segmentation-and-part-of-2","title":"Joint Chinese Word Segmentation and Part-of-speech Tagging via Two-stage Span Labeling","date":"2021-12-17","arxiv_id":"2112.09488","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-to-win-lottery-tickets-in-bert","title":"Learning to Win Lottery Tickets in BERT Transfer via Task-agnostic Mask Training","date":"2021-12-17","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"rank4class-a-ranking-formulation-for","title":"Rank4Class: A Ranking Formulation for Multiclass Classification","date":"2021-12-17","arxiv_id":"2112.09727","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-faithful-personalized-response","title":"Towards Faithful Personalized Response Selection in Retrieval Based Dialog Systems","date":"2021-12-17","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"an-empirical-study-on-transfer-learning-for","title":"An Empirical Study on Transfer Learning for Privilege Review","date":"2021-12-16","arxiv_id":"2112.08606","n_code_links":0,"syntology":null},{"paper":"/paper/commonsense-knowledge-augmented-pretrained-1","slug":"commonsense-knowledge-augmented-pretrained-1","title":"Knowledge-Augmented Language Models for Cause-Effect Relation Classification","date":"2021-12-16","arxiv_id":"2112.08615","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["phosseini/causal-reasoning"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"does-pre-training-induce-systematic-inference","title":"Does Pre-training Induce Systematic Inference? How Masked Language Models Acquire Commonsense Knowledge","date":"2021-12-16","arxiv_id":"2112.08583","n_code_links":0,"syntology":null},{"paper":null,"slug":"3d-question-answering","title":"3D Question Answering","date":"2021-12-15","arxiv_id":"2112.08359","n_code_links":0,"syntology":null},{"paper":null,"slug":"applying-softtriple-loss-for-supervised","title":"Applying SoftTriple Loss for Supervised Language Model Fine Tuning","date":"2021-12-15","arxiv_id":"2112.08462","n_code_links":0,"syntology":null},{"paper":null,"slug":"fine-tuning-large-neural-language-models-for","title":"Fine-Tuning Large Neural Language Models for Biomedical Natural Language Processing","date":"2021-12-15","arxiv_id":"2112.07869","n_code_links":0,"syntology":null},{"paper":"/paper/one-size-does-not-fit-all-investigating","slug":"one-size-does-not-fit-all-investigating","title":"One size does not fit all: Investigating strategies for differentially-private learning across NLP tasks","date":"2021-12-15","arxiv_id":"2112.08159","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["trusthlt/dp-across-nlp-tasks"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"one-system-to-rule-them-all-a-universal","title":"One System to Rule them All: a Universal Intent Recognition System for Customer Service Chatbots","date":"2021-12-15","arxiv_id":"2112.08261","n_code_links":0,"syntology":null},{"paper":null,"slug":"tracing-text-provenance-via-context-aware","title":"Tracing Text Provenance via Context-Aware Lexical Substitution","date":"2021-12-15","arxiv_id":"2112.07873","n_code_links":0,"syntology":null},{"paper":null,"slug":"ace-bert-adversarial-cross-modal-enhanced","title":"ACE-BERT: Adversarial Cross-modal Enhanced BERT for E-commerce Retrieval","date":"2021-12-14","arxiv_id":"2112.07209","n_code_links":0,"syntology":null},{"paper":null,"slug":"building-on-huang-et-al-glossbert-for-word","title":"Building on Huang et al. GlossBERT for Word Sense Disambiguation","date":"2021-12-14","arxiv_id":"2112.07089","n_code_links":0,"syntology":null},{"paper":null,"slug":"classifying-emails-into-human-vs-machine","title":"Classifying Emails into Human vs Machine Category","date":"2021-12-14","arxiv_id":"2112.07742","n_code_links":0,"syntology":null},{"paper":null,"slug":"coco-bert-improving-video-language-pre","title":"CoCo-BERT: Improving Video-Language Pre-training with Contrastive Cross-modal Matching and Denoising","date":"2021-12-14","arxiv_id":"2112.07515","n_code_links":0,"syntology":null},{"paper":null,"slug":"epigenomic-language-models-powered-by","title":"Epigenomic language models powered by Cerebras","date":"2021-12-14","arxiv_id":"2112.07571","n_code_links":0,"syntology":null},{"paper":"/paper/from-dense-to-sparse-contrastive-pruning-for","slug":"from-dense-to-sparse-contrastive-pruning-for","title":"From Dense to Sparse: Contrastive Pruning for Better Pre-trained Language Model Compression","date":"2021-12-14","arxiv_id":"2112.07198","n_code_links":2,"syntology":null},{"paper":"/paper/measuring-fairness-with-biased-rulers-a","slug":"measuring-fairness-with-biased-rulers-a","title":"Measuring Fairness with Biased Rulers: A Survey on Quantifying Biases in Pretrained Language Models","date":"2021-12-14","arxiv_id":"2112.07447","n_code_links":1,"syntology":null},{"paper":"/paper/text-classification-models-for-form-entity","slug":"text-classification-models-for-form-entity","title":"Text Classification Models for Form Entity Linking","date":"2021-12-14","arxiv_id":"2112.07443","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-a-unified-foundation-model-jointly","title":"Towards a Unified Foundation Model: Jointly Pre-Training Transformers on Unpaired Images and Text","date":"2021-12-14","arxiv_id":"2112.07074","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-study-on-token-pruning-for-colbert","title":"A Study on Token Pruning for ColBERT","date":"2021-12-13","arxiv_id":"2112.06540","n_code_links":0,"syntology":null},{"paper":null,"slug":"context-vs-target-word-quantifying-biases-in","title":"Measuring Context-Word Biases in Lexical Semantic Datasets","date":"2021-12-13","arxiv_id":"2112.06733","n_code_links":0,"syntology":null},{"paper":null,"slug":"do-data-based-curricula-work","title":"Do Data-based Curricula Work?","date":"2021-12-13","arxiv_id":"2112.06510","n_code_links":0,"syntology":null},{"paper":null,"slug":"roof-bert-divide-understanding-labour-and","title":"Roof-Transformer: Divided and Joined Understanding with Knowledge Enhancement","date":"2021-12-13","arxiv_id":"2112.06736","n_code_links":0,"syntology":null},{"paper":"/paper/wechsel-effective-initialization-of-subword-1","slug":"wechsel-effective-initialization-of-subword-1","title":"WECHSEL: Effective initialization of subword embeddings for cross-lingual transfer of monolingual language models","date":"2021-12-13","arxiv_id":"2112.06598","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["cpjku/wechsel"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"findings-on-conversation-disentanglement","title":"Findings on Conversation Disentanglement","date":"2021-12-10","arxiv_id":"2112.05346","n_code_links":0,"syntology":null},{"paper":"/paper/multimodal-interactions-using-pretrained","slug":"multimodal-interactions-using-pretrained","title":"Multimodal Interactions Using Pretrained Unimodal Models for SIMMC 2.0","date":"2021-12-10","arxiv_id":"2112.05328","n_code_links":1,"syntology":null},{"paper":"/paper/detecting-potentially-harmful-and-protective","slug":"detecting-potentially-harmful-and-protective","title":"Detecting potentially harmful and protective suicide-related content on twitter: A machine learning approach","date":"2021-12-09","arxiv_id":"2112.04796","n_code_links":2,"syntology":null},{"paper":null,"slug":"from-scattered-sources-to-comprehensive","title":"From Scattered Sources to Comprehensive Technology Landscape: A Recommendation-based Retrieval Approach","date":"2021-12-09","arxiv_id":"2112.04810","n_code_links":0,"syntology":null},{"paper":"/paper/semantic-search-as-extractive-paraphrase-span-1","slug":"semantic-search-as-extractive-paraphrase-span-1","title":"Semantic Search as Extractive Paraphrase Span Detection","date":"2021-12-09","arxiv_id":"2112.04886","n_code_links":1,"syntology":null},{"paper":"/paper/improving-language-models-by-retrieving-from","slug":"improving-language-models-by-retrieving-from","title":"Improving language models by retrieving from trillions of tokens","date":"2021-12-08","arxiv_id":"2112.04426","n_code_links":2,"syntology":{"ran":16,"of":23,"n_ran_checked":14,"n_instrument":2,"unverified":7,"pointer_only":3,"phrase":"16 ran (of which 5 constructed an object rather than computing a result; 14 with no instrument failure: 0 honoured, 3 violated, 11 with no contract checked; 2 where Syntology's instrument failed) · 7 unverified","official":null}},{"paper":"/paper/jaber-junior-arabic-bert","slug":"jaber-junior-arabic-bert","title":"JABER and SABER: Junior and Senior Arabic BERt","date":"2021-12-08","arxiv_id":"2112.04329","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-transferable-approach-for-partitioning","title":"A Transferable Approach for Partitioning Machine Learning Models on Multi-Chip-Modules","date":"2021-12-07","arxiv_id":"2112.04041","n_code_links":0,"syntology":null},{"paper":"/paper/racebert-a-transformer-based-model-for","slug":"racebert-a-transformer-based-model-for","title":"raceBERT -- A Transformer-based Model for Predicting Race and Ethnicity from Names","date":"2021-12-07","arxiv_id":"2112.03807","n_code_links":1,"syntology":null},{"paper":"/paper/bertmap-a-bert-based-ontology-alignment","slug":"bertmap-a-bert-based-ontology-alignment","title":"BERTMap: A BERT-based Ontology Alignment System","date":"2021-12-05","arxiv_id":"2112.02682","n_code_links":1,"syntology":null},{"paper":"/paper/causal-distillation-for-language-models","slug":"causal-distillation-for-language-models","title":"Causal Distillation for Language Models","date":"2021-12-05","arxiv_id":"2112.02505","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":3,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["frankaging/Causal-Distill"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/dibert-dependency-injected-bidirectional","slug":"dibert-dependency-injected-bidirectional","title":"DIBERT: Dependency Injected Bidirectional Encoder Representations from Transformers","date":"2021-12-05","arxiv_id":null,"n_code_links":1,"syntology":null}],"record_sha256":"02f23bd529a00f39a7a8a0c3929af2c612ba5b444ff07f9fcb0a653ff5f9b218","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}