{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/bert/papers/32","list_of":"/method/bert","method":"BERT","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":32,"pages_in_order":70,"rows_per_page":100,"rows":[3101,3200],"of":6938,"counts":{"archive_papers_tagged":6938,"with_a_code_link":2862,"where_syntology_ran_a_sample":640,"not_listed_spam_title":0,"listed":6938,"listed_where_code_ran":640,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":520,"every_run_a_failure_of_syntologys_instrument":120,"listed_with_a_run_with_no_instrument_failure":520,"listed_every_run_a_failure_of_syntologys_instrument":120,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/bert","prev":"/method/bert/papers/31","next":"/method/bert/papers/33","papers":[{"paper":null,"slug":"l3cube-hindbert-and-devbert-pre-trained-bert","title":"L3Cube-HindBERT and DevBERT: Pre-Trained BERT Transformer models for Devanagari based Hindi and Marathi Languages","date":"2022-11-21","arxiv_id":"2211.11418","n_code_links":0,"syntology":null},{"paper":"/paper/l3cube-mahasbert-and-hindsbert-sentence-bert","slug":"l3cube-mahasbert-and-hindsbert-sentence-bert","title":"L3Cube-MahaSBERT and HindSBERT: Sentence BERT Models and Benchmarking BERT Sentence Representations for Hindi and Marathi","date":"2022-11-21","arxiv_id":"2211.11187","n_code_links":1,"syntology":null},{"paper":null,"slug":"tcbert-a-technical-report-for-chinese-topic","title":"TCBERT: A Technical Report for Chinese Topic Classification BERT","date":"2022-11-21","arxiv_id":"2211.11304","n_code_links":0,"syntology":null},{"paper":null,"slug":"conceptor-aided-debiasing-of-contextualized","title":"Conceptor-Aided Debiasing of Large Language Models","date":"2022-11-20","arxiv_id":"2211.11087","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-conspiracy-theory-against-covid-19","title":"Detecting Conspiracy Theory Against COVID-19 Vaccines","date":"2022-11-20","arxiv_id":"2211.13003","n_code_links":0,"syntology":null},{"paper":null,"slug":"feature-weaken-vicinal-data-augmentation-for","title":"Feature Weaken: Vicinal Data Augmentation for Classification","date":"2022-11-20","arxiv_id":"2211.10944","n_code_links":0,"syntology":null},{"paper":"/paper/understanding-and-improving-knowledge-1","slug":"understanding-and-improving-knowledge-1","title":"Understanding and Improving Knowledge Distillation for Quantization-Aware Training of Large Transformer Encoders","date":"2022-11-20","arxiv_id":"2211.11014","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-survey-on-knowledge-enhanced-multimodal","title":"A survey on knowledge-enhanced multimodal learning","date":"2022-11-19","arxiv_id":"2211.12328","n_code_links":0,"syntology":null},{"paper":null,"slug":"entity-assisted-language-models-for","title":"Entity-Assisted Language Models for Identifying Check-worthy Sentences","date":"2022-11-19","arxiv_id":"2211.10678","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-users-social-network-embeddings","title":"Leveraging Users' Social Network Embeddings for Fake News Detection on Twitter","date":"2022-11-19","arxiv_id":"2211.10672","n_code_links":0,"syntology":null},{"paper":null,"slug":"metadata-might-make-language-models-better","title":"Metadata Might Make Language Models Better","date":"2022-11-18","arxiv_id":"2211.10086","n_code_links":0,"syntology":null},{"paper":null,"slug":"where-did-you-tweet-from-inferring-the-origin","title":"Where did you tweet from? Inferring the origin locations of tweets based on contextual information","date":"2022-11-18","arxiv_id":"2211.16506","n_code_links":0,"syntology":null},{"paper":null,"slug":"longfnt-long-form-speech-recognition-with","title":"LongFNT: Long-form Speech Recognition with Factorized Neural Transducer","date":"2022-11-17","arxiv_id":"2211.09412","n_code_links":0,"syntology":null},{"paper":"/paper/protsi-prototypical-siamese-network-with-data","slug":"protsi-prototypical-siamese-network-with-data","title":"ProtSi: Prototypical Siamese Network with Data Augmentation for Few-Shot Subjective Answer Evaluation","date":"2022-11-17","arxiv_id":"2211.09855","n_code_links":1,"syntology":null},{"paper":"/paper/random-ltd-random-and-layerwise-token","slug":"random-ltd-random-and-layerwise-token","title":"Random-LTD: Random and Layerwise Token Dropping Brings Efficient Training for Large-scale Transformers","date":"2022-11-17","arxiv_id":"2211.11586","n_code_links":1,"syntology":null},{"paper":null,"slug":"fast-and-accurate-fsa-system-using-elbert-an","title":"Fast and Accurate FSA System Using ELBERT: An Efficient and Lightweight BERT","date":"2022-11-16","arxiv_id":"2211.08842","n_code_links":0,"syntology":null},{"paper":"/paper/an-fnet-based-auto-encoder-for-long-sequence","slug":"an-fnet-based-auto-encoder-for-long-sequence","title":"An FNet based Auto Encoder for Long Sequence News Story Generation","date":"2022-11-15","arxiv_id":"2211.08295","n_code_links":1,"syntology":null},{"paper":null,"slug":"empowering-language-models-with-knowledge","title":"Empowering Language Models with Knowledge Graph Reasoning for Question Answering","date":"2022-11-15","arxiv_id":"2211.08380","n_code_links":0,"syntology":null},{"paper":null,"slug":"robbert-2022-updating-a-dutch-language-model","title":"RobBERT-2022: Updating a Dutch Language Model to Account for Evolving Language Use","date":"2022-11-15","arxiv_id":"2211.08192","n_code_links":0,"syntology":null},{"paper":"/paper/greenplm-cross-lingual-pre-trained-language","slug":"greenplm-cross-lingual-pre-trained-language","title":"GreenPLM: Cross-Lingual Transfer of Monolingual Pre-Trained Language Models at Almost No Cost","date":"2022-11-13","arxiv_id":"2211.06993","n_code_links":1,"syntology":null},{"paper":"/paper/xu-at-semeval-2022-task-4-pre-bert-neural-1","slug":"xu-at-semeval-2022-task-4-pre-bert-neural-1","title":"Xu at SemEval-2022 Task 4: Pre-BERT Neural Network Methods vs Post-BERT RoBERTa Approach for Patronizing and Condescending Language Detection","date":"2022-11-13","arxiv_id":"2211.06874","n_code_links":1,"syntology":null},{"paper":"/paper/dark-patterns-in-e-commerce-a-dataset-and-its","slug":"dark-patterns-in-e-commerce-a-dataset-and-its","title":"Dark patterns in e-commerce: a dataset and its baseline evaluations","date":"2022-11-12","arxiv_id":"2211.06543","n_code_links":1,"syntology":null},{"paper":"/paper/misinformation-detection-using-persuasive","slug":"misinformation-detection-using-persuasive","title":"Using Persuasive Writing Strategies to Explain and Detect Health Misinformation","date":"2022-11-11","arxiv_id":"2211.05985","n_code_links":1,"syntology":null},{"paper":null,"slug":"bert-based-combination-of-convolutional-and","title":"BERT-Based Combination of Convolutional and Recurrent Neural Network for Indonesian Sentiment Analysis","date":"2022-11-10","arxiv_id":"2211.05273","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-in-plutarch-s-shadows","title":"BERT in Plutarch's Shadows","date":"2022-11-10","arxiv_id":"2211.05673","n_code_links":0,"syntology":null},{"paper":null,"slug":"biomedical-multi-hop-question-answering-using","title":"Biomedical Multi-hop Question Answering Using Knowledge Graph Embeddings and Language Models","date":"2022-11-10","arxiv_id":"2211.05351","n_code_links":0,"syntology":null},{"paper":"/paper/cherry-hypothesis-identifying-the-cherry-on","slug":"cherry-hypothesis-identifying-the-cherry-on","title":"PAD-Net: An Efficient Framework for Dynamic Networks","date":"2022-11-10","arxiv_id":"2211.05528","n_code_links":1,"syntology":null},{"paper":null,"slug":"syntax-guided-domain-adaptation-for-aspect","title":"Syntax-Guided Domain Adaptation for Aspect-based Sentiment Analysis","date":"2022-11-10","arxiv_id":"2211.05457","n_code_links":0,"syntology":null},{"paper":"/paper/collateral-facilitation-in-humans-and","slug":"collateral-facilitation-in-humans-and","title":"Collateral facilitation in humans and language models","date":"2022-11-09","arxiv_id":"2211.05198","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["jmichaelov/collateral-facilitation"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"cross-lingual-transfer-learning-for-check","title":"Cross-lingual Transfer Learning for Check-worthy Claim Identification over Twitter","date":"2022-11-09","arxiv_id":"2211.05087","n_code_links":0,"syntology":null},{"paper":"/paper/mask-more-and-mask-later-efficient-pre","slug":"mask-more-and-mask-later-efficient-pre","title":"Mask More and Mask Later: Efficient Pre-training of Masked Language Models by Disentangling the [MASK] Token","date":"2022-11-09","arxiv_id":"2211.04898","n_code_links":1,"syntology":null},{"paper":null,"slug":"sentiment-analysis-of-persian-language-review","title":"Sentiment Analysis of Persian Language: Review of Algorithms, Approaches and Datasets","date":"2022-11-09","arxiv_id":"2212.06041","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-multimodal-approach-for-dementia-detection","title":"A Multimodal Approach for Dementia Detection from Spontaneous Speech with Tensor Fusion Layer","date":"2022-11-08","arxiv_id":"2211.04368","n_code_links":0,"syntology":null},{"paper":null,"slug":"discover-explanation-improvement-automatic","title":"Discover, Explanation, Improvement: An Automatic Slice Detection Framework for Natural Language Processing","date":"2022-11-08","arxiv_id":"2211.04476","n_code_links":0,"syntology":null},{"paper":null,"slug":"ad-bert-using-pre-trained-contextualized","title":"AD-BERT: Using Pre-trained contextualized embeddings to Predict the Progression from Mild Cognitive Impairment to Alzheimer's Disease","date":"2022-11-07","arxiv_id":"2212.06042","n_code_links":0,"syntology":null},{"paper":"/paper/suffix-retrieval-augmented-language-modeling","slug":"suffix-retrieval-augmented-language-modeling","title":"Suffix Retrieval-Augmented Language Modeling","date":"2022-11-06","arxiv_id":"2211.03053","n_code_links":1,"syntology":null},{"paper":null,"slug":"bert-deep-cnn-state-of-the-art-for-sentiment","title":"BERT-Deep CNN: State-of-the-Art for Sentiment Analysis of COVID-19 Tweets","date":"2022-11-04","arxiv_id":"2211.09733","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-for-long-documents-a-case-study-of","title":"BERT for Long Documents: A Case Study of Automated ICD Coding","date":"2022-11-04","arxiv_id":"2211.02519","n_code_links":0,"syntology":null},{"paper":"/paper/continuous-prompt-tuning-based-textual","slug":"continuous-prompt-tuning-based-textual","title":"Continuous Prompt Tuning Based Textual Entailment Model for E-commerce Entity Typing","date":"2022-11-04","arxiv_id":"2211.02483","n_code_links":1,"syntology":null},{"paper":"/paper/fine-tuning-language-models-via-epistemic","slug":"fine-tuning-language-models-via-epistemic","title":"Fine-Tuning Language Models via Epistemic Neural Networks","date":"2022-11-03","arxiv_id":"2211.01568","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["deepmind/neural_testbed"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"bectra-transducer-based-end-to-end-asr-with","title":"BECTRA: Transducer-based End-to-End ASR with BERT-Enhanced Encoder","date":"2022-11-02","arxiv_id":"2211.00792","n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-level-distillation-of-semantic","title":"Multi-level Distillation of Semantic Knowledge for Pre-training Multilingual Language Model","date":"2022-11-02","arxiv_id":"2211.01200","n_code_links":0,"syntology":null},{"paper":null,"slug":"investigating-content-aware-neural-text-to","title":"Investigating Content-Aware Neural Text-To-Speech MOS Prediction Using Prosodic and Linguistic Features","date":"2022-11-01","arxiv_id":"2211.00342","n_code_links":0,"syntology":null},{"paper":null,"slug":"reduce-reuse-recycle-improving-training","title":"Reduce, Reuse, Recycle: Improving Training Efficiency with Distillation","date":"2022-11-01","arxiv_id":"2211.00683","n_code_links":0,"syntology":null},{"paper":"/paper/efficient-document-retrieval-by-end-to-end","slug":"efficient-document-retrieval-by-end-to-end","title":"Efficient Document Retrieval by End-to-End Refining and Quantizing BERT Embedding with Contrastive Product Quantization","date":"2022-10-31","arxiv_id":"2210.17170","n_code_links":1,"syntology":null},{"paper":null,"slug":"leveraging-pre-trained-models-for-failure","title":"Leveraging Pre-trained Models for Failure Analysis Triplets Generation","date":"2022-10-31","arxiv_id":"2210.17497","n_code_links":0,"syntology":null},{"paper":"/paper/quala-minilm-a-quantized-length-adaptive","slug":"quala-minilm-a-quantized-length-adaptive","title":"QuaLA-MiniLM: a Quantized Length Adaptive MiniLM","date":"2022-10-31","arxiv_id":"2210.17114","n_code_links":2,"syntology":null},{"paper":null,"slug":"sdcl-self-distillation-contrastive-learning","title":"SDCL: Self-Distillation Contrastive Learning for Chinese Spell Checking","date":"2022-10-31","arxiv_id":"2210.17168","n_code_links":0,"syntology":null},{"paper":"/paper/parameter-efficient-tuning-makes-a-good","slug":"parameter-efficient-tuning-makes-a-good","title":"Parameter-Efficient Tuning Makes a Good Classification Head","date":"2022-10-30","arxiv_id":"2210.16771","n_code_links":1,"syntology":null},{"paper":null,"slug":"bert-meets-ctc-new-formulation-of-end-to-end","title":"BERT Meets CTC: New Formulation of End-to-End Speech Recognition with Pre-trained Masked Language Model","date":"2022-10-29","arxiv_id":"2210.16663","n_code_links":0,"syntology":null},{"paper":null,"slug":"empirical-evaluation-of-post-training","title":"Empirical Evaluation of Post-Training Quantization Methods for Language Tasks","date":"2022-10-29","arxiv_id":"2210.16621","n_code_links":0,"syntology":null},{"paper":"/paper/exploiting-prompt-learning-with-pre-trained","slug":"exploiting-prompt-learning-with-pre-trained","title":"Exploiting prompt learning with pre-trained language models for Alzheimer's Disease detection","date":"2022-10-29","arxiv_id":"2210.16539","n_code_links":1,"syntology":null},{"paper":"/paper/bebert-efficient-and-robust-binary-ensemble","slug":"bebert-efficient-and-robust-binary-ensemble","title":"BEBERT: Efficient and Robust Binary Ensemble BERT","date":"2022-10-28","arxiv_id":"2210.15976","n_code_links":1,"syntology":null},{"paper":null,"slug":"feature-engineering-vs-bert-on-twitter-data","title":"Feature Engineering vs BERT on Twitter Data","date":"2022-10-28","arxiv_id":"2210.16168","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-the-use-of-modality-specific-large-scale","title":"On the Use of Modality-Specific Large-Scale Pre-Trained Encoders for Multimodal Sentiment Analysis","date":"2022-10-28","arxiv_id":"2210.15937","n_code_links":0,"syntology":null},{"paper":"/paper/probing-for-targeted-syntactic-knowledge","slug":"probing-for-targeted-syntactic-knowledge","title":"Probing for targeted syntactic knowledge through grammatical error detection","date":"2022-10-28","arxiv_id":"2210.16228","n_code_links":1,"syntology":null},{"paper":null,"slug":"bert-flow-vae-a-weakly-supervised-model-for-1","title":"BERT-Flow-VAE: A Weakly-supervised Model for Multi-Label Text Classification","date":"2022-10-27","arxiv_id":"2210.15225","n_code_links":0,"syntology":null},{"paper":"/paper/coco-dr-combating-distribution-shifts-in-zero","slug":"coco-dr-combating-distribution-shifts-in-zero","title":"COCO-DR: Combating Distribution Shifts in Zero-Shot Dense Retrieval with Contrastive and Distributionally Robust Learning","date":"2022-10-27","arxiv_id":"2210.15212","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["openmatch/coco-dr"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/cost-eff-collaborative-optimization-of","slug":"cost-eff-collaborative-optimization-of","title":"COST-EFF: Collaborative Optimization of Spatial and Temporal Efficiency with Slenderized Multi-exit Language Models","date":"2022-10-27","arxiv_id":"2210.15523","n_code_links":1,"syntology":null},{"paper":"/paper/fast-distilbert-on-cpus","slug":"fast-distilbert-on-cpus","title":"Fast DistilBERT on CPUs","date":"2022-10-27","arxiv_id":"2211.07715","n_code_links":1,"syntology":null},{"paper":"/paper/fctalker-fine-and-coarse-grained-context","slug":"fctalker-fine-and-coarse-grained-context","title":"FCTalker: Fine and Coarse Grained Context Modeling for Expressive Conversational Speech Synthesis","date":"2022-10-27","arxiv_id":"2210.15360","n_code_links":1,"syntology":null},{"paper":"/paper/masked-vision-language-transformer-in-fashion","slug":"masked-vision-language-transformer-in-fashion","title":"Masked Vision-Language Transformer in Fashion","date":"2022-10-27","arxiv_id":"2210.15110","n_code_links":1,"syntology":null},{"paper":"/paper/unsupervised-boundary-aware-language-model","slug":"unsupervised-boundary-aware-language-model","title":"Unsupervised Boundary-Aware Language Model Pretraining for Chinese Sequence Labeling","date":"2022-10-27","arxiv_id":"2210.15231","n_code_links":2,"syntology":null},{"paper":"/paper/automatic-extraction-of-materials-and","slug":"automatic-extraction-of-materials-and","title":"Automatic extraction of materials and properties from superconductors scientific literature","date":"2022-10-26","arxiv_id":"2210.15600","n_code_links":2,"syntology":null},{"paper":null,"slug":"bi-link-bridging-inductive-link-predictions","title":"Bi-Link: Bridging Inductive Link Predictions from Text via Contrastive Learning of Transformers and Prompts","date":"2022-10-26","arxiv_id":"2210.14463","n_code_links":0,"syntology":null},{"paper":null,"slug":"don-t-prompt-search-mining-based-zero-shot","title":"Don't Prompt, Search! Mining-based Zero-Shot Learning with Language Models","date":"2022-10-26","arxiv_id":"2210.14803","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-robustness-of-prefix-tuning-in","title":"Exploring Robustness of Prefix Tuning in Noisy Data: A Case Study in Financial Sentiment Analysis","date":"2022-10-26","arxiv_id":"2211.05584","n_code_links":0,"syntology":null},{"paper":null,"slug":"ielm-an-open-information-extraction-benchmark","title":"IELM: An Open Information Extraction Benchmark for Pre-Trained Language Models","date":"2022-10-25","arxiv_id":"2210.14128","n_code_links":0,"syntology":null},{"paper":null,"slug":"entity-level-sentiment-analysis-in-contact","title":"Entity-level Sentiment Analysis in Contact Center Telephone Conversations","date":"2022-10-24","arxiv_id":"2210.13401","n_code_links":0,"syntology":null},{"paper":null,"slug":"explaining-translationese-why-are-neural","title":"Explaining Translationese: why are Neural Classifiers Better and what do they Learn?","date":"2022-10-24","arxiv_id":"2210.13391","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-euphemism-detection-in-few-shot-and","slug":"exploring-euphemism-detection-in-few-shot-and","title":"Exploring Euphemism Detection in Few-Shot and Zero-Shot Settings","date":"2022-10-24","arxiv_id":"2210.12926","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-better-your-syntax-the-better-your","title":"The Better Your Syntax, the Better Your Semantics? Probing Pretrained Language Models for the English Comparative Correlative","date":"2022-10-24","arxiv_id":"2210.13181","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-bert-based-deep-learning-approach-for","title":"A BERT-based Deep Learning Approach for Reputation Analysis in Social Media","date":"2022-10-23","arxiv_id":"2211.01954","n_code_links":0,"syntology":null},{"paper":null,"slug":"automated-essay-scoring-using-transformers","title":"Data Augmentation for Automated Essay Scoring using Transformer Models","date":"2022-10-23","arxiv_id":"2210.12809","n_code_links":0,"syntology":null},{"paper":null,"slug":"meta-learning-pathologies-from-radiology","title":"Meta-learning Pathologies from Radiology Reports using Variance Aware Prototypical Networks","date":"2022-10-22","arxiv_id":"2210.13979","n_code_links":0,"syntology":null},{"paper":"/paper/amos-an-adam-style-optimizer-with-adaptive","slug":"amos-an-adam-style-optimizer-with-adaptive","title":"Amos: An Adam-style Optimizer with Adaptive Weight Decay towards Model-Oriented Scale","date":"2022-10-21","arxiv_id":"2210.11693","n_code_links":1,"syntology":{"ran":14,"of":18,"n_ran_checked":14,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"14 ran (of which 0 constructed an object rather than computing a result; 14 with no instrument failure: 0 honoured, 0 violated, 14 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["google-research/jestimator"],"state":"official (archive's flag): 14 ran","n_ran":14,"n_constructed":0,"n_ran_no_instrument_failure":14,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/discovering-differences-in-the-representation","slug":"discovering-differences-in-the-representation","title":"Discovering Differences in the Representation of People using Contextualized Semantic Axes","date":"2022-10-21","arxiv_id":"2210.12170","n_code_links":1,"syntology":null},{"paper":"/paper/probing-with-noise-unpicking-the-warp-and","slug":"probing-with-noise-unpicking-the-warp-and","title":"Probing with Noise: Unpicking the Warp and Weft of Embeddings","date":"2022-10-21","arxiv_id":"2210.12206","n_code_links":1,"syntology":null},{"paper":null,"slug":"spabert-a-pretrained-language-model-from","title":"SpaBERT: A Pretrained Language Model from Geographic Data for Geo-Entity Representation","date":"2022-10-21","arxiv_id":"2210.12213","n_code_links":0,"syntology":null},{"paper":"/paper/a-unified-neural-network-model-for-1","slug":"a-unified-neural-network-model-for-1","title":"A Unified Neural Network Model for Readability Assessment with Feature Projection and Length-Balanced Loss","date":"2022-10-19","arxiv_id":"2210.10305","n_code_links":1,"syntology":null},{"paper":"/paper/biogpt-generative-pre-trained-transformer-for","slug":"biogpt-generative-pre-trained-transformer-for","title":"BioGPT: Generative Pre-trained Transformer for Biomedical Text Generation and Mining","date":"2022-10-19","arxiv_id":"2210.10341","n_code_links":4,"syntology":null},{"paper":"/paper/language-model-decomposition-quantifying-the","slug":"language-model-decomposition-quantifying-the","title":"Language Model Decomposition: Quantifying the Dependency and Correlation of Language Models","date":"2022-10-19","arxiv_id":"2210.10289","n_code_links":1,"syntology":null},{"paper":"/paper/tempo-accelerating-transformer-based-model","slug":"tempo-accelerating-transformer-based-model","title":"Tempo: Accelerating Transformer-Based Model Training through Memory Footprint Reduction","date":"2022-10-19","arxiv_id":"2210.10246","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":0,"n_instrument":1,"unverified":1,"pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["uoft-ecosystem/tempo"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["unlocated"]}}},{"paper":"/paper/elastic-numerical-reasoning-with-adaptive","slug":"elastic-numerical-reasoning-with-adaptive","title":"ELASTIC: Numerical Reasoning with Adaptive Symbolic Compiler","date":"2022-10-18","arxiv_id":"2210.10105","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 1 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["neurasearch/neurips-2022-submission-3358"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":1,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"can-bert-do-it-controller-area-network","title":"CAN-BERT do it? Controller Area Network Intrusion Detection System based on BERT Language Model","date":"2022-10-17","arxiv_id":"2210.09439","n_code_links":0,"syntology":null},{"paper":"/paper/idna-abf-multi-scale-deep-biological-language","slug":"idna-abf-multi-scale-deep-biological-language","title":"iDNA-ABF: multi-scale deep biological language learning model for the interpretable prediction of DNA methylations","date":"2022-10-17","arxiv_id":null,"n_code_links":2,"syntology":null},{"paper":"/paper/using-bottleneck-adapters-to-identify-cancer","slug":"using-bottleneck-adapters-to-identify-cancer","title":"Using Bottleneck Adapters to Identify Cancer in Clinical Notes under Low-Resource Constraints","date":"2022-10-17","arxiv_id":"2210.09440","n_code_links":1,"syntology":null},{"paper":null,"slug":"zero-shot-ranking-socio-political-texts-with","title":"Zero-Shot Ranking Socio-Political Texts with Transformer Language Models to Reduce Close Reading Time","date":"2022-10-17","arxiv_id":"2210.09179","n_code_links":0,"syntology":null},{"paper":null,"slug":"acoustic-aware-non-autoregressive-spell","title":"Acoustic-aware Non-autoregressive Spell Correction with Mask Sample Decoding","date":"2022-10-16","arxiv_id":"2210.08665","n_code_links":0,"syntology":null},{"paper":null,"slug":"ctcbert-advancing-hidden-unit-bert-with-ctc","title":"CTCBERT: Advancing Hidden-unit BERT with CTC Objectives","date":"2022-10-16","arxiv_id":"2210.08603","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-semantic-matching-through","title":"Improving Semantic Matching through Dependency-Enhanced Pre-trained Model with Adaptive Fusion","date":"2022-10-16","arxiv_id":"2210.08471","n_code_links":0,"syntology":null},{"paper":null,"slug":"aralegal-bert-a-pretrained-language-model-for","title":"AraLegal-BERT: A pretrained language model for Arabic Legal text","date":"2022-10-15","arxiv_id":"2210.08284","n_code_links":0,"syntology":null},{"paper":"/paper/dylora-parameter-efficient-tuning-of-pre","slug":"dylora-parameter-efficient-tuning-of-pre","title":"DyLoRA: Parameter Efficient Tuning of Pre-trained Models using Dynamic Search-Free Low-Rank Adaptation","date":"2022-10-14","arxiv_id":"2210.07558","n_code_links":2,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["huawei-noah/kd-nlp"],"state":"official: not harvested","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":[]}}},{"paper":"/paper/kernel-whitening-overcome-dataset-bias-with","slug":"kernel-whitening-overcome-dataset-bias-with","title":"Kernel-Whitening: Overcome Dataset Bias with Isotropic Sentence Embedding","date":"2022-10-14","arxiv_id":"2210.07547","n_code_links":2,"syntology":null},{"paper":"/paper/overlooked-video-classification-in-weakly","slug":"overlooked-video-classification-in-weakly","title":"Overlooked Video Classification in Weakly Supervised Video Anomaly Detection","date":"2022-10-13","arxiv_id":"2210.06688","n_code_links":1,"syntology":null},{"paper":null,"slug":"squat-sharpness-and-quantization-aware","title":"SQuAT: Sharpness- and Quantization-Aware Training for BERT","date":"2022-10-13","arxiv_id":"2210.07171","n_code_links":0,"syntology":null},{"paper":null,"slug":"tone-prediction-and-orthographic-conversion","title":"Tone prediction and orthographic conversion for Basaa","date":"2022-10-13","arxiv_id":"2210.06986","n_code_links":0,"syntology":null},{"paper":"/paper/foundation-transformers","slug":"foundation-transformers","title":"Foundation Transformers","date":"2022-10-12","arxiv_id":"2210.06423","n_code_links":4,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["microsoft/unilm"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"gmp-well-tuned-global-magnitude-pruning-can","title":"GMP*: Well-Tuned Gradual Magnitude Pruning Can Outperform Most BERT-Pruning Methods","date":"2022-10-12","arxiv_id":"2210.06384","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-text-style-transfer-via-style-masked","title":"On Text Style Transfer via Style Masked Language Models","date":"2022-10-12","arxiv_id":"2210.06394","n_code_links":0,"syntology":null}],"record_sha256":"2336d94c5c8b96c6a21cd89a6c2b62c2e66a9ec938df6d6c4645ed63e7f020fc","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}