{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/wordpiece/papers/32","list_of":"/method/wordpiece","method":"WordPiece","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":32,"pages_in_order":71,"rows_per_page":100,"rows":[3101,3200],"of":7063,"counts":{"archive_papers_tagged":7063,"with_a_code_link":2910,"where_syntology_ran_a_sample":650,"not_listed_spam_title":0,"listed":7063,"listed_where_code_ran":650,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":529,"every_run_a_failure_of_syntologys_instrument":121,"listed_with_a_run_with_no_instrument_failure":529,"listed_every_run_a_failure_of_syntologys_instrument":121,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/wordpiece","prev":"/method/wordpiece/papers/31","next":"/method/wordpiece/papers/33","papers":[{"paper":null,"slug":"index-indonesian-idiom-and-expression-dataset","title":"InDEX: Indonesian Idiom and Expression Dataset for Cloze Test","date":"2022-11-24","arxiv_id":"2211.13376","n_code_links":0,"syntology":null},{"paper":null,"slug":"using-selective-masking-as-a-bridge-between","title":"Using Selective Masking as a Bridge between Pre-training and Fine-tuning","date":"2022-11-24","arxiv_id":"2211.13815","n_code_links":0,"syntology":null},{"paper":"/paper/improving-visual-textual-sentiment-analysis","slug":"improving-visual-textual-sentiment-analysis","title":"Holistic Visual-Textual Sentiment Analysis with Prior Models","date":"2022-11-23","arxiv_id":"2211.12981","n_code_links":1,"syntology":null},{"paper":null,"slug":"seat-stable-and-explainable-attention","title":"SEAT: Stable and Explainable Attention","date":"2022-11-23","arxiv_id":"2211.13290","n_code_links":0,"syntology":null},{"paper":null,"slug":"word-level-representation-from-bytes-for","title":"Word-Level Representation From Bytes For Language Modeling","date":"2022-11-23","arxiv_id":"2211.12677","n_code_links":0,"syntology":null},{"paper":null,"slug":"olga-an-ontology-and-lstm-based-approach-for","title":"OLGA : An Ontology and LSTM-based approach for generating Arithmetic Word Problems (AWPs) of transfer type","date":"2022-11-22","arxiv_id":"2211.12164","n_code_links":0,"syntology":null},{"paper":"/paper/cbeaf-adapting-enhanced-continual-pretraining","slug":"cbeaf-adapting-enhanced-continual-pretraining","title":"AF Adapter: Continual Pretraining for Building Chinese Biomedical Language Model","date":"2022-11-21","arxiv_id":"2211.11363","n_code_links":1,"syntology":null},{"paper":"/paper/exploring-the-efficacy-of-pre-trained","slug":"exploring-the-efficacy-of-pre-trained","title":"Exploring the Efficacy of Pre-trained Checkpoints in Text-to-Music Generation Task","date":"2022-11-21","arxiv_id":"2211.11216","n_code_links":2,"syntology":null},{"paper":null,"slug":"l3cube-hindbert-and-devbert-pre-trained-bert","title":"L3Cube-HindBERT and DevBERT: Pre-Trained BERT Transformer models for Devanagari based Hindi and Marathi Languages","date":"2022-11-21","arxiv_id":"2211.11418","n_code_links":0,"syntology":null},{"paper":"/paper/l3cube-mahasbert-and-hindsbert-sentence-bert","slug":"l3cube-mahasbert-and-hindsbert-sentence-bert","title":"L3Cube-MahaSBERT and HindSBERT: Sentence BERT Models and Benchmarking BERT Sentence Representations for Hindi and Marathi","date":"2022-11-21","arxiv_id":"2211.11187","n_code_links":1,"syntology":null},{"paper":null,"slug":"tcbert-a-technical-report-for-chinese-topic","title":"TCBERT: A Technical Report for Chinese Topic Classification BERT","date":"2022-11-21","arxiv_id":"2211.11304","n_code_links":0,"syntology":null},{"paper":null,"slug":"conceptor-aided-debiasing-of-contextualized","title":"Conceptor-Aided Debiasing of Large Language Models","date":"2022-11-20","arxiv_id":"2211.11087","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-conspiracy-theory-against-covid-19","title":"Detecting Conspiracy Theory Against COVID-19 Vaccines","date":"2022-11-20","arxiv_id":"2211.13003","n_code_links":0,"syntology":null},{"paper":null,"slug":"feature-weaken-vicinal-data-augmentation-for","title":"Feature Weaken: Vicinal Data Augmentation for Classification","date":"2022-11-20","arxiv_id":"2211.10944","n_code_links":0,"syntology":null},{"paper":"/paper/understanding-and-improving-knowledge-1","slug":"understanding-and-improving-knowledge-1","title":"Understanding and Improving Knowledge Distillation for Quantization-Aware Training of Large Transformer Encoders","date":"2022-11-20","arxiv_id":"2211.11014","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-survey-on-knowledge-enhanced-multimodal","title":"A survey on knowledge-enhanced multimodal learning","date":"2022-11-19","arxiv_id":"2211.12328","n_code_links":0,"syntology":null},{"paper":null,"slug":"entity-assisted-language-models-for","title":"Entity-Assisted Language Models for Identifying Check-worthy Sentences","date":"2022-11-19","arxiv_id":"2211.10678","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-users-social-network-embeddings","title":"Leveraging Users' Social Network Embeddings for Fake News Detection on Twitter","date":"2022-11-19","arxiv_id":"2211.10672","n_code_links":0,"syntology":null},{"paper":null,"slug":"metadata-might-make-language-models-better","title":"Metadata Might Make Language Models Better","date":"2022-11-18","arxiv_id":"2211.10086","n_code_links":0,"syntology":null},{"paper":null,"slug":"where-did-you-tweet-from-inferring-the-origin","title":"Where did you tweet from? Inferring the origin locations of tweets based on contextual information","date":"2022-11-18","arxiv_id":"2211.16506","n_code_links":0,"syntology":null},{"paper":null,"slug":"longfnt-long-form-speech-recognition-with","title":"LongFNT: Long-form Speech Recognition with Factorized Neural Transducer","date":"2022-11-17","arxiv_id":"2211.09412","n_code_links":0,"syntology":null},{"paper":"/paper/protsi-prototypical-siamese-network-with-data","slug":"protsi-prototypical-siamese-network-with-data","title":"ProtSi: Prototypical Siamese Network with Data Augmentation for Few-Shot Subjective Answer Evaluation","date":"2022-11-17","arxiv_id":"2211.09855","n_code_links":1,"syntology":null},{"paper":"/paper/random-ltd-random-and-layerwise-token","slug":"random-ltd-random-and-layerwise-token","title":"Random-LTD: Random and Layerwise Token Dropping Brings Efficient Training for Large-scale Transformers","date":"2022-11-17","arxiv_id":"2211.11586","n_code_links":1,"syntology":null},{"paper":null,"slug":"fast-and-accurate-fsa-system-using-elbert-an","title":"Fast and Accurate FSA System Using ELBERT: An Efficient and Lightweight BERT","date":"2022-11-16","arxiv_id":"2211.08842","n_code_links":0,"syntology":null},{"paper":"/paper/an-fnet-based-auto-encoder-for-long-sequence","slug":"an-fnet-based-auto-encoder-for-long-sequence","title":"An FNet based Auto Encoder for Long Sequence News Story Generation","date":"2022-11-15","arxiv_id":"2211.08295","n_code_links":1,"syntology":null},{"paper":null,"slug":"empowering-language-models-with-knowledge","title":"Empowering Language Models with Knowledge Graph Reasoning for Question Answering","date":"2022-11-15","arxiv_id":"2211.08380","n_code_links":0,"syntology":null},{"paper":null,"slug":"robbert-2022-updating-a-dutch-language-model","title":"RobBERT-2022: Updating a Dutch Language Model to Account for Evolving Language Use","date":"2022-11-15","arxiv_id":"2211.08192","n_code_links":0,"syntology":null},{"paper":"/paper/greenplm-cross-lingual-pre-trained-language","slug":"greenplm-cross-lingual-pre-trained-language","title":"GreenPLM: Cross-Lingual Transfer of Monolingual Pre-Trained Language Models at Almost No Cost","date":"2022-11-13","arxiv_id":"2211.06993","n_code_links":1,"syntology":null},{"paper":"/paper/xu-at-semeval-2022-task-4-pre-bert-neural-1","slug":"xu-at-semeval-2022-task-4-pre-bert-neural-1","title":"Xu at SemEval-2022 Task 4: Pre-BERT Neural Network Methods vs Post-BERT RoBERTa Approach for Patronizing and Condescending Language Detection","date":"2022-11-13","arxiv_id":"2211.06874","n_code_links":1,"syntology":null},{"paper":"/paper/dark-patterns-in-e-commerce-a-dataset-and-its","slug":"dark-patterns-in-e-commerce-a-dataset-and-its","title":"Dark patterns in e-commerce: a dataset and its baseline evaluations","date":"2022-11-12","arxiv_id":"2211.06543","n_code_links":1,"syntology":null},{"paper":"/paper/misinformation-detection-using-persuasive","slug":"misinformation-detection-using-persuasive","title":"Using Persuasive Writing Strategies to Explain and Detect Health Misinformation","date":"2022-11-11","arxiv_id":"2211.05985","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-architectural-bottleneck-principle","title":"The Architectural Bottleneck Principle","date":"2022-11-11","arxiv_id":"2211.06420","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-based-combination-of-convolutional-and","title":"BERT-Based Combination of Convolutional and Recurrent Neural Network for Indonesian Sentiment Analysis","date":"2022-11-10","arxiv_id":"2211.05273","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-in-plutarch-s-shadows","title":"BERT in Plutarch's Shadows","date":"2022-11-10","arxiv_id":"2211.05673","n_code_links":0,"syntology":null},{"paper":null,"slug":"biomedical-multi-hop-question-answering-using","title":"Biomedical Multi-hop Question Answering Using Knowledge Graph Embeddings and Language Models","date":"2022-11-10","arxiv_id":"2211.05351","n_code_links":0,"syntology":null},{"paper":"/paper/cherry-hypothesis-identifying-the-cherry-on","slug":"cherry-hypothesis-identifying-the-cherry-on","title":"PAD-Net: An Efficient Framework for Dynamic Networks","date":"2022-11-10","arxiv_id":"2211.05528","n_code_links":1,"syntology":null},{"paper":null,"slug":"syntax-guided-domain-adaptation-for-aspect","title":"Syntax-Guided Domain Adaptation for Aspect-based Sentiment Analysis","date":"2022-11-10","arxiv_id":"2211.05457","n_code_links":0,"syntology":null},{"paper":"/paper/collateral-facilitation-in-humans-and","slug":"collateral-facilitation-in-humans-and","title":"Collateral facilitation in humans and language models","date":"2022-11-09","arxiv_id":"2211.05198","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["jmichaelov/collateral-facilitation"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"cross-lingual-transfer-learning-for-check","title":"Cross-lingual Transfer Learning for Check-worthy Claim Identification over Twitter","date":"2022-11-09","arxiv_id":"2211.05087","n_code_links":0,"syntology":null},{"paper":"/paper/mask-more-and-mask-later-efficient-pre","slug":"mask-more-and-mask-later-efficient-pre","title":"Mask More and Mask Later: Efficient Pre-training of Masked Language Models by Disentangling the [MASK] Token","date":"2022-11-09","arxiv_id":"2211.04898","n_code_links":1,"syntology":null},{"paper":null,"slug":"sentiment-analysis-of-persian-language-review","title":"Sentiment Analysis of Persian Language: Review of Algorithms, Approaches and Datasets","date":"2022-11-09","arxiv_id":"2212.06041","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-multimodal-approach-for-dementia-detection","title":"A Multimodal Approach for Dementia Detection from Spontaneous Speech with Tensor Fusion Layer","date":"2022-11-08","arxiv_id":"2211.04368","n_code_links":0,"syntology":null},{"paper":null,"slug":"discover-explanation-improvement-automatic","title":"Discover, Explanation, Improvement: An Automatic Slice Detection Framework for Natural Language Processing","date":"2022-11-08","arxiv_id":"2211.04476","n_code_links":0,"syntology":null},{"paper":null,"slug":"ad-bert-using-pre-trained-contextualized","title":"AD-BERT: Using Pre-trained contextualized embeddings to Predict the Progression from Mild Cognitive Impairment to Alzheimer's Disease","date":"2022-11-07","arxiv_id":"2212.06042","n_code_links":0,"syntology":null},{"paper":"/paper/suffix-retrieval-augmented-language-modeling","slug":"suffix-retrieval-augmented-language-modeling","title":"Suffix Retrieval-Augmented Language Modeling","date":"2022-11-06","arxiv_id":"2211.03053","n_code_links":1,"syntology":null},{"paper":null,"slug":"bert-deep-cnn-state-of-the-art-for-sentiment","title":"BERT-Deep CNN: State-of-the-Art for Sentiment Analysis of COVID-19 Tweets","date":"2022-11-04","arxiv_id":"2211.09733","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-for-long-documents-a-case-study-of","title":"BERT for Long Documents: A Case Study of Automated ICD Coding","date":"2022-11-04","arxiv_id":"2211.02519","n_code_links":0,"syntology":null},{"paper":"/paper/continuous-prompt-tuning-based-textual","slug":"continuous-prompt-tuning-based-textual","title":"Continuous Prompt Tuning Based Textual Entailment Model for E-commerce Entity Typing","date":"2022-11-04","arxiv_id":"2211.02483","n_code_links":1,"syntology":null},{"paper":"/paper/fine-tuning-language-models-via-epistemic","slug":"fine-tuning-language-models-via-epistemic","title":"Fine-Tuning Language Models via Epistemic Neural Networks","date":"2022-11-03","arxiv_id":"2211.01568","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["deepmind/neural_testbed"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"bectra-transducer-based-end-to-end-asr-with","title":"BECTRA: Transducer-based End-to-End ASR with BERT-Enhanced Encoder","date":"2022-11-02","arxiv_id":"2211.00792","n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-level-distillation-of-semantic","title":"Multi-level Distillation of Semantic Knowledge for Pre-training Multilingual Language Model","date":"2022-11-02","arxiv_id":"2211.01200","n_code_links":0,"syntology":null},{"paper":null,"slug":"processing-long-legal-documents-with-pre","title":"Processing Long Legal Documents with Pre-trained Transformers: Modding LegalBERT and Longformer","date":"2022-11-02","arxiv_id":"2211.00974","n_code_links":0,"syntology":null},{"paper":"/paper/classactionprediction-a-challenging-benchmark","slug":"classactionprediction-a-challenging-benchmark","title":"ClassActionPrediction: A Challenging Benchmark for Legal Judgment Prediction of Class Action Cases in the US","date":"2022-11-01","arxiv_id":"2211.00582","n_code_links":1,"syntology":null},{"paper":null,"slug":"investigating-content-aware-neural-text-to","title":"Investigating Content-Aware Neural Text-To-Speech MOS Prediction Using Prosodic and Linguistic Features","date":"2022-11-01","arxiv_id":"2211.00342","n_code_links":0,"syntology":null},{"paper":null,"slug":"reduce-reuse-recycle-improving-training","title":"Reduce, Reuse, Recycle: Improving Training Efficiency with Distillation","date":"2022-11-01","arxiv_id":"2211.00683","n_code_links":0,"syntology":null},{"paper":"/paper/efficient-document-retrieval-by-end-to-end","slug":"efficient-document-retrieval-by-end-to-end","title":"Efficient Document Retrieval by End-to-End Refining and Quantizing BERT Embedding with Contrastive Product Quantization","date":"2022-10-31","arxiv_id":"2210.17170","n_code_links":1,"syntology":null},{"paper":null,"slug":"leveraging-pre-trained-models-for-failure","title":"Leveraging Pre-trained Models for Failure Analysis Triplets Generation","date":"2022-10-31","arxiv_id":"2210.17497","n_code_links":0,"syntology":null},{"paper":"/paper/quala-minilm-a-quantized-length-adaptive","slug":"quala-minilm-a-quantized-length-adaptive","title":"QuaLA-MiniLM: a Quantized Length Adaptive MiniLM","date":"2022-10-31","arxiv_id":"2210.17114","n_code_links":2,"syntology":null},{"paper":null,"slug":"sdcl-self-distillation-contrastive-learning","title":"SDCL: Self-Distillation Contrastive Learning for Chinese Spell Checking","date":"2022-10-31","arxiv_id":"2210.17168","n_code_links":0,"syntology":null},{"paper":"/paper/parameter-efficient-tuning-makes-a-good","slug":"parameter-efficient-tuning-makes-a-good","title":"Parameter-Efficient Tuning Makes a Good Classification Head","date":"2022-10-30","arxiv_id":"2210.16771","n_code_links":1,"syntology":null},{"paper":null,"slug":"bert-meets-ctc-new-formulation-of-end-to-end","title":"BERT Meets CTC: New Formulation of End-to-End Speech Recognition with Pre-trained Masked Language Model","date":"2022-10-29","arxiv_id":"2210.16663","n_code_links":0,"syntology":null},{"paper":null,"slug":"empirical-evaluation-of-post-training","title":"Empirical Evaluation of Post-Training Quantization Methods for Language Tasks","date":"2022-10-29","arxiv_id":"2210.16621","n_code_links":0,"syntology":null},{"paper":"/paper/exploiting-prompt-learning-with-pre-trained","slug":"exploiting-prompt-learning-with-pre-trained","title":"Exploiting prompt learning with pre-trained language models for Alzheimer's Disease detection","date":"2022-10-29","arxiv_id":"2210.16539","n_code_links":1,"syntology":null},{"paper":"/paper/bebert-efficient-and-robust-binary-ensemble","slug":"bebert-efficient-and-robust-binary-ensemble","title":"BEBERT: Efficient and Robust Binary Ensemble BERT","date":"2022-10-28","arxiv_id":"2210.15976","n_code_links":1,"syntology":null},{"paper":null,"slug":"feature-engineering-vs-bert-on-twitter-data","title":"Feature Engineering vs BERT on Twitter Data","date":"2022-10-28","arxiv_id":"2210.16168","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-the-use-of-modality-specific-large-scale","title":"On the Use of Modality-Specific Large-Scale Pre-Trained Encoders for Multimodal Sentiment Analysis","date":"2022-10-28","arxiv_id":"2210.15937","n_code_links":0,"syntology":null},{"paper":"/paper/probing-for-targeted-syntactic-knowledge","slug":"probing-for-targeted-syntactic-knowledge","title":"Probing for targeted syntactic knowledge through grammatical error detection","date":"2022-10-28","arxiv_id":"2210.16228","n_code_links":1,"syntology":null},{"paper":null,"slug":"bert-flow-vae-a-weakly-supervised-model-for-1","title":"BERT-Flow-VAE: A Weakly-supervised Model for Multi-Label Text Classification","date":"2022-10-27","arxiv_id":"2210.15225","n_code_links":0,"syntology":null},{"paper":"/paper/coco-dr-combating-distribution-shifts-in-zero","slug":"coco-dr-combating-distribution-shifts-in-zero","title":"COCO-DR: Combating Distribution Shifts in Zero-Shot Dense Retrieval with Contrastive and Distributionally Robust Learning","date":"2022-10-27","arxiv_id":"2210.15212","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["openmatch/coco-dr"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/cost-eff-collaborative-optimization-of","slug":"cost-eff-collaborative-optimization-of","title":"COST-EFF: Collaborative Optimization of Spatial and Temporal Efficiency with Slenderized Multi-exit Language Models","date":"2022-10-27","arxiv_id":"2210.15523","n_code_links":1,"syntology":null},{"paper":"/paper/fast-distilbert-on-cpus","slug":"fast-distilbert-on-cpus","title":"Fast DistilBERT on CPUs","date":"2022-10-27","arxiv_id":"2211.07715","n_code_links":1,"syntology":null},{"paper":"/paper/fctalker-fine-and-coarse-grained-context","slug":"fctalker-fine-and-coarse-grained-context","title":"FCTalker: Fine and Coarse Grained Context Modeling for Expressive Conversational Speech Synthesis","date":"2022-10-27","arxiv_id":"2210.15360","n_code_links":1,"syntology":null},{"paper":"/paper/masked-vision-language-transformer-in-fashion","slug":"masked-vision-language-transformer-in-fashion","title":"Masked Vision-Language Transformer in Fashion","date":"2022-10-27","arxiv_id":"2210.15110","n_code_links":1,"syntology":null},{"paper":"/paper/unsupervised-boundary-aware-language-model","slug":"unsupervised-boundary-aware-language-model","title":"Unsupervised Boundary-Aware Language Model Pretraining for Chinese Sequence Labeling","date":"2022-10-27","arxiv_id":"2210.15231","n_code_links":2,"syntology":null},{"paper":"/paper/automatic-extraction-of-materials-and","slug":"automatic-extraction-of-materials-and","title":"Automatic extraction of materials and properties from superconductors scientific literature","date":"2022-10-26","arxiv_id":"2210.15600","n_code_links":2,"syntology":null},{"paper":null,"slug":"beyond-english-centric-bitexts-for-better","title":"Beyond English-Centric Bitexts for Better Multilingual Language Representation Learning","date":"2022-10-26","arxiv_id":"2210.14867","n_code_links":0,"syntology":null},{"paper":null,"slug":"bi-link-bridging-inductive-link-predictions","title":"Bi-Link: Bridging Inductive Link Predictions from Text via Contrastive Learning of Transformers and Prompts","date":"2022-10-26","arxiv_id":"2210.14463","n_code_links":0,"syntology":null},{"paper":null,"slug":"don-t-prompt-search-mining-based-zero-shot","title":"Don't Prompt, Search! Mining-based Zero-Shot Learning with Language Models","date":"2022-10-26","arxiv_id":"2210.14803","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-robustness-of-prefix-tuning-in","title":"Exploring Robustness of Prefix Tuning in Noisy Data: A Case Study in Financial Sentiment Analysis","date":"2022-10-26","arxiv_id":"2211.05584","n_code_links":0,"syntology":null},{"paper":null,"slug":"interpreting-chemical-words-of-a-data-driven","title":"Exploring Data-Driven Chemical SMILES Tokenization Approaches to Identify Key Protein-Ligand Binding Moieties","date":"2022-10-26","arxiv_id":"2210.14642","n_code_links":0,"syntology":null},{"paper":"/paper/how-long-is-enough-exploring-the-optimal","slug":"how-long-is-enough-exploring-the-optimal","title":"How Long Is Enough? Exploring the Optimal Intervals of Long-Range Clinical Note Language Modeling","date":"2022-10-25","arxiv_id":"2211.07713","n_code_links":1,"syntology":null},{"paper":null,"slug":"ielm-an-open-information-extraction-benchmark","title":"IELM: An Open Information Extraction Benchmark for Pre-Trained Language Models","date":"2022-10-25","arxiv_id":"2210.14128","n_code_links":0,"syntology":null},{"paper":null,"slug":"effective-pre-training-objectives-for","title":"Effective Pre-Training Objectives for Transformer-based Autoencoders","date":"2022-10-24","arxiv_id":"2210.13536","n_code_links":0,"syntology":null},{"paper":null,"slug":"entity-level-sentiment-analysis-in-contact","title":"Entity-level Sentiment Analysis in Contact Center Telephone Conversations","date":"2022-10-24","arxiv_id":"2210.13401","n_code_links":0,"syntology":null},{"paper":null,"slug":"explaining-translationese-why-are-neural","title":"Explaining Translationese: why are Neural Classifiers Better and what do they Learn?","date":"2022-10-24","arxiv_id":"2210.13391","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-euphemism-detection-in-few-shot-and","slug":"exploring-euphemism-detection-in-few-shot-and","title":"Exploring Euphemism Detection in Few-Shot and Zero-Shot Settings","date":"2022-10-24","arxiv_id":"2210.12926","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-better-your-syntax-the-better-your","title":"The Better Your Syntax, the Better Your Semantics? Probing Pretrained Language Models for the English Comparative Correlative","date":"2022-10-24","arxiv_id":"2210.13181","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-bert-based-deep-learning-approach-for","title":"A BERT-based Deep Learning Approach for Reputation Analysis in Social Media","date":"2022-10-23","arxiv_id":"2211.01954","n_code_links":0,"syntology":null},{"paper":null,"slug":"automated-essay-scoring-using-transformers","title":"Data Augmentation for Automated Essay Scoring using Transformer Models","date":"2022-10-23","arxiv_id":"2210.12809","n_code_links":0,"syntology":null},{"paper":null,"slug":"discriminative-language-model-as-semantic","title":"Discriminative Language Model as Semantic Consistency Scorer for Prompt-based Few-Shot Text Classification","date":"2022-10-23","arxiv_id":"2210.12763","n_code_links":0,"syntology":null},{"paper":null,"slug":"meta-learning-pathologies-from-radiology","title":"Meta-learning Pathologies from Radiology Reports using Variance Aware Prototypical Networks","date":"2022-10-22","arxiv_id":"2210.13979","n_code_links":0,"syntology":null},{"paper":"/paper/amos-an-adam-style-optimizer-with-adaptive","slug":"amos-an-adam-style-optimizer-with-adaptive","title":"Amos: An Adam-style Optimizer with Adaptive Weight Decay towards Model-Oriented Scale","date":"2022-10-21","arxiv_id":"2210.11693","n_code_links":1,"syntology":{"ran":14,"of":18,"n_ran_checked":14,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"14 ran (of which 0 constructed an object rather than computing a result; 14 with no instrument failure: 0 honoured, 0 violated, 14 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["google-research/jestimator"],"state":"official (archive's flag): 14 ran","n_ran":14,"n_constructed":0,"n_ran_no_instrument_failure":14,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/discovering-differences-in-the-representation","slug":"discovering-differences-in-the-representation","title":"Discovering Differences in the Representation of People using Contextualized Semantic Axes","date":"2022-10-21","arxiv_id":"2210.12170","n_code_links":1,"syntology":null},{"paper":null,"slug":"littlebird-efficient-faster-longer","title":"LittleBird: Efficient Faster & Longer Transformer for Question Answering","date":"2022-10-21","arxiv_id":"2210.11870","n_code_links":0,"syntology":null},{"paper":"/paper/probing-with-noise-unpicking-the-warp-and","slug":"probing-with-noise-unpicking-the-warp-and","title":"Probing with Noise: Unpicking the Warp and Weft of Embeddings","date":"2022-10-21","arxiv_id":"2210.12206","n_code_links":1,"syntology":null},{"paper":null,"slug":"spabert-a-pretrained-language-model-from","title":"SpaBERT: A Pretrained Language Model from Geographic Data for Geo-Entity Representation","date":"2022-10-21","arxiv_id":"2210.12213","n_code_links":0,"syntology":null},{"paper":"/paper/a-unified-neural-network-model-for-1","slug":"a-unified-neural-network-model-for-1","title":"A Unified Neural Network Model for Readability Assessment with Feature Projection and Length-Balanced Loss","date":"2022-10-19","arxiv_id":"2210.10305","n_code_links":1,"syntology":null},{"paper":"/paper/biogpt-generative-pre-trained-transformer-for","slug":"biogpt-generative-pre-trained-transformer-for","title":"BioGPT: Generative Pre-trained Transformer for Biomedical Text Generation and Mining","date":"2022-10-19","arxiv_id":"2210.10341","n_code_links":4,"syntology":null},{"paper":"/paper/language-model-decomposition-quantifying-the","slug":"language-model-decomposition-quantifying-the","title":"Language Model Decomposition: Quantifying the Dependency and Correlation of Language Models","date":"2022-10-19","arxiv_id":"2210.10289","n_code_links":1,"syntology":null},{"paper":"/paper/tempo-accelerating-transformer-based-model","slug":"tempo-accelerating-transformer-based-model","title":"Tempo: Accelerating Transformer-Based Model Training through Memory Footprint Reduction","date":"2022-10-19","arxiv_id":"2210.10246","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":0,"n_instrument":1,"unverified":1,"pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["uoft-ecosystem/tempo"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["unlocated"]}}}],"record_sha256":"5f7d875046d9a48b99f94354f9cf92b3a4b8e5b0818bc52d97ffa8191f27f4ec","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}