{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/bert/papers/28","list_of":"/method/bert","method":"BERT","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":28,"pages_in_order":70,"rows_per_page":100,"rows":[2701,2800],"of":6938,"counts":{"archive_papers_tagged":6938,"with_a_code_link":2862,"where_syntology_ran_a_sample":640,"not_listed_spam_title":0,"listed":6938,"listed_where_code_ran":640,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":520,"every_run_a_failure_of_syntologys_instrument":120,"listed_with_a_run_with_no_instrument_failure":520,"listed_every_run_a_failure_of_syntologys_instrument":120,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/bert","prev":"/method/bert/papers/27","next":"/method/bert/papers/29","papers":[{"paper":null,"slug":"research-on-multilingual-news-clustering","title":"Research on Multilingual News Clustering Based on Cross-Language Word Embeddings","date":"2023-05-30","arxiv_id":"2305.18880","n_code_links":0,"syntology":null},{"paper":"/paper/scone-benchmarking-negation-reasoning-in","slug":"scone-benchmarking-negation-reasoning-in","title":"ScoNe: Benchmarking Negation Reasoning in Language Models With Fine-Tuning and In-Context Learning","date":"2023-05-30","arxiv_id":"2305.19426","n_code_links":1,"syntology":null},{"paper":null,"slug":"abstractive-summarization-as-augmentation-for","title":"Abstractive Summarization as Augmentation for Document-Level Event Detection","date":"2023-05-29","arxiv_id":"2305.18023","n_code_links":0,"syntology":null},{"paper":"/paper/from-adversarial-arms-race-to-model-centric","slug":"from-adversarial-arms-race-to-model-centric","title":"From Adversarial Arms Race to Model-centric Evaluation: Motivating a Unified Automatic Robustness Evaluation Framework","date":"2023-05-29","arxiv_id":"2305.18503","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["thunlp/robtest"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"slimfit-memory-efficient-fine-tuning-of","title":"SlimFit: Memory-Efficient Fine-Tuning of Transformer-based Models Using Training Dynamics","date":"2023-05-29","arxiv_id":"2305.18513","n_code_links":0,"syntology":null},{"paper":null,"slug":"transformer-language-models-handle-word","title":"Transformer Language Models Handle Word Frequency in Prediction Head","date":"2023-05-29","arxiv_id":"2305.18294","n_code_links":0,"syntology":null},{"paper":"/paper/rethinking-masked-language-modeling-for","slug":"rethinking-masked-language-modeling-for","title":"Rethinking Masked Language Modeling for Chinese Spelling Correction","date":"2023-05-28","arxiv_id":"2305.17721","n_code_links":1,"syntology":null},{"paper":null,"slug":"transfer-learning-for-power-outage-detection","title":"Transfer Learning for Power Outage Detection Task with Limited Training Data","date":"2023-05-28","arxiv_id":"2305.17817","n_code_links":0,"syntology":null},{"paper":"/paper/an-investigation-into-the-effects-of-pre","slug":"an-investigation-into-the-effects-of-pre","title":"Diagnosing Transformers: Illuminating Feature Spaces for Clinical Decision-Making","date":"2023-05-27","arxiv_id":"2305.17588","n_code_links":1,"syntology":null},{"paper":null,"slug":"complementary-and-integrative-health-lexicon","title":"Complementary and Integrative Health Lexicon (CIHLex) and Entity Recognition in the Literature","date":"2023-05-27","arxiv_id":"2305.17353","n_code_links":0,"syntology":null},{"paper":"/paper/modeling-adversarial-attack-on-pre-trained","slug":"modeling-adversarial-attack-on-pre-trained","title":"Modeling Adversarial Attack on Pre-trained Language Models as Sequential Decision Making","date":"2023-05-27","arxiv_id":"2305.17440","n_code_links":1,"syntology":null},{"paper":null,"slug":"calibration-of-transformer-based-models-for","title":"Calibration of Transformer-based Models for Identifying Stress and Depression in Social Media","date":"2023-05-26","arxiv_id":"2305.16797","n_code_links":0,"syntology":null},{"paper":"/paper/geovln-learning-geometry-enhanced-visual-1","slug":"geovln-learning-geometry-enhanced-visual-1","title":"GeoVLN: Learning Geometry-Enhanced Visual Representation with Slot Attention for Vision-and-Language Navigation","date":"2023-05-26","arxiv_id":"2305.17102","n_code_links":1,"syntology":null},{"paper":null,"slug":"knse-a-knowledge-aware-natural-language","title":"KNSE: A Knowledge-aware Natural Language Inference Framework for Dialogue Symptom Status Recognition","date":"2023-05-26","arxiv_id":"2305.16833","n_code_links":0,"syntology":null},{"paper":null,"slug":"theoretical-and-practical-perspectives-on","title":"Theoretical and Practical Perspectives on what Influence Functions Do","date":"2023-05-26","arxiv_id":"2305.16971","n_code_links":0,"syntology":null},{"paper":"/paper/zero-is-not-hero-yet-benchmarking-zero-shot","slug":"zero-is-not-hero-yet-benchmarking-zero-shot","title":"Zero is Not Hero Yet: Benchmarking Zero-Shot Performance of LLMs for Financial Tasks","date":"2023-05-26","arxiv_id":"2305.16633","n_code_links":1,"syntology":null},{"paper":null,"slug":"comparative-study-of-pre-trained-bert-models","title":"Comparative Study of Pre-Trained BERT Models for Code-Mixed Hindi-English Data","date":"2023-05-25","arxiv_id":"2305.15722","n_code_links":0,"syntology":null},{"paper":null,"slug":"context-aware-attention-layers-coupled-with","title":"Context-aware attention layers coupled with optimal transport domain adaptation and multimodal fusion methods for recognizing dementia from spontaneous speech","date":"2023-05-25","arxiv_id":"2305.16406","n_code_links":0,"syntology":null},{"paper":null,"slug":"not-wacky-vs-definitely-wacky-a-study-of","title":"Not wacky vs. definitely wacky: A study of scalar adverbs in pretrained language models","date":"2023-05-25","arxiv_id":"2305.16426","n_code_links":0,"syntology":null},{"paper":"/paper/text-to-motion-retrieval-towards-joint","slug":"text-to-motion-retrieval-towards-joint","title":"Text-to-Motion Retrieval: Towards Joint Understanding of Human Motion Data and Natural Language","date":"2023-05-25","arxiv_id":"2305.15842","n_code_links":1,"syntology":null},{"paper":"/paper/a-causal-view-of-entity-bias-in-large","slug":"a-causal-view-of-entity-bias-in-large","title":"A Causal View of Entity Bias in (Large) Language Models","date":"2023-05-24","arxiv_id":"2305.14695","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["luka-group/causal-view-of-entity-bias"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/complex-mathematical-symbol-definition","slug":"complex-mathematical-symbol-definition","title":"Complex Mathematical Symbol Definition Structures: A Dataset and Model for Coordination Resolution in Definition Extraction","date":"2023-05-24","arxiv_id":"2305.14660","n_code_links":1,"syntology":null},{"paper":"/paper/context-aware-transformer-pre-training-for","slug":"context-aware-transformer-pre-training-for","title":"Context-Aware Transformer Pre-Training for Answer Sentence Selection","date":"2023-05-24","arxiv_id":"2305.15358","n_code_links":0,"syntology":null},{"paper":null,"slug":"dynamic-masking-rate-schedules-for-mlm","title":"Dynamic Masking Rate Schedules for MLM Pretraining","date":"2023-05-24","arxiv_id":"2305.15096","n_code_links":0,"syntology":null},{"paper":null,"slug":"extracting-psychological-indicators-using","title":"Extracting Psychological Indicators Using Question Answering","date":"2023-05-24","arxiv_id":"2305.14891","n_code_links":0,"syntology":null},{"paper":"/paper/ghostbuster-detecting-text-ghostwritten-by","slug":"ghostbuster-detecting-text-ghostwritten-by","title":"Ghostbuster: Detecting Text Ghostwritten by Large Language Models","date":"2023-05-24","arxiv_id":"2305.15047","n_code_links":2,"syntology":null},{"paper":"/paper/how-to-distill-your-bert-an-empirical-study","slug":"how-to-distill-your-bert-an-empirical-study","title":"How to Distill your BERT: An Empirical Study on the Impact of Weight Initialisation and Distillation Objectives","date":"2023-05-24","arxiv_id":"2305.15032","n_code_links":1,"syntology":{"ran":6,"of":11,"n_ran_checked":5,"n_instrument":1,"unverified":5,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 5 unverified","official":{"repos":["mainlp/how-to-distill-your-bert"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":"/paper/revisiting-token-dropping-strategy-in","slug":"revisiting-token-dropping-strategy-in","title":"Revisiting Token Dropping Strategy in Efficient BERT Pretraining","date":"2023-05-24","arxiv_id":"2305.15273","n_code_links":1,"syntology":null},{"paper":null,"slug":"2305-14521","title":"Few-shot Adaptation to Distribution Shifts By Mixing Source and Target Embeddings","date":"2023-05-23","arxiv_id":"2305.14521","n_code_links":0,"syntology":null},{"paper":"/paper/a-simple-method-for-unsupervised-bilingual","slug":"a-simple-method-for-unsupervised-bilingual","title":"When your Cousin has the Right Connections: Unsupervised Bilingual Lexicon Induction for Related Data-Imbalanced Languages","date":"2023-05-23","arxiv_id":"2305.14012","n_code_links":1,"syntology":null},{"paper":"/paper/all-roads-lead-to-rome-exploring-the","slug":"all-roads-lead-to-rome-exploring-the","title":"All Roads Lead to Rome? Exploring the Invariance of Transformers' Representations","date":"2023-05-23","arxiv_id":"2305.14555","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["twinkle0331/bert-similarity"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"assessing-linguistic-generalisation-in","title":"Assessing Linguistic Generalisation in Language Models: A Dataset for Brazilian Portuguese","date":"2023-05-23","arxiv_id":"2305.14070","n_code_links":0,"syntology":null},{"paper":"/paper/axomiyaberta-a-phonologically-aware","slug":"axomiyaberta-a-phonologically-aware","title":"AxomiyaBERTa: A Phonologically-aware Transformer Model for Assamese","date":"2023-05-23","arxiv_id":"2305.13641","n_code_links":1,"syntology":null},{"paper":"/paper/connecting-the-dots-what-graph-based-text","slug":"connecting-the-dots-what-graph-based-text","title":"Connecting the Dots: What Graph-Based Text Representations Work Best for Text Classification Using Graph Neural Networks?","date":"2023-05-23","arxiv_id":"2305.14578","n_code_links":1,"syntology":null},{"paper":"/paper/exploring-large-language-models-for-classical","slug":"exploring-large-language-models-for-classical","title":"Exploring Large Language Models for Classical Philology","date":"2023-05-23","arxiv_id":"2305.13698","n_code_links":1,"syntology":null},{"paper":null,"slug":"handling-realistic-label-noise-in-bert-text","title":"Handling Realistic Label Noise in BERT Text Classification","date":"2023-05-23","arxiv_id":"2305.16337","n_code_links":0,"syntology":null},{"paper":"/paper/on-robustness-of-finetuned-transformer-based","slug":"on-robustness-of-finetuned-transformer-based","title":"On Robustness of Finetuned Transformer-based NLP Models","date":"2023-05-23","arxiv_id":"2305.14453","n_code_links":1,"syntology":null},{"paper":null,"slug":"training-transitive-and-commutative","title":"Training Transitive and Commutative Multimodal Transformers with LoReTTa","date":"2023-05-23","arxiv_id":"2305.14243","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-energy-based-language-models-with","slug":"exploring-energy-based-language-models-with","title":"Exploring Energy-based Language Models with Different Architectures and Training Methods for Speech Recognition","date":"2023-05-22","arxiv_id":"2305.12676","n_code_links":2,"syntology":null},{"paper":null,"slug":"gatology-for-linguistics-what-syntactic","title":"GATology for Linguistics: What Syntactic Dependencies It Knows","date":"2023-05-22","arxiv_id":"2305.13403","n_code_links":0,"syntology":null},{"paper":null,"slug":"imsimcse-improving-contrastive-learning-for","title":"SimCSE++: Improving Contrastive Learning for Sentence Embeddings from Two Perspectives","date":"2023-05-22","arxiv_id":"2305.13192","n_code_links":0,"syntology":null},{"paper":"/paper/language-agnostic-bias-detection-in-language","slug":"language-agnostic-bias-detection-in-language","title":"Language-Agnostic Bias Detection in Language Models with Bias Probing","date":"2023-05-22","arxiv_id":"2305.13302","n_code_links":1,"syntology":null},{"paper":"/paper/logical-reasoning-for-natural-language","slug":"logical-reasoning-for-natural-language","title":"Atomic Inference for NLI with Generated Facts as Atoms","date":"2023-05-22","arxiv_id":"2305.13214","n_code_links":1,"syntology":null},{"paper":null,"slug":"stock-and-market-index-prediction-using","title":"Stock and market index prediction using Informer network","date":"2023-05-22","arxiv_id":"2305.14382","n_code_links":0,"syntology":null},{"paper":null,"slug":"syntactic-knowledge-via-graph-attention-with","title":"Syntactic Knowledge via Graph Attention with BERT in Machine Translation","date":"2023-05-22","arxiv_id":"2305.13413","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-deeper-autoregressive-approach-to-non","title":"A Deeper (Autoregressive) Approach to Non-Convergent Discourse Parsing","date":"2023-05-21","arxiv_id":"2305.12510","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-symbolic-framework-for-systematic","title":"A Symbolic Framework for Evaluating Mathematical Reasoning and Generalisation with Transformers","date":"2023-05-21","arxiv_id":"2305.12563","n_code_links":0,"syntology":null},{"paper":"/paper/bertrlfuzzer-a-bert-and-reinforcement","slug":"bertrlfuzzer-a-bert-and-reinforcement","title":"BertRLFuzzer: A BERT and Reinforcement Learning Based Fuzzer","date":"2023-05-21","arxiv_id":"2305.12534","n_code_links":1,"syntology":null},{"paper":null,"slug":"f-pabee-flexible-patience-based-early-exiting","title":"F-PABEE: Flexible-patience-based Early Exiting for Single-label and Multi-label text Classification Tasks","date":"2023-05-21","arxiv_id":"2305.11916","n_code_links":0,"syntology":null},{"paper":null,"slug":"infor-coef-information-bottleneck-based","title":"Infor-Coef: Information Bottleneck-based Dynamic Token Downsampling for Compact and Efficient language model","date":"2023-05-21","arxiv_id":"2305.12458","n_code_links":0,"syntology":null},{"paper":null,"slug":"ir-models-and-the-covid-19-pandemic-a","title":"IR Models and the COVID-19 Pandemic: A Comparative Study of Performance and Challenges","date":"2023-05-21","arxiv_id":"2305.12528","n_code_links":0,"syntology":null},{"paper":null,"slug":"cdjur-br-a-golden-collection-of-legal","title":"CDJUR-BR -- A Golden Collection of Legal Document from Brazilian Justice with Fine-Grained Named Entities","date":"2023-05-20","arxiv_id":"2305.18315","n_code_links":0,"syntology":null},{"paper":null,"slug":"sentfin-1-0-entity-aware-sentiment-analysis","title":"SEntFiN 1.0: Entity-Aware Sentiment Analysis for Financial News","date":"2023-05-20","arxiv_id":"2305.12257","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-sequence-to-sequence-approach-for-arabic","title":"A Sequence-to-Sequence Approach for Arabic Pronoun Resolution","date":"2023-05-19","arxiv_id":"2305.11529","n_code_links":0,"syntology":null},{"paper":null,"slug":"eye-spatialnet-spatial-information-extraction","title":"Eye-SpatialNet: Spatial Information Extraction from Ophthalmology Notes","date":"2023-05-19","arxiv_id":"2305.11948","n_code_links":0,"syntology":null},{"paper":null,"slug":"federated-foundation-models-privacy","title":"Federated Foundation Models: Privacy-Preserving and Collaborative Learning for Large Models","date":"2023-05-19","arxiv_id":"2305.11414","n_code_links":0,"syntology":null},{"paper":null,"slug":"language-universal-phonetic-representation-in","title":"Language-Universal Phonetic Representation in Multilingual Speech Pretraining for Low-Resource Speech Recognition","date":"2023-05-19","arxiv_id":"2305.11569","n_code_links":0,"syntology":null},{"paper":null,"slug":"ahead-of-time-p-tuning","title":"Ahead-of-Time P-Tuning","date":"2023-05-18","arxiv_id":"2305.10835","n_code_links":0,"syntology":null},{"paper":"/paper/ditto-a-simple-and-efficient-approach-to","slug":"ditto-a-simple-and-efficient-approach-to","title":"Ditto: A Simple and Efficient Approach to Improve Sentence Embeddings","date":"2023-05-18","arxiv_id":"2305.10786","n_code_links":1,"syntology":null},{"paper":null,"slug":"pdp-parameter-free-differentiable-pruning-is","title":"PDP: Parameter-free Differentiable Pruning is All You Need","date":"2023-05-18","arxiv_id":"2305.11203","n_code_links":0,"syntology":null},{"paper":null,"slug":"trading-syntax-trees-for-wordpieces-target","title":"Trading Syntax Trees for Wordpieces: Target-oriented Opinion Words Extraction with Wordpieces and Aspect Enhancement","date":"2023-05-18","arxiv_id":"2305.11034","n_code_links":0,"syntology":null},{"paper":"/paper/a-quantitative-study-of-nlp-approaches-to","slug":"a-quantitative-study-of-nlp-approaches-to","title":"A quantitative study of NLP approaches to question difficulty estimation","date":"2023-05-17","arxiv_id":"2305.10236","n_code_links":1,"syntology":null},{"paper":"/paper/ad-kd-attribution-driven-knowledge","slug":"ad-kd-attribution-driven-knowledge","title":"AD-KD: Attribution-Driven Knowledge Distillation for Language Model Compression","date":"2023-05-17","arxiv_id":"2305.10010","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":2,"n_instrument":1,"unverified":1,"pointer_only":4,"phrase":"3 ran (of which 1 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["brucewsy/ad-kd"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":1,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/explaining-black-box-text-modules-in-natural","slug":"explaining-black-box-text-modules-in-natural","title":"Explaining black box text modules in natural language with language models","date":"2023-05-17","arxiv_id":"2305.09863","n_code_links":2,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["csinva/imodelsX","microsoft/automated-explanations"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/solving-cosine-similarity-underestimation","slug":"solving-cosine-similarity-underestimation","title":"Solving Cosine Similarity Underestimation between High Frequency Words by L2 Norm Discounting","date":"2023-05-17","arxiv_id":"2305.10610","n_code_links":1,"syntology":null},{"paper":"/paper/berttm-leveraging-contextualized-word","slug":"berttm-leveraging-contextualized-word","title":"CWTM: Leveraging Contextualized Word Embeddings from BERT for Neural Topic Modeling","date":"2023-05-16","arxiv_id":"2305.09329","n_code_links":1,"syntology":null},{"paper":"/paper/measuring-stereotypes-using-entity-centric","slug":"measuring-stereotypes-using-entity-centric","title":"Measuring Dimensions of Self-Presentation in Twitter Bios and their Links to Misinformation Sharing","date":"2023-05-16","arxiv_id":"2305.09548","n_code_links":1,"syntology":null},{"paper":"/paper/weight-inherited-distillation-for-task","slug":"weight-inherited-distillation-for-task","title":"Weight-Inherited Distillation for Task-Agnostic BERT Compression","date":"2023-05-16","arxiv_id":"2305.09098","n_code_links":1,"syntology":null},{"paper":"/paper/coreference-aware-double-channel-attention","slug":"coreference-aware-double-channel-attention","title":"Coreference-aware Double-channel Attention Network for Multi-party Dialogue Reading Comprehension","date":"2023-05-15","arxiv_id":"2305.08348","n_code_links":1,"syntology":null},{"paper":"/paper/knowledge-rumination-for-pre-trained-language","slug":"knowledge-rumination-for-pre-trained-language","title":"Knowledge Rumination for Pre-trained Language Models","date":"2023-05-15","arxiv_id":"2305.08732","n_code_links":1,"syntology":null},{"paper":null,"slug":"private-training-set-inspection-in-mlaas","title":"Private Training Set Inspection in MLaaS","date":"2023-05-15","arxiv_id":"2305.09058","n_code_links":0,"syntology":null},{"paper":null,"slug":"text2gender-a-deep-learning-architecture-for","title":"Text2Gender: A Deep Learning Architecture for Analysis of Blogger's Age and Gender","date":"2023-05-15","arxiv_id":"2305.08633","n_code_links":0,"syntology":null},{"paper":"/paper/matsci-nlp-evaluating-scientific-language","slug":"matsci-nlp-evaluating-scientific-language","title":"MatSci-NLP: Evaluating Scientific Language Models on Materials Science Language Tasks Using Text-to-Schema Modeling","date":"2023-05-14","arxiv_id":"2305.08264","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["banglab-udem-mila/nlp4matsci-acl23"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/gpt-sentinel-distinguishing-human-and-chatgpt","slug":"gpt-sentinel-distinguishing-human-and-chatgpt","title":"GPT-Sentinel: Distinguishing Human and ChatGPT Generated Content","date":"2023-05-13","arxiv_id":"2305.07969","n_code_links":2,"syntology":null},{"paper":null,"slug":"pests-persian-english-cross-lingual-corpus","title":"PESTS: Persian_English Cross Lingual Corpus for Semantic Textual Similarity","date":"2023-05-13","arxiv_id":"2305.07893","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-method-to-automate-the-discharge-summary","title":"A Method to Automate the Discharge Summary Hospital Course for Neurology Patients","date":"2023-05-10","arxiv_id":"2305.06416","n_code_links":0,"syntology":null},{"paper":"/paper/enriching-language-models-with-graph-based","slug":"enriching-language-models-with-graph-based","title":"Enriching language models with graph-based context information to better understand textual data","date":"2023-05-10","arxiv_id":"2305.11070","n_code_links":1,"syntology":null},{"paper":"/paper/a-review-of-vision-language-models-and-their","slug":"a-review-of-vision-language-models-and-their","title":"A Review of Vision-Language Models and their Performance on the Hateful Memes Challenge","date":"2023-05-09","arxiv_id":"2305.06159","n_code_links":1,"syntology":null},{"paper":"/paper/alleviating-over-smoothing-for-unsupervised","slug":"alleviating-over-smoothing-for-unsupervised","title":"Alleviating Over-smoothing for Unsupervised Sentence Representation","date":"2023-05-09","arxiv_id":"2305.06154","n_code_links":1,"syntology":{"ran":0,"of":4,"n_ran_checked":0,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"0 ran · 4 unverified","official":{"repos":["nuochenpku/sscl"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":4,"ran_from_kinds":[]}}},{"paper":null,"slug":"attack-named-entity-recognition-by-entity","title":"Attack Named Entity Recognition by Entity Boundary Interference","date":"2023-05-09","arxiv_id":"2305.05253","n_code_links":0,"syntology":null},{"paper":"/paper/detection-of-depression-on-social-networks","slug":"detection-of-depression-on-social-networks","title":"Detection of depression on social networks using transformers and ensembles","date":"2023-05-09","arxiv_id":"2305.05325","n_code_links":1,"syntology":null},{"paper":null,"slug":"investigating-the-effect-of-sub-word","title":"Effects of sub-word segmentation on performance of transformer language models","date":"2023-05-09","arxiv_id":"2305.05480","n_code_links":0,"syntology":null},{"paper":null,"slug":"strae-autoencoding-for-pre-trained-embeddings","title":"StrAE: Autoencoding for Pre-Trained Embeddings using Explicit Structure","date":"2023-05-09","arxiv_id":"2305.05588","n_code_links":0,"syntology":null},{"paper":null,"slug":"gersteinlab-at-mediqa-chat-2023-clinical-note","title":"GersteinLab at MEDIQA-Chat 2023: Clinical Note Summarization from Doctor-Patient Conversations through Fine-tuning and In-context Learning","date":"2023-05-08","arxiv_id":"2305.05001","n_code_links":0,"syntology":null},{"paper":null,"slug":"precog-exploring-the-relation-between","title":"PreCog: Exploring the Relation between Memorization and Performance in Pre-trained Language Models","date":"2023-05-08","arxiv_id":"2305.04673","n_code_links":0,"syntology":null},{"paper":null,"slug":"vulnerability-detection-using-two-stage-deep","title":"Vulnerability Detection Using Two-Stage Deep Learning Models","date":"2023-05-08","arxiv_id":"2305.09673","n_code_links":0,"syntology":null},{"paper":null,"slug":"stanford-mlab-at-semeval-2023-task-10","title":"Stanford MLab at SemEval-2023 Task 10: Exploring GloVe- and Transformer-Based Methods for the Explainable Detection of Online Sexism","date":"2023-05-07","arxiv_id":"2305.04356","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-the-usage-of-continual-learning-for-out-of","title":"On the Usage of Continual Learning for Out-of-Distribution Generalization in Pre-trained Language Models of Code","date":"2023-05-06","arxiv_id":"2305.04106","n_code_links":0,"syntology":null},{"paper":null,"slug":"rhetorical-role-labeling-of-legal-documents","title":"Rhetorical Role Labeling of Legal Documents using Transformers and Graph Neural Networks","date":"2023-05-06","arxiv_id":"2305.04100","n_code_links":0,"syntology":null},{"paper":null,"slug":"block-the-label-and-noise-an-n-gram-masked","title":"Block the Label and Noise: An N-Gram Masked Speller for Chinese Spell Checking","date":"2023-05-05","arxiv_id":"2305.03314","n_code_links":0,"syntology":null},{"paper":null,"slug":"clac-at-semeval-2023-task-2-comparing-span","title":"CLaC at SemEval-2023 Task 2: Comparing Span-Prediction and Sequence-Labeling approaches for NER","date":"2023-05-05","arxiv_id":"2305.03845","n_code_links":0,"syntology":null},{"paper":null,"slug":"harnessing-the-power-of-bert-in-the-turkish","title":"Harnessing the Power of BERT in the Turkish Clinical Domain: Pretraining Approaches for Limited Data Scenarios","date":"2023-05-05","arxiv_id":"2305.03788","n_code_links":0,"syntology":null},{"paper":"/paper/using-chatgpt-for-entity-matching","slug":"using-chatgpt-for-entity-matching","title":"Using ChatGPT for Entity Matching","date":"2023-05-05","arxiv_id":"2305.03423","n_code_links":1,"syntology":null},{"paper":null,"slug":"enhancing-pashto-text-classification-using","title":"Enhancing Pashto Text Classification using Language Processing Techniques for Single And Multi-Label Analysis","date":"2023-05-04","arxiv_id":"2305.03201","n_code_links":0,"syntology":null},{"paper":"/paper/improving-code-example-recommendations-on","slug":"improving-code-example-recommendations-on","title":"Improving Code Example Recommendations on Informal Documentation Using BERT and Query-Aware LSH: A Comparative Study","date":"2023-05-04","arxiv_id":"2305.03017","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-novel-plagiarism-detection-approach","title":"A Novel Plagiarism Detection Approach Combining BERT-based Word Embedding, Attention-based LSTMs and an Improved Differential Evolution Algorithm","date":"2023-05-03","arxiv_id":"2305.02374","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-bert-and-parsbert-for-analyzing","title":"evaluating bert and parsbert for analyzing persian advertisement data","date":"2023-05-03","arxiv_id":"2305.02426","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-bert-based-scientific-relation","title":"Evaluating BERT-based Scientific Relation Classifiers for Scholarly Knowledge Graph Construction on Digital Library Collections","date":"2023-05-03","arxiv_id":"2305.02291","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-linguistic-properties-of","title":"Exploring Linguistic Properties of Monolingual BERTs with Typological Classification among Languages","date":"2023-05-03","arxiv_id":"2305.02215","n_code_links":0,"syntology":null},{"paper":null,"slug":"cancer-hallmark-classification-using","title":"Improving Cancer Hallmark Classification with BERT-based Deep Learning Approach","date":"2023-05-02","arxiv_id":"2305.03501","n_code_links":0,"syntology":null}],"record_sha256":"5d562950fb1feff866532613fc3d88a13a028f50cfe2e727b6088a871afbe9db","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}