{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/roberta/papers/2","list_of":"/method/roberta","method":"RoBERTa","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":2,"pages_in_order":10,"rows_per_page":100,"rows":[101,200],"of":913,"counts":{"archive_papers_tagged":913,"with_a_code_link":399,"where_syntology_ran_a_sample":87,"not_listed_spam_title":0,"listed":913,"listed_where_code_ran":87,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":66,"every_run_a_failure_of_syntologys_instrument":21,"listed_with_a_run_with_no_instrument_failure":66,"listed_every_run_a_failure_of_syntologys_instrument":21,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/roberta","prev":"/method/roberta","next":"/method/roberta/papers/3","papers":[{"paper":null,"slug":"comparing-unidirectional-bidirectional-and","title":"Comparing Unidirectional, Bidirectional, and Word2vec Models for Discovering Vulnerabilities in Compiled Lifted Code","date":"2024-09-26","arxiv_id":"2409.17513","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-academic-skills-assessment-with-nlp","title":"Improving Academic Skills Assessment with NLP and Ensemble Learning","date":"2024-09-23","arxiv_id":"2409.19013","n_code_links":0,"syntology":null},{"paper":null,"slug":"drift-to-remember","title":"Drift to Remember","date":"2024-09-21","arxiv_id":"2409.13997","n_code_links":0,"syntology":null},{"paper":null,"slug":"hut-a-more-computation-efficient-fine-tuning","title":"HUT: A More Computation Efficient Fine-Tuning Method With Hadamard Updated Transformation","date":"2024-09-20","arxiv_id":"2409.13501","n_code_links":0,"syntology":null},{"paper":"/paper/towards-understanding-evolution-of-science","slug":"towards-understanding-evolution-of-science","title":"Towards understanding evolution of science through language model series","date":"2024-09-15","arxiv_id":"2409.09636","n_code_links":1,"syntology":null},{"paper":null,"slug":"enhanced-online-grooming-detection-employing","title":"Enhanced Online Grooming Detection Employing Context Determination and Message-Level Analysis","date":"2024-09-12","arxiv_id":"2409.07958","n_code_links":0,"syntology":null},{"paper":null,"slug":"constrained-multi-layer-contrastive-learning","title":"Constrained Multi-Layer Contrastive Learning for Implicit Discourse Relationship Recognition","date":"2024-09-07","arxiv_id":"2409.13716","n_code_links":0,"syntology":null},{"paper":"/paper/cbf-llm-safe-control-for-llm-alignment","slug":"cbf-llm-safe-control-for-llm-alignment","title":"CBF-LLM: Safe Control for LLM Alignment","date":"2024-08-28","arxiv_id":"2408.15625","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-self-contained-negation-test-set","title":"The Self-Contained Negation Test Set","date":"2024-08-21","arxiv_id":"2408.11469","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-strategy-to-combine-1stgen-transformers-and","title":"A Strategy to Combine 1stGen Transformers and Open LLMs for Automatic Text Classification","date":"2024-08-19","arxiv_id":"2408.09629","n_code_links":0,"syntology":null},{"paper":null,"slug":"active-learning-for-identifying-disaster","title":"Active Learning for Identifying Disaster-Related Tweets: A Comparison with Keyword Filtering and Generic Fine-Tuning","date":"2024-08-19","arxiv_id":"2408.09914","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-language-of-trauma-modeling-traumatic","title":"The Language of Trauma: Modeling Traumatic Event Descriptions Across Domains with Explainable AI","date":"2024-08-12","arxiv_id":"2408.05977","n_code_links":0,"syntology":null},{"paper":"/paper/is-child-directed-speech-effective-training","slug":"is-child-directed-speech-effective-training","title":"Is Child-Directed Speech Effective Training Data for Language Models?","date":"2024-08-07","arxiv_id":"2408.03617","n_code_links":1,"syntology":{"ran":7,"of":9,"n_ran_checked":5,"n_instrument":2,"unverified":2,"pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 1 violated, 4 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["styfeng/tinydialogues"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"2408-01935","title":"Defining and Evaluating Decision and Composite Risk in Language Models Applied to Natural Language Inference","date":"2024-08-04","arxiv_id":"2408.01935","n_code_links":0,"syntology":null},{"paper":null,"slug":"mistralbsm-leveraging-mistral-7b-for","title":"MistralBSM: Leveraging Mistral-7B for Vehicular Networks Misbehavior Detection","date":"2024-07-26","arxiv_id":"2407.18462","n_code_links":0,"syntology":null},{"paper":null,"slug":"banyan-improved-representation-learning-with","title":"Banyan: Improved Representation Learning with Explicit Structure","date":"2024-07-25","arxiv_id":"2407.17771","n_code_links":0,"syntology":null},{"paper":null,"slug":"roberta-resnext-and-bilstm-with-self","title":"RoBERTa, ResNeXt and BiLSTM with self-attention: The ultimate trio for customer sentiment analysis","date":"2024-07-25","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"artificial-intelligence-in-extracting","title":"Artificial Intelligence in Extracting Diagnostic Data from Dental Records","date":"2024-07-23","arxiv_id":"2407.21050","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-large-language-models-for-anxiety","slug":"evaluating-large-language-models-for-anxiety","title":"Evaluating Large Language Models for Anxiety and Depression Classification using Counseling and Psychotherapy Transcripts","date":"2024-07-18","arxiv_id":"2407.13228","n_code_links":1,"syntology":null},{"paper":null,"slug":"qalam-a-multimodal-llm-for-arabic-optical","title":"Qalam : A Multimodal LLM for Arabic Optical Character and Handwriting Recognition","date":"2024-07-18","arxiv_id":"2407.13559","n_code_links":0,"syntology":null},{"paper":"/paper/sharif-str-at-semeval-2024-task-1-transformer","slug":"sharif-str-at-semeval-2024-task-1-transformer","title":"Sharif-STR at SemEval-2024 Task 1: Transformer as a Regression Model for Fine-Grained Scoring of Textual Semantic Relations","date":"2024-07-17","arxiv_id":"2407.12426","n_code_links":1,"syntology":null},{"paper":null,"slug":"r-sfllm-jamming-resilient-framework-for-split","title":"R-SFLLM: Jamming Resilient Framework for Split Federated Learning with Large Language Models","date":"2024-07-16","arxiv_id":"2407.11654","n_code_links":0,"syntology":null},{"paper":null,"slug":"resource-management-for-low-latency","title":"Resource Management for Low-latency Cooperative Fine-tuning of Foundation Models at the Network Edge","date":"2024-07-13","arxiv_id":"2407.09873","n_code_links":0,"syntology":null},{"paper":null,"slug":"robustness-of-llms-to-perturbations-in-text","title":"Robustness of LLMs to Perturbations in Text","date":"2024-07-12","arxiv_id":"2407.08989","n_code_links":0,"syntology":null},{"paper":"/paper/rosa-random-subspace-adaptation-for-efficient","slug":"rosa-random-subspace-adaptation-for-efficient","title":"ROSA: Random Subspace Adaptation for Efficient Fine-Tuning","date":"2024-07-10","arxiv_id":"2407.07802","n_code_links":1,"syntology":null},{"paper":"/paper/an-empirical-comparison-of-vocabulary","slug":"an-empirical-comparison-of-vocabulary","title":"An Empirical Comparison of Vocabulary Expansion and Initialization Approaches for Language Models","date":"2024-07-08","arxiv_id":"2407.05841","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["AI4Bharat/VocabAdaptation_LLM"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/using-llms-to-label-medical-papers-according","slug":"using-llms-to-label-medical-papers-according","title":"Using LLMs to label medical papers according to the CIViC evidence model","date":"2024-07-05","arxiv_id":"2407.04466","n_code_links":1,"syntology":null},{"paper":null,"slug":"hybrinfox-at-checkthat-2024-task-1-enhancing","title":"HYBRINFOX at CheckThat! 2024 -- Task 1: Enhancing Language Models with Structured Information for Check-Worthiness Estimation","date":"2024-07-04","arxiv_id":"2407.03850","n_code_links":0,"syntology":null},{"paper":null,"slug":"hybrinfox-at-checkthat-2024-task-2-enriching","title":"HYBRINFOX at CheckThat! 2024 -- Task 2: Enriching BERT Models with the Expert System VAGO for Subjectivity Detection","date":"2024-07-04","arxiv_id":"2407.03770","n_code_links":0,"syntology":null},{"paper":null,"slug":"llm-generated-natural-language-meets-scaling","title":"LLM-Generated Natural Language Meets Scaling Laws: New Explorations and Data Augmentation Methods","date":"2024-06-29","arxiv_id":"2407.00322","n_code_links":0,"syntology":null},{"paper":null,"slug":"ragbench-explainable-benchmark-for-retrieval","title":"RAGBench: Explainable Benchmark for Retrieval-Augmented Generation Systems","date":"2024-06-25","arxiv_id":"2407.11005","n_code_links":0,"syntology":null},{"paper":"/paper/this-paper-had-the-smartest-reviewers","slug":"this-paper-had-the-smartest-reviewers","title":"This Paper Had the Smartest Reviewers -- Flattery Detection Utilising an Audio-Textual Transformer-Based Approach","date":"2024-06-25","arxiv_id":"2406.17667","n_code_links":1,"syntology":null},{"paper":"/paper/unlocking-continual-learning-abilities-in","slug":"unlocking-continual-learning-abilities-in","title":"Unlocking Continual Learning Abilities in Language Models","date":"2024-06-25","arxiv_id":"2406.17245","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":5,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["wenyudu/migu"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/unambiguous-recognition-should-not-rely","slug":"unambiguous-recognition-should-not-rely","title":"MixTex: Unambiguous Recognition Should Not Rely Solely on Real Data","date":"2024-06-24","arxiv_id":"2406.17148","n_code_links":1,"syntology":null},{"paper":"/paper/welldunn-on-the-robustness-and-explainability","slug":"welldunn-on-the-robustness-and-explainability","title":"WellDunn: On the Robustness and Explainability of Language Models and Large Language Models in Identifying Wellness Dimensions","date":"2024-06-17","arxiv_id":"2406.12058","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["vedantpalit/WellDunn"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/sharelora-parameter-efficient-and-robust","slug":"sharelora-parameter-efficient-and-robust","title":"ShareLoRA: Parameter Efficient and Robust Large Language Model Fine-tuning via Shared Low-Rank Adaptation","date":"2024-06-16","arxiv_id":"2406.10785","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":0,"n_instrument":5,"unverified":1,"pointer_only":6,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 5 where Syntology's instrument failed) · 1 unverified","official":{"repos":["Rain9876/ShareLoRA"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/the-devil-is-in-the-neurons-interpreting-and","slug":"the-devil-is-in-the-neurons-interpreting-and","title":"The Devil is in the Neurons: Interpreting and Mitigating Social Biases in Pre-trained Language Models","date":"2024-06-14","arxiv_id":"2406.10130","n_code_links":1,"syntology":null},{"paper":"/paper/bpe-knockout-pruning-pre-existing-bpe","slug":"bpe-knockout-pruning-pre-existing-bpe","title":"BPE-knockout: Pruning Pre-existing BPE Tokenisers with Backwards-compatible Morphological Semi-supervision","date":"2024-06-13","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/label-aware-hard-negative-sampling-strategies","slug":"label-aware-hard-negative-sampling-strategies","title":"Label-aware Hard Negative Sampling Strategies with Momentum Contrastive Learning for Implicit Hate Speech Detection","date":"2024-06-12","arxiv_id":"2406.07886","n_code_links":1,"syntology":null},{"paper":"/paper/advancing-semantic-textual-similarity","slug":"advancing-semantic-textual-similarity","title":"Advancing Semantic Textual Similarity Modeling: A Regression Framework with Translated ReLU and Smooth K2 Loss","date":"2024-06-08","arxiv_id":"2406.05326","n_code_links":2,"syntology":null},{"paper":"/paper/bamo-at-semeval-2024-task-9-brainteaser-a","slug":"bamo-at-semeval-2024-task-9-brainteaser-a","title":"BAMO at SemEval-2024 Task 9: BRAINTEASER: A Novel Task Defying Common Sense","date":"2024-06-07","arxiv_id":"2406.04947","n_code_links":1,"syntology":null},{"paper":"/paper/rico-reddit-ideological-communities","slug":"rico-reddit-ideological-communities","title":"RICo: Reddit ideological communities","date":"2024-06-05","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"probing-the-category-of-verbal-aspect-in","title":"Probing the Category of Verbal Aspect in Transformer Language Models","date":"2024-06-04","arxiv_id":"2406.02335","n_code_links":0,"syntology":null},{"paper":null,"slug":"annotation-guidelines-based-knowledge","title":"Annotation Guidelines-Based Knowledge Augmentation: Towards Enhancing Large Language Models for Educational Text Classification","date":"2024-06-03","arxiv_id":"2406.00954","n_code_links":0,"syntology":null},{"paper":"/paper/factgenius-combining-zero-shot-prompting-and","slug":"factgenius-combining-zero-shot-prompting-and","title":"FactGenius: Combining Zero-Shot Prompting and Fuzzy Relation Mining to Improve Fact Verification with Knowledge Graphs","date":"2024-06-03","arxiv_id":"2406.01311","n_code_links":1,"syntology":null},{"paper":"/paper/roberta-bilstm-a-context-aware-hybrid-model","slug":"roberta-bilstm-a-context-aware-hybrid-model","title":"RoBERTa-BiLSTM: A Context-Aware Hybrid Model for Sentiment Analysis","date":"2024-06-01","arxiv_id":"2406.00367","n_code_links":1,"syntology":null},{"paper":null,"slug":"bi-directional-transformers-vs-word2vec","title":"Bi-Directional Transformers vs. word2vec: Discovering Vulnerabilities in Lifted Compiled Code","date":"2024-05-31","arxiv_id":"2405.20611","n_code_links":0,"syntology":null},{"paper":null,"slug":"ensemble-model-with-bert-roberta-and-xlnet","title":"Ensemble Model With Bert,Roberta and Xlnet For Molecular property prediction","date":"2024-05-30","arxiv_id":"2406.06553","n_code_links":0,"syntology":null},{"paper":null,"slug":"performance-evaluation-of-reddit-comments","title":"Performance evaluation of Reddit Comments using Machine Learning and Natural Language Processing methods in Sentiment Analysis","date":"2024-05-27","arxiv_id":"2405.16810","n_code_links":0,"syntology":null},{"paper":null,"slug":"incremental-comprehension-of-garden-path","title":"Incremental Comprehension of Garden-Path Sentences by Large Language Models: Semantic Interpretation, Syntactic Re-Analysis, and Attention","date":"2024-05-25","arxiv_id":"2405.16042","n_code_links":0,"syntology":null},{"paper":null,"slug":"textit-comet-a-underline-com-munication","title":"Comet: A Communication-efficient and Performant Approximation for Private Transformer Inference","date":"2024-05-24","arxiv_id":"2405.17485","n_code_links":0,"syntology":null},{"paper":null,"slug":"crema-crisis-response-through-computational","title":"CReMa: Crisis Response through Computational Identification and Matching of Cross-Lingual Requests and Offers Shared on Social Media","date":"2024-05-20","arxiv_id":"2405.11897","n_code_links":0,"syntology":null},{"paper":null,"slug":"explainabledetector-exploring-transformer","title":"ExplainableDetector: Exploring Transformer-based Language Modeling Approach for SMS Spam Detection with Explainability Analysis","date":"2024-05-12","arxiv_id":"2405.08026","n_code_links":0,"syntology":null},{"paper":null,"slug":"tacoere-cluster-aware-compression-for-event","title":"TacoERE: Cluster-aware Compression for Event Relation Extraction","date":"2024-05-11","arxiv_id":"2405.06890","n_code_links":0,"syntology":null},{"paper":null,"slug":"reddit-impacts-a-named-entity-recognition","title":"Reddit-Impacts: A Named Entity Recognition Dataset for Analyzing Clinical and Social Effects of Substance Use Derived from Social Media","date":"2024-05-09","arxiv_id":"2405.06145","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-anti-semitic-hate-speech-using","title":"Detecting Anti-Semitic Hate Speech using Transformer-based Large Language Models","date":"2024-05-06","arxiv_id":"2405.03794","n_code_links":0,"syntology":null},{"paper":"/paper/structural-pruning-of-pre-trained-language","slug":"structural-pruning-of-pre-trained-language","title":"Structural Pruning of Pre-trained Language Models via Neural Architecture Search","date":"2024-05-03","arxiv_id":"2405.02267","n_code_links":1,"syntology":null},{"paper":null,"slug":"early-transformers-a-study-on-efficient","title":"Early Transformers: A study on Efficient Training of Transformer Models through Early-Bird Lottery Tickets","date":"2024-05-02","arxiv_id":"2405.02353","n_code_links":0,"syntology":null},{"paper":"/paper/investigating-wit-creativity-and","slug":"investigating-wit-creativity-and","title":"Investigating Wit, Creativity, and Detectability of Large Language Models in Domain-Specific Writing Style Adaptation of Reddit's Showerthoughts","date":"2024-05-02","arxiv_id":"2405.01660","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-named-entity-recognition-and-topic-modeling","title":"A Named Entity Recognition and Topic Modeling-based Solution for Locating and Better Assessment of Natural Disasters in Social Media","date":"2024-05-01","arxiv_id":"2405.00903","n_code_links":0,"syntology":null},{"paper":null,"slug":"federa-efficient-fine-tuning-of-language","title":"FeDeRA:Efficient Fine-tuning of Language Models in Federated Learning Leveraging Weight Decomposition","date":"2024-04-29","arxiv_id":"2404.18848","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-perplexity-predict-fine-tuning","title":"Can Perplexity Predict Fine-Tuning Performance? An Investigation of Tokenization Effects on Sequential Language Models for Nepali","date":"2024-04-28","arxiv_id":"2404.18071","n_code_links":0,"syntology":null},{"paper":"/paper/calc-cmu-at-semeval-2024-task-7-pre-calc","slug":"calc-cmu-at-semeval-2024-task-7-pre-calc","title":"Pre-Calc: Learning to Use the Calculator Improves Numeracy in Language Models","date":"2024-04-22","arxiv_id":"2404.14355","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":0,"n_instrument":3,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","official":{"repos":["calc-cmu/pre-calc"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"marking-visual-grading-with-highlighting","title":"Marking: Visual Grading with Highlighting Errors and Annotating Missing Bits","date":"2024-04-22","arxiv_id":"2404.14301","n_code_links":0,"syntology":null},{"paper":"/paper/do-english-named-entity-recognizers-work-well","slug":"do-english-named-entity-recognizers-work-well","title":"Do \"English\" Named Entity Recognizers Work Well on Global Englishes?","date":"2024-04-20","arxiv_id":"2404.13465","n_code_links":2,"syntology":null},{"paper":"/paper/evaluating-subword-tokenization-alien-subword","slug":"evaluating-subword-tokenization-alien-subword","title":"Evaluating Subword Tokenization: Alien Subword Composition and OOV Generalization Challenge","date":"2024-04-20","arxiv_id":"2404.13292","n_code_links":1,"syntology":null},{"paper":null,"slug":"enabling-natural-zero-shot-prompting-on","title":"Enabling Natural Zero-Shot Prompting on Encoder Models via Statement-Tuning","date":"2024-04-19","arxiv_id":"2404.12897","n_code_links":0,"syntology":null},{"paper":null,"slug":"emrqa-msquad-a-medical-dataset-structured","title":"emrQA-msquad: A Medical Dataset Structured with the SQuAD V2.0 Framework, Enriched with emrQA Medical Information","date":"2024-04-18","arxiv_id":"2404.12050","n_code_links":0,"syntology":null},{"paper":null,"slug":"demystifying-legalese-an-automated-approach","title":"Demystifying Legalese: An Automated Approach for Summarizing and Analyzing Overlaps in Privacy Policies and Terms of Service","date":"2024-04-17","arxiv_id":"2404.13087","n_code_links":0,"syntology":null},{"paper":null,"slug":"relational-graph-convolutional-networks-for-1","title":"Relational Graph Convolutional Networks for Sentiment Analysis","date":"2024-04-16","arxiv_id":"2404.13079","n_code_links":0,"syntology":null},{"paper":null,"slug":"comprehensive-study-on-german-language-models","title":"Comprehensive Study on German Language Models for Clinical and Biomedical Text Understanding","date":"2024-04-08","arxiv_id":"2404.05694","n_code_links":0,"syntology":null},{"paper":"/paper/investigating-the-robustness-of-modelling","slug":"investigating-the-robustness-of-modelling","title":"Investigating the Robustness of Modelling Decisions for Few-Shot Cross-Topic Stance Detection: A Preregistered Study","date":"2024-04-05","arxiv_id":"2404.03987","n_code_links":1,"syntology":null},{"paper":"/paper/bcamirs-at-semeval-2024-task-4-beyond-words-a","slug":"bcamirs-at-semeval-2024-task-4-beyond-words-a","title":"BCAmirs at SemEval-2024 Task 4: Beyond Words: A Multimodal and Multilingual Exploration of Persuasion in Memes","date":"2024-04-03","arxiv_id":"2404.03022","n_code_links":1,"syntology":null},{"paper":"/paper/explainable-deep-learning-a-visual-analytics","slug":"explainable-deep-learning-a-visual-analytics","title":"Explainable Deep Learning: A Visual Analytics Approach with Transition Matrices","date":"2024-03-29","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"alloybert-alloy-property-prediction-with","title":"AlloyBERT: Alloy Property Prediction with Large Language Models","date":"2024-03-28","arxiv_id":"2403.19783","n_code_links":0,"syntology":null},{"paper":null,"slug":"acted-automatic-acquisition-of-typical-event","title":"AcTED: Automatic Acquisition of Typical Event Duration for Semi-supervised Temporal Commonsense QA","date":"2024-03-27","arxiv_id":"2403.18504","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-large-language-models-for-health-1","title":"Evaluating Large Language Models for Health-Related Text Classification Tasks with Public Social Media Data","date":"2024-03-27","arxiv_id":"2403.19031","n_code_links":0,"syntology":null},{"paper":"/paper/semrode-macro-adversarial-training-to-learn","slug":"semrode-macro-adversarial-training-to-learn","title":"SemRoDe: Macro Adversarial Training to Learn Representations That are Robust to Word-Level Attacks","date":"2024-03-27","arxiv_id":"2403.18423","n_code_links":1,"syntology":{"ran":1,"of":3,"n_ran_checked":1,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["aniloid2/semrode-macroadversarialtraining"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/fingerprinting-web-servers-through","slug":"fingerprinting-web-servers-through","title":"Fingerprinting web servers through Transformer-encoded HTTP response headers","date":"2024-03-26","arxiv_id":"2404.00056","n_code_links":1,"syntology":null},{"paper":"/paper/llambert-large-scale-low-cost-data-annotation","slug":"llambert-large-scale-low-cost-data-annotation","title":"LlamBERT: Large-scale low-cost data annotation in NLP","date":"2024-03-23","arxiv_id":"2403.15938","n_code_links":1,"syntology":null},{"paper":"/paper/ax-to-grind-urdu-benchmark-dataset-for-urdu","slug":"ax-to-grind-urdu-benchmark-dataset-for-urdu","title":"Ax-to-Grind Urdu: Benchmark Dataset for Urdu Fake News Detection","date":"2024-03-20","arxiv_id":"2403.14037","n_code_links":1,"syntology":null},{"paper":"/paper/cicle-conformal-in-context-learning-for","slug":"cicle-conformal-in-context-learning-for","title":"CICLe: Conformal In-Context Learning for Largescale Multi-Class Food Risk Classification","date":"2024-03-18","arxiv_id":"2403.11904","n_code_links":1,"syntology":null},{"paper":"/paper/lookupffn-making-transformers-compute-lite","slug":"lookupffn-making-transformers-compute-lite","title":"LookupFFN: Making Transformers Compute-lite for CPU inference","date":"2024-03-12","arxiv_id":"2403.07221","n_code_links":1,"syntology":{"ran":6,"of":6,"n_ran_checked":6,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["mlpen/lookupffn"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/galore-memory-efficient-llm-training-by","slug":"galore-memory-efficient-llm-training-by","title":"GaLore: Memory-Efficient LLM Training by Gradient Low-Rank Projection","date":"2024-03-06","arxiv_id":"2403.03507","n_code_links":3,"syntology":null},{"paper":"/paper/eee-qa-exploring-effective-and-efficient","slug":"eee-qa-exploring-effective-and-efficient","title":"EEE-QA: Exploring Effective and Efficient Question-Answer Representations","date":"2024-03-04","arxiv_id":"2403.02176","n_code_links":1,"syntology":null},{"paper":"/paper/analysis-of-privacy-leakage-in-federated","slug":"analysis-of-privacy-leakage-in-federated","title":"Analysis of Privacy Leakage in Federated Large Language Models","date":"2024-03-02","arxiv_id":"2403.04784","n_code_links":1,"syntology":{"ran":4,"of":6,"n_ran_checked":3,"n_instrument":1,"unverified":2,"pointer_only":6,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["vunhatminh/fl_attacks"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/leveraging-pre-trained-language-models-for-3","slug":"leveraging-pre-trained-language-models-for-3","title":"Leveraging pre-trained language models for code generation","date":"2024-02-29","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"pelle-encoder-based-language-models-for","title":"PeLLE: Encoder-based language models for Brazilian Portuguese based on open data","date":"2024-02-29","arxiv_id":"2402.19204","n_code_links":0,"syntology":null},{"paper":"/paper/multi-task-media-bias-analysis-generalization","slug":"multi-task-media-bias-analysis-generalization","title":"MAGPIE: Multi-Task Media-Bias Analysis Generalization for Pre-Trained Identification of Expressions","date":"2024-02-27","arxiv_id":"2403.07910","n_code_links":1,"syntology":null},{"paper":"/paper/asymmetry-in-low-rank-adapters-of-foundation","slug":"asymmetry-in-low-rank-adapters-of-foundation","title":"Asymmetry in Low-Rank Adapters of Foundation Models","date":"2024-02-26","arxiv_id":"2402.16842","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":5,"n_instrument":0,"unverified":1,"pointer_only":6,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["Jiacheng-Zhu-AIML/AsymmetryLoRA"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/semeval-2024-task-8-weighted-layer-averaging","slug":"semeval-2024-task-8-weighted-layer-averaging","title":"SemEval-2024 Task 8: Weighted Layer Averaging RoBERTa for Black-Box Machine-Generated Text Detection","date":"2024-02-24","arxiv_id":"2402.15873","n_code_links":1,"syntology":null},{"paper":"/paper/advancing-parameter-efficiency-in-fine-tuning","slug":"advancing-parameter-efficiency-in-fine-tuning","title":"Advancing Parameter Efficiency in Fine-tuning via Representation Editing","date":"2024-02-23","arxiv_id":"2402.15179","n_code_links":2,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"2 ran (of which 2 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified; every one of the 2 samples that ran constructed an object rather than computing a result","official":{"repos":["mlwu22/red"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"towards-efficient-active-learning-in-nlp-via","title":"Towards Efficient Active Learning in NLP via Pretrained Representations","date":"2024-02-23","arxiv_id":"2402.15613","n_code_links":0,"syntology":null},{"paper":"/paper/can-gnn-be-good-adapter-for-llms","slug":"can-gnn-be-good-adapter-for-llms","title":"Can GNN be Good Adapter for LLMs?","date":"2024-02-20","arxiv_id":"2402.12984","n_code_links":2,"syntology":{"ran":6,"of":8,"n_ran_checked":6,"n_instrument":0,"unverified":2,"pointer_only":8,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["zjunet/graphadapter","hxttkl/GraphAdapter"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/acquiring-clean-language-models-from-backdoor","slug":"acquiring-clean-language-models-from-backdoor","title":"Acquiring Clean Language Models from Backdoor Poisoned Datasets by Downscaling Frequency Space","date":"2024-02-19","arxiv_id":"2402.12026","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["zrw00/musclelora"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/a-curious-case-of-searching-for-the","slug":"a-curious-case-of-searching-for-the","title":"A Curious Case of Searching for the Correlation between Training Data and Adversarial Robustness of Transformer Textual Models","date":"2024-02-18","arxiv_id":"2402.11469","n_code_links":1,"syntology":null},{"paper":null,"slug":"detecting-a-proxy-for-potential-comorbid-adhd","title":"Detecting a Proxy for Potential Comorbid ADHD in People Reporting Anxiety Symptoms from Social Media Data","date":"2024-02-17","arxiv_id":"2403.05561","n_code_links":0,"syntology":null},{"paper":null,"slug":"efficient-models-for-the-detection-of-hate","title":"Efficient Models for the Detection of Hate, Abuse and Profanity","date":"2024-02-08","arxiv_id":"2402.05624","n_code_links":0,"syntology":null},{"paper":null,"slug":"neural-models-for-source-code-synthesis-and","title":"Neural Models for Source Code Synthesis and Completion","date":"2024-02-08","arxiv_id":"2402.06690","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-use-of-a-large-language-model-for","title":"The Use of a Large Language Model for Cyberbullying Detection","date":"2024-02-06","arxiv_id":"2402.04088","n_code_links":0,"syntology":null}],"record_sha256":"292517f8ea19e866aea7f2f60687992aab52dc7319b24039adf85d43574e2811","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}