{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/wordpiece/papers/25","list_of":"/method/wordpiece","method":"WordPiece","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":25,"pages_in_order":71,"rows_per_page":100,"rows":[2401,2500],"of":7063,"counts":{"archive_papers_tagged":7063,"with_a_code_link":2910,"where_syntology_ran_a_sample":650,"not_listed_spam_title":0,"listed":7063,"listed_where_code_ran":650,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":529,"every_run_a_failure_of_syntologys_instrument":121,"listed_with_a_run_with_no_instrument_failure":529,"listed_every_run_a_failure_of_syntologys_instrument":121,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/wordpiece","prev":"/method/wordpiece/papers/24","next":"/method/wordpiece/papers/26","papers":[{"paper":null,"slug":"protoner-few-shot-incremental-learning-for","title":"ProtoNER: Few shot Incremental Learning for Named Entity Recognition using Prototypical Networks","date":"2023-10-03","arxiv_id":"2310.02372","n_code_links":0,"syntology":null},{"paper":"/paper/label-supervised-llama-finetuning","slug":"label-supervised-llama-finetuning","title":"Label Supervised LLaMA Finetuning","date":"2023-10-02","arxiv_id":"2310.01208","n_code_links":2,"syntology":null},{"paper":null,"slug":"natural-language-models-for-data","title":"Natural Language Models for Data Visualization Utilizing nvBench Dataset","date":"2023-10-02","arxiv_id":"2310.00832","n_code_links":0,"syntology":null},{"paper":null,"slug":"target-aware-contextual-political-bias","title":"Target-Aware Contextual Political Bias Detection in News","date":"2023-10-02","arxiv_id":"2310.01138","n_code_links":0,"syntology":null},{"paper":"/paper/question-answering-model-for-schizophrenia","slug":"question-answering-model-for-schizophrenia","title":"Question-Answering Model for Schizophrenia Symptoms and Their Impact on Daily Life using Mental Health Forums Data","date":"2023-09-30","arxiv_id":"2310.00448","n_code_links":0,"syntology":null},{"paper":"/paper/relbert-embedding-relations-with-language","slug":"relbert-embedding-relations-with-language","title":"RelBERT: Embedding Relations with Language Models","date":"2023-09-30","arxiv_id":"2310.00299","n_code_links":1,"syntology":null},{"paper":null,"slug":"intuitive-or-dependent-investigating-llms","title":"Intuitive or Dependent? Investigating LLMs' Behavior Style to Conflicting Prompts","date":"2023-09-29","arxiv_id":"2309.17415","n_code_links":0,"syntology":null},{"paper":"/paper/hallucination-reduction-in-long-input-text","slug":"hallucination-reduction-in-long-input-text","title":"Hallucination Reduction in Long Input Text Summarization","date":"2023-09-28","arxiv_id":"2309.16781","n_code_links":1,"syntology":null},{"paper":null,"slug":"mededit-model-editing-for-medical-question","title":"MKRAG: Medical Knowledge Retrieval Augmented Generation for Medical Question Answering","date":"2023-09-27","arxiv_id":"2309.16035","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-nlp-benchmark-dataset-for-assessing","title":"An NLP Benchmark Dataset for Assessing Corporate Climate Policy Engagement","date":"2023-09-26","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/capp-130-a-corpus-of-chinese-application","slug":"capp-130-a-corpus-of-chinese-application","title":"CAPP-130: A Corpus of Chinese Application Privacy Policy Summarization and Interpretation","date":"2023-09-26","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"low-rank-adaptation-of-large-language-model","title":"Low-rank Adaptation of Large Language Model Rescoring for Parameter-Efficient Speech Recognition","date":"2023-09-26","arxiv_id":"2309.15223","n_code_links":0,"syntology":null},{"paper":"/paper/ragas-automated-evaluation-of-retrieval","slug":"ragas-automated-evaluation-of-retrieval","title":"RAGAS: Automated Evaluation of Retrieval Augmented Generation","date":"2023-09-26","arxiv_id":"2309.15217","n_code_links":3,"syntology":{"ran":4,"of":6,"n_ran_checked":4,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":null}},{"paper":null,"slug":"comprehensive-overview-of-named-entity","title":"Comprehensive Overview of Named Entity Recognition: Models, Domain-Specific Applications and Challenges","date":"2023-09-25","arxiv_id":"2309.14084","n_code_links":0,"syntology":null},{"paper":null,"slug":"accelerating-large-batch-training-via","title":"Accelerating Large Batch Training via Gradient Signal to Noise Ratio (GSNR)","date":"2023-09-24","arxiv_id":"2309.13681","n_code_links":0,"syntology":null},{"paper":"/paper/seeing-is-not-always-believing-invisible","slug":"seeing-is-not-always-believing-invisible","title":"Seeing Is Not Always Believing: Invisible Collision Attack and Defence on Pre-Trained Models","date":"2023-09-24","arxiv_id":"2309.13579","n_code_links":1,"syntology":null},{"paper":"/paper/lexical-squad-multimodal-hate-speech-event","slug":"lexical-squad-multimodal-hate-speech-event","title":"Lexical Squad@Multimodal Hate Speech Event Detection 2023: Multimodal Hate Speech Detection using Fused Ensemble Approach","date":"2023-09-23","arxiv_id":"2309.13354","n_code_links":1,"syntology":null},{"paper":"/paper/amplify-attention-based-mixup-for-performance","slug":"amplify-attention-based-mixup-for-performance","title":"AMPLIFY:Attention-based Mixup for Performance Improvement and Label Smoothing in Transformer","date":"2023-09-22","arxiv_id":"2309.12689","n_code_links":1,"syntology":null},{"paper":"/paper/toproberta-topology-aware-authorship","slug":"toproberta-topology-aware-authorship","title":"TOPFORMER: Topology-Aware Authorship Attribution of Deepfake Texts with Diverse Writing Styles","date":"2023-09-22","arxiv_id":"2309.12934","n_code_links":1,"syntology":null},{"paper":"/paper/bad-actor-good-advisor-exploring-the-role-of","slug":"bad-actor-good-advisor-exploring-the-role-of","title":"Bad Actor, Good Advisor: Exploring the Role of Large Language Models in Fake News Detection","date":"2023-09-21","arxiv_id":"2309.12247","n_code_links":1,"syntology":null},{"paper":"/paper/bayestune-bayesian-sparse-deep-model-fine","slug":"bayestune-bayesian-sparse-deep-model-fine","title":"BayesTune: Bayesian Sparse Deep Model Fine-tuning","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/implicit-differentiable-outlier-detection","slug":"implicit-differentiable-outlier-detection","title":"Implicit Differentiable Outlier Detection Enable Robust Deep Multimodal Analysis","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/making-scalable-meta-learning-practical","slug":"making-scalable-meta-learning-practical","title":"Making Scalable Meta Learning Practical","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/marich-a-query-efficient-distributionally-1","slug":"marich-a-query-efficient-distributionally-1","title":"Marich: A Query-efficient Distributionally Equivalent Model Extraction Attack","date":"2023-09-21","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/on-the-relationship-between-skill-neurons-and","slug":"on-the-relationship-between-skill-neurons-and","title":"On the Relationship between Skill Neurons and Robustness in Prompt Tuning","date":"2023-09-21","arxiv_id":"2309.12263","n_code_links":1,"syntology":null},{"paper":null,"slug":"re-exploring-the-role-of-grammar-and-word","title":"[Re] Exploring the Role of Grammar and Word Choice in Bias Toward African American English (AAE) in Hate Speech Classification","date":"2023-09-21","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"slhcat-mapping-wikipedia-categories-and-lists","title":"SLHCat: Mapping Wikipedia Categories and Lists to DBpedia by Leveraging Semantic, Lexical, and Hierarchical Features","date":"2023-09-21","arxiv_id":"2309.11791","n_code_links":0,"syntology":null},{"paper":null,"slug":"spiced-news-similarity-detection-dataset-with","title":"SPICED: News Similarity Detection Dataset with Multiple Topics and Complexity Levels","date":"2023-09-21","arxiv_id":"2309.13080","n_code_links":0,"syntology":null},{"paper":null,"slug":"stock-market-sentiment-classification-and","title":"Stock Market Sentiment Classification and Backtesting via Fine-tuned BERT","date":"2023-09-21","arxiv_id":"2309.11979","n_code_links":0,"syntology":null},{"paper":"/paper/the-cambridge-law-corpus-a-corpus-for-legal-1","slug":"the-cambridge-law-corpus-a-corpus-for-legal-1","title":"The Cambridge Law Corpus: A Dataset for Legal AI Research","date":"2023-09-21","arxiv_id":"2309.12269","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-efficient-pre-trained-language-model","title":"Towards Efficient Pre-Trained Language Model via Feature Correlation Distillation","date":"2023-09-21","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"attentionmix-data-augmentation-method-that","title":"AttentionMix: Data augmentation method that relies on BERT attention mechanism","date":"2023-09-20","arxiv_id":"2309.11104","n_code_links":0,"syntology":null},{"paper":"/paper/cot-bert-enhancing-unsupervised-sentence","slug":"cot-bert-enhancing-unsupervised-sentence","title":"CoT-BERT: Enhancing Unsupervised Sentence Representation through Chain-of-Thought","date":"2023-09-20","arxiv_id":"2309.11143","n_code_links":2,"syntology":null},{"paper":null,"slug":"gpt-molberta-gpt-molecular-features-language","title":"GPT-MolBERTa: GPT Molecular Features Language Model for molecular property prediction","date":"2023-09-20","arxiv_id":"2310.03030","n_code_links":0,"syntology":null},{"paper":"/paper/sequence-to-sequence-spanish-pre-trained","slug":"sequence-to-sequence-spanish-pre-trained","title":"Sequence-to-Sequence Spanish Pre-trained Language Models","date":"2023-09-20","arxiv_id":"2309.11259","n_code_links":1,"syntology":null},{"paper":null,"slug":"mixed-distil-bert-code-mixed-language","title":"Mixed-Distil-BERT: Code-mixed Language Modeling for Bangla, English, and Hindi","date":"2023-09-19","arxiv_id":"2309.10272","n_code_links":0,"syntology":null},{"paper":"/paper/facilitating-nsfw-text-detection-in-open","slug":"facilitating-nsfw-text-detection-in-open","title":"Facilitating NSFW Text Detection in Open-Domain Dialogue Systems via Knowledge Distillation","date":"2023-09-18","arxiv_id":"2309.09749","n_code_links":1,"syntology":null},{"paper":null,"slug":"proposition-from-the-perspective-of-chinese","title":"Proposition from the Perspective of Chinese Language: A Chinese Proposition Classification Evaluation Benchmark","date":"2023-09-18","arxiv_id":"2309.09602","n_code_links":0,"syntology":null},{"paper":"/paper/detecting-covariate-drift-in-text-data-using","slug":"detecting-covariate-drift-in-text-data-using","title":"Detecting covariate drift in text data using document embeddings and dimensionality reduction","date":"2023-09-17","arxiv_id":"2309.10000","n_code_links":1,"syntology":null},{"paper":"/paper/has-sentiment-returned-to-the-pre-pandemic","slug":"has-sentiment-returned-to-the-pre-pandemic","title":"Has Sentiment Returned to the Pre-pandemic Level? A Sentiment Analysis Using U.S. College Subreddit Data from 2019 to 2022","date":"2023-09-16","arxiv_id":"2309.08845","n_code_links":1,"syntology":null},{"paper":"/paper/albner-a-corpus-for-named-entity-recognition","slug":"albner-a-corpus-for-named-entity-recognition","title":"AlbNER: A Corpus for Named Entity Recognition in Albanian","date":"2023-09-15","arxiv_id":"2309.08741","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-relevant-information-in-high-volume","title":"Detecting Relevant Information in High-Volume Chat Logs: Keyphrase Extraction for Grooming and Drug Dealing Forensic Analysis","date":"2023-09-15","arxiv_id":"2311.04905","n_code_links":0,"syntology":null},{"paper":"/paper/structural-self-supervised-objectives-for","slug":"structural-self-supervised-objectives-for","title":"Structural Self-Supervised Objectives for Transformers","date":"2023-09-15","arxiv_id":"2309.08272","n_code_links":1,"syntology":null},{"paper":"/paper/transformer-based-punctuation-restoration-for","slug":"transformer-based-punctuation-restoration-for","title":"Transformer Based Punctuation Restoration for Turkish","date":"2023-09-15","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"vulnsense-efficient-vulnerability-detection","title":"VulnSense: Efficient Vulnerability Detection in Ethereum Smart Contracts by Multimodal Learning with Graph Neural Network and Language Model","date":"2023-09-15","arxiv_id":"2309.08474","n_code_links":0,"syntology":null},{"paper":null,"slug":"automatic-data-visualization-generation-from","title":"Automatic Data Visualization Generation from Chinese Natural Language Questions","date":"2023-09-14","arxiv_id":"2309.07650","n_code_links":0,"syntology":null},{"paper":null,"slug":"debcse-rethinking-unsupervised-contrastive","title":"DebCSE: Rethinking Unsupervised Contrastive Sentence Embedding Learning in the Debiasing Perspective","date":"2023-09-14","arxiv_id":"2309.07396","n_code_links":0,"syntology":null},{"paper":"/paper/encodecmae-leveraging-neural-codecs-for","slug":"encodecmae-leveraging-neural-codecs-for","title":"EnCodecMAE: Leveraging neural codecs for universal audio representation learning","date":"2023-09-14","arxiv_id":"2309.07391","n_code_links":2,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["habla-liaa/encodecmae"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"text-classification-of-cancer-clinical-trial","title":"Text Classification of Cancer Clinical Trial Eligibility Criteria","date":"2023-09-14","arxiv_id":"2309.07812","n_code_links":0,"syntology":null},{"paper":"/paper/2309-05951","slug":"2309-05951","title":"Balanced and Explainable Social Media Analysis for Public Health with Large Language Models","date":"2023-09-12","arxiv_id":"2309.05951","n_code_links":1,"syntology":null},{"paper":null,"slug":"automatically-estimating-the-effort-required","title":"PRESTI: Predicting Repayment Effort of Self-Admitted Technical Debt Using Textual Information","date":"2023-09-12","arxiv_id":"2309.06020","n_code_links":0,"syntology":null},{"paper":null,"slug":"overview-of-memotion-3-sentiment-and-emotion","title":"Overview of Memotion 3: Sentiment and Emotion Analysis of Codemixed Hinglish Memes","date":"2023-09-12","arxiv_id":"2309.06517","n_code_links":0,"syntology":null},{"paper":"/paper/applying-biobert-to-extract-germline-gene","slug":"applying-biobert-to-extract-germline-gene","title":"Applying BioBERT to Extract Germline Gene-Disease Associations for Building a Knowledge Graph from the Biomedical Literature","date":"2023-09-11","arxiv_id":"2309.13061","n_code_links":1,"syntology":null},{"paper":null,"slug":"crisistransformers-pre-trained-language","title":"CrisisTransformers: Pre-trained language models and sentence encoders for crisis-related social media texts","date":"2023-09-11","arxiv_id":"2309.05494","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-natural-language-biases-with-prompt","title":"Detecting Natural Language Biases with Prompt-based Learning","date":"2023-09-11","arxiv_id":"2309.05227","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-personalized-user-preference-from","title":"Learning Personalized User Preference from Cold Start in Multi-turn Conversations","date":"2023-09-10","arxiv_id":"2309.05127","n_code_links":0,"syntology":null},{"paper":"/paper/neural-hidden-crf-a-robust-weakly-supervised","slug":"neural-hidden-crf-a-robust-weakly-supervised","title":"Neural-Hidden-CRF: A Robust Weakly-Supervised Sequence Labeler","date":"2023-09-10","arxiv_id":"2309.05086","n_code_links":1,"syntology":null},{"paper":null,"slug":"rgat-a-deeper-look-into-syntactic-dependency","title":"RGAT: A Deeper Look into Syntactic Dependency Information for Coreference Resolution","date":"2023-09-10","arxiv_id":"2309.04977","n_code_links":0,"syntology":null},{"paper":"/paper/encoding-multi-domain-scientific-papers-by","slug":"encoding-multi-domain-scientific-papers-by","title":"Encoding Multi-Domain Scientific Papers by Ensembling Multiple CLS Tokens","date":"2023-09-08","arxiv_id":"2309.04333","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["ronaldseoh/multi2spe"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":"/paper/fuzzy-fingerprinting-transformer-language","slug":"fuzzy-fingerprinting-transformer-language","title":"Fuzzy Fingerprinting Transformer Language-Models for Emotion Recognition in Conversations","date":"2023-09-08","arxiv_id":"2309.04292","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-pretrained-image-text-models-for","title":"Leveraging Pretrained Image-text Models for Improving Audio-Visual Learning","date":"2023-09-08","arxiv_id":"2309.04628","n_code_links":0,"syntology":null},{"paper":"/paper/uq-at-smm4h-2023-alex-for-public-health","slug":"uq-at-smm4h-2023-alex-for-public-health","title":"UQ at #SMM4H 2023: ALEX for Public Health Analysis with Social Media","date":"2023-09-08","arxiv_id":"2309.04213","n_code_links":1,"syntology":null},{"paper":"/paper/certifying-llm-safety-against-adversarial","slug":"certifying-llm-safety-against-adversarial","title":"Certifying LLM Safety against Adversarial Prompting","date":"2023-09-06","arxiv_id":"2309.02705","n_code_links":1,"syntology":null},{"paper":null,"slug":"leave-no-place-behind-improved-geolocation-in","title":"Leave no Place Behind: Improved Geolocation in Humanitarian Documents","date":"2023-09-06","arxiv_id":"2309.02914","n_code_links":0,"syntology":null},{"paper":"/paper/offensive-hebrew-corpus-and-detection-using","slug":"offensive-hebrew-corpus-and-detection-using","title":"Offensive Hebrew Corpus and Detection using BERT","date":"2023-09-06","arxiv_id":"2309.02724","n_code_links":1,"syntology":null},{"paper":null,"slug":"self-supervised-masked-digital-elevation","title":"Self-Supervised Masked Digital Elevation Models Encoding for Low-Resource Downstream Tasks","date":"2023-09-06","arxiv_id":"2309.03367","n_code_links":0,"syntology":null},{"paper":null,"slug":"incorporating-dictionaries-into-a-neural","title":"Incorporating Dictionaries into a Neural Network Architecture to Extract COVID-19 Medical Concepts From Social Media","date":"2023-09-05","arxiv_id":"2309.02188","n_code_links":0,"syntology":null},{"paper":null,"slug":"language-models-for-novelty-detection-in","title":"Language Models for Novelty Detection in System Call Traces","date":"2023-09-05","arxiv_id":"2309.02206","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-bert-language-models-for-multi","title":"Leveraging BERT Language Models for Multi-Lingual ESG Issue Identification","date":"2023-09-05","arxiv_id":"2309.02189","n_code_links":0,"syntology":null},{"paper":null,"slug":"sample-size-in-natural-language-processing","title":"Sample Size in Natural Language Processing within Healthcare Research","date":"2023-09-05","arxiv_id":"2309.02237","n_code_links":0,"syntology":null},{"paper":"/paper/benchmarking-large-language-models-in","slug":"benchmarking-large-language-models-in","title":"Benchmarking Large Language Models in Retrieval-Augmented Generation","date":"2023-09-04","arxiv_id":"2309.01431","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":2,"n_instrument":1,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 1 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["chen700564/RGB"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"a-study-on-the-implementation-of-generative","title":"A Study on the Implementation of Generative AI Services Using an Enterprise Data-Based LLM Application Architecture","date":"2023-09-03","arxiv_id":"2309.01105","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-visual-interpretation-based-self-improved","title":"A Visual Interpretation-Based Self-Improved Classification System Using Virtual Adversarial Training","date":"2023-09-03","arxiv_id":"2309.01196","n_code_links":0,"syntology":null},{"paper":null,"slug":"knowledge-graph-embeddings-for-multi-lingual","title":"Knowledge Graph Embeddings for Multi-Lingual Structured Representations of Radiology Reports","date":"2023-09-02","arxiv_id":"2309.00917","n_code_links":0,"syntology":null},{"paper":null,"slug":"studying-the-impacts-of-pre-training-using","title":"Studying the impacts of pre-training using ChatGPT-generated text on downstream tasks","date":"2023-09-02","arxiv_id":"2309.05668","n_code_links":0,"syntology":null},{"paper":"/paper/batchprompt-accomplish-more-with-less","slug":"batchprompt-accomplish-more-with-less","title":"BatchPrompt: Accomplish more with less","date":"2023-09-01","arxiv_id":"2309.00384","n_code_links":1,"syntology":null},{"paper":null,"slug":"sortednet-a-place-for-every-network-and-every","title":"SortedNet: A Scalable and Generalized Framework for Training Modular Deep Neural Networks","date":"2023-09-01","arxiv_id":"2309.00255","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-humans-help-bert-gain-confidence","title":"Can humans help BERT gain \"confidence\"?","date":"2023-08-31","arxiv_id":"2309.06580","n_code_links":0,"syntology":null},{"paper":null,"slug":"dictabert-a-state-of-the-art-bert-suite-for","title":"DictaBERT: A State-of-the-Art BERT Suite for Modern Hebrew","date":"2023-08-31","arxiv_id":"2308.16687","n_code_links":0,"syntology":null},{"paper":null,"slug":"linking-microblogging-sentiments-to-stock","title":"Linking microblogging sentiments to stock price movement: An application of GPT-4","date":"2023-08-31","arxiv_id":"2308.16771","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-improving-the-expressiveness-of","title":"Towards Improving the Expressiveness of Singing Voice Synthesis with BERT Derived Semantic Information","date":"2023-08-31","arxiv_id":"2308.16836","n_code_links":0,"syntology":null},{"paper":"/paper/transformer-compression-via-subspace","slug":"transformer-compression-via-subspace","title":"$\\rm SP^3$: Enhancing Structured Pruning via PCA Projection","date":"2023-08-31","arxiv_id":"2308.16475","n_code_links":1,"syntology":{"ran":12,"of":14,"n_ran_checked":11,"n_instrument":1,"unverified":2,"pointer_only":14,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 1 honoured, 1 violated, 9 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["hyx1999/sp3"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"analyzing-character-and-consciousness-in-ai","title":"Analyzing Character and Consciousness in AI-Generated Social Content: A Case Study of Chirper, the AI Social Network","date":"2023-08-30","arxiv_id":"2309.08614","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-deepzen-speech-synthesis-system-for","title":"The DeepZen Speech Synthesis System for Blizzard Challenge 2023","date":"2023-08-30","arxiv_id":"2308.15945","n_code_links":0,"syntology":null},{"paper":"/paper/spikebert-a-language-spikformer-trained-with","slug":"spikebert-a-language-spikformer-trained-with","title":"SpikeBERT: A Language Spikformer Learned from BERT with Knowledge Distillation","date":"2023-08-29","arxiv_id":"2308.15122","n_code_links":1,"syntology":{"ran":8,"of":9,"n_ran_checked":8,"n_instrument":0,"unverified":1,"pointer_only":9,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["Lvchangze/SpikeBERT"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/aner-arabic-and-arabizi-named-entity","slug":"aner-arabic-and-arabizi-named-entity","title":"ANER: Arabic and Arabizi Named Entity Recognition using Transformer-Based Approach","date":"2023-08-28","arxiv_id":"2308.14669","n_code_links":1,"syntology":null},{"paper":null,"slug":"target-independent-xla-optimization-using","title":"Target-independent XLA optimization using Reinforcement Learning","date":"2023-08-28","arxiv_id":"2308.14364","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-knowledge-distillation-for-bert","title":"Improving Knowledge Distillation for BERT Models: Loss Functions, Mapping Methods, and Weight Tuning","date":"2023-08-26","arxiv_id":"2308.13958","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-knowledge-and-reinforcement","title":"Leveraging Knowledge and Reinforcement Learning for Enhanced Reliability of Language Models","date":"2023-08-25","arxiv_id":"2308.13467","n_code_links":0,"syntology":null},{"paper":"/paper/a-small-and-fast-bert-for-chinese-medical","slug":"a-small-and-fast-bert-for-chinese-medical","title":"A Small and Fast BERT for Chinese Medical Punctuation Restoration","date":"2023-08-24","arxiv_id":"2308.12568","n_code_links":1,"syntology":null},{"paper":"/paper/advancing-hungarian-text-processing-with","slug":"advancing-hungarian-text-processing-with","title":"Advancing Hungarian Text Processing with HuSpaCy: Efficient and Accurate NLP Pipelines","date":"2023-08-24","arxiv_id":"2308.12635","n_code_links":2,"syntology":null},{"paper":null,"slug":"multi-bert-for-embeddings-for-recommendation","title":"Multi-BERT for Embeddings for Recommendation System","date":"2023-08-24","arxiv_id":"2308.13050","n_code_links":0,"syntology":null},{"paper":"/paper/sentence-embedding-models-for-ancient-greek","slug":"sentence-embedding-models-for-ancient-greek","title":"Sentence Embedding Models for Ancient Greek Using Multilingual Knowledge Distillation","date":"2023-08-24","arxiv_id":"2308.13116","n_code_links":2,"syntology":null},{"paper":null,"slug":"text-similarity-from-image-contents-using","title":"Text Similarity from Image Contents using Statistical and Semantic Analysis Techniques","date":"2023-08-24","arxiv_id":"2308.12842","n_code_links":0,"syntology":null},{"paper":null,"slug":"simple-is-better-and-large-is-not-enough","title":"Simple is Better and Large is Not Enough: Towards Ensembling of Foundational Language Models","date":"2023-08-23","arxiv_id":"2308.12272","n_code_links":0,"syntology":null},{"paper":"/paper/spikingbert-distilling-bert-to-train-spiking","slug":"spikingbert-distilling-bert-to-train-spiking","title":"SpikingBERT: Distilling BERT to Train Spiking Language Models Using Implicit Differentiation","date":"2023-08-21","arxiv_id":"2308.10873","n_code_links":1,"syntology":{"ran":13,"of":17,"n_ran_checked":8,"n_instrument":5,"unverified":4,"pointer_only":17,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 2 honoured, 0 violated, 6 with no contract checked; 5 where Syntology's instrument failed) · 4 unverified","official":{"repos":["neurocomplab-psu/spikingbert"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/how-good-are-large-language-models-at-out-of","slug":"how-good-are-large-language-models-at-out-of","title":"How Good Are LLMs at Out-of-Distribution Detection?","date":"2023-08-20","arxiv_id":"2308.10261","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["awenbocc/llm-ood"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/improving-adversarial-robustness-of-masked","slug":"improving-adversarial-robustness-of-masked","title":"Improving Adversarial Robustness of Masked Autoencoders via Test-time Frequency-domain Prompting","date":"2023-08-20","arxiv_id":"2308.10315","n_code_links":1,"syntology":null},{"paper":null,"slug":"east-efficient-and-accurate-secure","title":"East: Efficient and Accurate Secure Transformer Framework for Inference","date":"2023-08-19","arxiv_id":"2308.09923","n_code_links":0,"syntology":null},{"paper":null,"slug":"open-closed-or-small-language-models-for-text","title":"Open, Closed, or Small Language Models for Text Classification?","date":"2023-08-19","arxiv_id":"2308.10092","n_code_links":0,"syntology":null}],"record_sha256":"3aa8be970793d212f5ec9c74c932e004e3b59413ddaabb868d4e10f54e861244","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}