{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/linear-layer/papers/112","list_of":"/method/linear-layer","method":"Linear Layer","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":112,"pages_in_order":255,"rows_per_page":100,"rows":[11101,11200],"of":25421,"counts":{"archive_papers_tagged":25421,"with_a_code_link":11479,"where_syntology_ran_a_sample":3523,"not_listed_spam_title":0,"listed":25421,"listed_where_code_ran":3523,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":2976,"every_run_a_failure_of_syntologys_instrument":547,"listed_with_a_run_with_no_instrument_failure":2976,"listed_every_run_a_failure_of_syntologys_instrument":547,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/linear-layer","prev":"/method/linear-layer/papers/111","next":"/method/linear-layer/papers/113","papers":[{"paper":null,"slug":"utilising-a-large-language-model-to-annotate","title":"Utilising a Large Language Model to Annotate Subject Metadata: A Case Study in an Australian National Research Data Catalogue","date":"2023-10-17","arxiv_id":"2310.11318","n_code_links":0,"syntology":null},{"paper":"/paper/vct-visual-change-transformer-for-remote","slug":"vct-visual-change-transformer-for-remote","title":"VcT: Visual change Transformer for Remote Sensing Image Change Detection","date":"2023-10-17","arxiv_id":"2310.11417","n_code_links":1,"syntology":null},{"paper":null,"slug":"what-is-a-good-question-task-oriented-asking","title":"Alexpaca: Learning Factual Clarification Question Generation Without Examples","date":"2023-10-17","arxiv_id":"2310.11571","n_code_links":0,"syntology":null},{"paper":"/paper/a-cross-transformer-for-image-denoising","slug":"a-cross-transformer-for-image-denoising","title":"A cross Transformer for image denoising","date":"2023-10-16","arxiv_id":"2310.10408","n_code_links":1,"syntology":null},{"paper":"/paper/approximating-two-layer-feedforward-networks","slug":"approximating-two-layer-feedforward-networks","title":"Approximating Two-Layer Feedforward Networks for Efficient Transformers","date":"2023-10-16","arxiv_id":"2310.10837","n_code_links":2,"syntology":{"ran":3,"of":4,"n_ran_checked":1,"n_instrument":2,"unverified":1,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["robertcsordas/moe","robertcsordas/moe_layer"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"battle-of-the-large-language-models-dolly-vs","title":"Battle of the Large Language Models: Dolly vs LLaMA vs Vicuna vs Guanaco vs Bard vs ChatGPT -- A Text-to-SQL Parsing Comparison","date":"2023-10-16","arxiv_id":"2310.10190","n_code_links":0,"syntology":null},{"paper":null,"slug":"biomedjourney-counterfactual-biomedical-image","title":"BiomedJourney: Counterfactual Biomedical Image Generation by Instruction-Learning from Multimodal Patient Journeys","date":"2023-10-16","arxiv_id":"2310.10765","n_code_links":0,"syntology":null},{"paper":"/paper/bioplanner-automatic-evaluation-of-llms-on","slug":"bioplanner-automatic-evaluation-of-llms-on","title":"BioPlanner: Automatic Evaluation of LLMs on Protocol Planning in Biology","date":"2023-10-16","arxiv_id":"2310.10632","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["bioplanner/bioplanner"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/data-contamination-through-the-lens-of-time","slug":"data-contamination-through-the-lens-of-time","title":"Data Contamination Through the Lens of Time","date":"2023-10-16","arxiv_id":"2310.10628","n_code_links":1,"syntology":{"ran":2,"of":4,"n_ran_checked":2,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["abacusai/to-the-cutoff"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/factored-verification-detecting-and-reducing","slug":"factored-verification-detecting-and-reducing","title":"Factored Verification: Detecting and Reducing Hallucination in Summaries of Academic Papers","date":"2023-10-16","arxiv_id":"2310.10627","n_code_links":1,"syntology":null},{"paper":null,"slug":"few-shot-action-recognition-with-captioning","title":"Few-shot Action Recognition with Captioning Foundation Models","date":"2023-10-16","arxiv_id":"2310.10125","n_code_links":0,"syntology":null},{"paper":null,"slug":"fine-tuning-chatgpt-for-automatic-scoring","title":"Fine-tuning ChatGPT for Automatic Scoring","date":"2023-10-16","arxiv_id":"2310.10072","n_code_links":0,"syntology":null},{"paper":"/paper/interpreting-and-exploiting-functional","slug":"interpreting-and-exploiting-functional","title":"Interpreting and Exploiting Functional Specialization in Multi-Head Attention under Multi-task Learning","date":"2023-10-16","arxiv_id":"2310.10318","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["znlp/functionalspecializationinmha"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":"/paper/investigating-bias-in-multilingual-language","slug":"investigating-bias-in-multilingual-language","title":"Investigating Bias in Multilingual Language Models: Cross-Lingual Transfer of Debiasing Techniques","date":"2023-10-16","arxiv_id":"2310.10310","n_code_links":1,"syntology":null},{"paper":"/paper/learning-to-rank-context-for-named-entity","slug":"learning-to-rank-context-for-named-entity","title":"Learning to Rank Context for Named Entity Recognition Using a Synthetic Dataset","date":"2023-10-16","arxiv_id":"2310.10118","n_code_links":1,"syntology":null},{"paper":null,"slug":"mask-wearing-object-detection-algorithm-based","title":"Mask wearing object detection algorithm based on improved YOLOv5","date":"2023-10-16","arxiv_id":"2310.10245","n_code_links":0,"syntology":null},{"paper":null,"slug":"moconvq-unified-physics-based-motion-control","title":"MoConVQ: Unified Physics-Based Motion Control via Scalable Discrete Representations","date":"2023-10-16","arxiv_id":"2310.10198","n_code_links":0,"syntology":null},{"paper":"/paper/on-the-relevance-of-temporal-features-for","slug":"on-the-relevance-of-temporal-features-for","title":"On the Relevance of Temporal Features for Medical Ultrasound Video Recognition","date":"2023-10-16","arxiv_id":"2310.10453","n_code_links":1,"syntology":null},{"paper":"/paper/pela-learning-parameter-efficient-models-with","slug":"pela-learning-parameter-efficient-models-with","title":"PELA: Learning Parameter-Efficient Models with Low-Rank Approximation","date":"2023-10-16","arxiv_id":"2310.10700","n_code_links":1,"syntology":null},{"paper":null,"slug":"prediction-of-arabic-legal-rulings-using","title":"Prediction of Arabic Legal Rulings using Large Language Models","date":"2023-10-16","arxiv_id":"2310.10260","n_code_links":0,"syntology":null},{"paper":null,"slug":"prompt-packer-deceiving-llms-through","title":"Prompt Packer: Deceiving LLMs through Compositional Instruction with Hidden Attacks","date":"2023-10-16","arxiv_id":"2310.10077","n_code_links":0,"syntology":null},{"paper":null,"slug":"seunet-trans-a-simple-yet-effective-unet","title":"SeUNet-Trans: A Simple yet Effective UNet-Transformer Model for Medical Image Segmentation","date":"2023-10-16","arxiv_id":"2310.09998","n_code_links":0,"syntology":null},{"paper":"/paper/transom-an-efficient-fault-tolerant-system","slug":"transom-an-efficient-fault-tolerant-system","title":"TRANSOM: An Efficient Fault-Tolerant System for Training LLMs","date":"2023-10-16","arxiv_id":"2310.10046","n_code_links":1,"syntology":{"ran":0,"of":3,"n_ran_checked":0,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"0 ran · 3 unverified","official":{"repos":["SenseCore/transom-checkpoint-engine"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":3,"ran_from_kinds":[]}}},{"paper":"/paper/trigo-benchmarking-formal-mathematical-proof","slug":"trigo-benchmarking-formal-mathematical-proof","title":"TRIGO: Benchmarking Formal Mathematical Proof Reduction for Generative Language Models","date":"2023-10-16","arxiv_id":"2310.10180","n_code_links":1,"syntology":{"ran":7,"of":11,"n_ran_checked":7,"n_instrument":0,"unverified":4,"pointer_only":11,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["menik1126/TRIGO"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"verbosity-bias-in-preference-labeling-by","title":"Verbosity Bias in Preference Labeling by Large Language Models","date":"2023-10-16","arxiv_id":"2310.10076","n_code_links":0,"syntology":null},{"paper":"/paper/cocoformer-a-controllable-feature-rich","slug":"cocoformer-a-controllable-feature-rich","title":"CoCoFormer: A controllable feature-rich polyphonic music generation method","date":"2023-10-15","arxiv_id":"2310.09843","n_code_links":1,"syntology":null},{"paper":null,"slug":"configuration-validation-with-large-language","title":"Configuration Validation with Large Language Models","date":"2023-10-15","arxiv_id":"2310.09690","n_code_links":0,"syntology":null},{"paper":null,"slug":"diversifying-the-mixture-of-experts","title":"Diversifying the Mixture-of-Experts Representation for Language Models with Orthogonal Optimizer","date":"2023-10-15","arxiv_id":"2310.09762","n_code_links":0,"syntology":null},{"paper":"/paper/domain-specific-language-model-post-training","slug":"domain-specific-language-model-post-training","title":"Domain-Specific Language Model Post-Training for Indonesian Financial NLP","date":"2023-10-15","arxiv_id":"2310.09736","n_code_links":1,"syntology":null},{"paper":null,"slug":"empirical-study-of-pretrained-multilingual","title":"Empirical study of pretrained multilingual language models for zero-shot cross-lingual knowledge transfer in generation","date":"2023-10-15","arxiv_id":"2310.09917","n_code_links":0,"syntology":null},{"paper":null,"slug":"image-augmentation-with-controlled-diffusion","title":"Image Augmentation with Controlled Diffusion for Weakly-Supervised Semantic Segmentation","date":"2023-10-15","arxiv_id":"2310.09760","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-model-aware-in-context","title":"Large Language Model-Aware In-Context Learning for Code Generation","date":"2023-10-15","arxiv_id":"2310.09748","n_code_links":0,"syntology":null},{"paper":"/paper/large-language-models-for-in-context-student","slug":"large-language-models-for-in-context-student","title":"Large Language Models for In-Context Student Modeling: Synthesizing Student's Behavior in Visual Programming","date":"2023-10-15","arxiv_id":"2310.10690","n_code_links":1,"syntology":null},{"paper":"/paper/moemo-vision-transformer-integrating-cross","slug":"moemo-vision-transformer-integrating-cross","title":"MoEmo Vision Transformer: Integrating Cross-Attention and Movement Vectors in 3D Pose Estimation for HRI Emotion Detection","date":"2023-10-15","arxiv_id":"2310.09757","n_code_links":1,"syntology":null},{"paper":"/paper/spatio-temporal-graph-mixformer-for-traffic","slug":"spatio-temporal-graph-mixformer-for-traffic","title":"Spatio-Temporal Graph Mixformer for Traffic Forecasting","date":"2023-10-15","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"top-k-pooling-with-patch-contrastive-learning","title":"Top-K Pooling with Patch Contrastive Learning for Weakly-Supervised Semantic Segmentation","date":"2023-10-15","arxiv_id":"2310.09828","n_code_links":0,"syntology":null},{"paper":"/paper/unitime-a-language-empowered-unified-model","slug":"unitime-a-language-empowered-unified-model","title":"UniTime: A Language-Empowered Unified Model for Cross-Domain Time Series Forecasting","date":"2023-10-15","arxiv_id":"2310.09751","n_code_links":1,"syntology":{"ran":4,"of":8,"n_ran_checked":4,"n_instrument":0,"unverified":4,"pointer_only":4,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["liuxu77/unitime"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"arm-refining-multivariate-forecasting-with","title":"ARM: Refining Multivariate Forecasting with Adaptive Temporal-Contextual Learning","date":"2023-10-14","arxiv_id":"2310.09488","n_code_links":0,"syntology":null},{"paper":"/paper/dpzero-dimension-independent-and","slug":"dpzero-dimension-independent-and","title":"DPZero: Private Fine-Tuning of Language Models without Backpropagation","date":"2023-10-14","arxiv_id":"2310.09639","n_code_links":1,"syntology":{"ran":8,"of":17,"n_ran_checked":6,"n_instrument":2,"unverified":9,"pointer_only":3,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 2 where Syntology's instrument failed) · 9 unverified","official":{"repos":["liang137/dpzero"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":9,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"efficient-model-agnostic-multi-group","title":"Efficient Model-Agnostic Multi-Group Equivariant Networks","date":"2023-10-14","arxiv_id":"2310.09675","n_code_links":0,"syntology":null},{"paper":"/paper/instruction-tuning-with-human-curriculum","slug":"instruction-tuning-with-human-curriculum","title":"Instruction Tuning with Human Curriculum","date":"2023-10-14","arxiv_id":"2310.09518","n_code_links":1,"syntology":null},{"paper":null,"slug":"leveraging-generative-ai-improving-software","title":"Leveraging Generative AI: Improving Software Metadata Classification with Generated Code-Comment Pairs","date":"2023-10-14","arxiv_id":"2311.03365","n_code_links":0,"syntology":null},{"paper":null,"slug":"topology-guided-hypergraph-transformer","title":"Topology-guided Hypergraph Transformer Network: Unveiling Structural Insights for Improved Representation","date":"2023-10-14","arxiv_id":"2310.09657","n_code_links":0,"syntology":null},{"paper":"/paper/uniqa-a-unified-framework-for-both-full","slug":"uniqa-a-unified-framework-for-both-full","title":"You Only Train Once: A Unified Framework for Both Full-Reference and No-Reference Image Quality Assessment","date":"2023-10-14","arxiv_id":"2310.09560","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-comparative-analysis-of-task-agnostic","title":"A Comparative Analysis of Task-Agnostic Distillation Methods for Compressing Transformer Language Models","date":"2023-10-13","arxiv_id":"2310.08797","n_code_links":0,"syntology":null},{"paper":"/paper/a-systematic-evaluation-of-large-language-1","slug":"a-systematic-evaluation-of-large-language-1","title":"Assessing and Enhancing the Robustness of Large Language Models with Task Structure Variations for Logical Reasoning","date":"2023-10-13","arxiv_id":"2310.09430","n_code_links":1,"syntology":{"ran":6,"of":8,"n_ran_checked":6,"n_instrument":0,"unverified":2,"pointer_only":1,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["strong-ai-lab/logical-and-abstract-reasoning"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"automated-claim-matching-with-large-language","title":"Automated Claim Matching with Large Language Models: Empowering Fact-Checkers in the Fight Against Misinformation","date":"2023-10-13","arxiv_id":"2310.09223","n_code_links":0,"syntology":null},{"paper":null,"slug":"ddmt-denoising-diffusion-mask-transformer","title":"DDMT: Denoising Diffusion Mask Transformer Models for Multivariate Time Series Anomaly Detection","date":"2023-10-13","arxiv_id":"2310.08800","n_code_links":0,"syntology":null},{"paper":null,"slug":"distance-rank-aware-sequential-reward","title":"Distance-rank Aware Sequential Reward Learning for Inverse Reinforcement Learning with Sub-optimal Demonstrations","date":"2023-10-13","arxiv_id":"2310.08823","n_code_links":0,"syntology":null},{"paper":"/paper/dont-add-dont-miss-effective-content","slug":"dont-add-dont-miss-effective-content","title":"Dont Add, dont Miss: Effective Content Preserving Generation from Pre-Selected Text Spans","date":"2023-10-13","arxiv_id":"2310.09017","n_code_links":1,"syntology":{"ran":4,"of":5,"n_ran_checked":4,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["lovodkin93/cdr_ctr"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"enhancing-bert-based-visual-question","title":"Enhancing BERT-Based Visual Question Answering through Keyword-Driven Sentence Selection","date":"2023-10-13","arxiv_id":"2310.09432","n_code_links":0,"syntology":null},{"paper":"/paper/faster-3d-cardiac-ct-segmentation-with-vision","slug":"faster-3d-cardiac-ct-segmentation-with-vision","title":"Vision Transformers increase efficiency of 3D cardiac CT multi-label segmentation","date":"2023-10-13","arxiv_id":"2310.09099","n_code_links":1,"syntology":null},{"paper":"/paper/from-clip-to-dino-visual-encoders-shout-in","slug":"from-clip-to-dino-visual-encoders-shout-in","title":"From CLIP to DINO: Visual Encoders Shout in Multi-modal Large Language Models","date":"2023-10-13","arxiv_id":"2310.08825","n_code_links":1,"syntology":null},{"paper":null,"slug":"from-words-and-exercises-to-wellness-farsi","title":"From Words and Exercises to Wellness: Farsi Chatbot for Self-Attachment Technique","date":"2023-10-13","arxiv_id":"2310.09362","n_code_links":0,"syntology":null},{"paper":"/paper/glore-evaluating-logical-reasoning-of-large","slug":"glore-evaluating-logical-reasoning-of-large","title":"GLoRE: Evaluating Logical Reasoning of Large Language Models","date":"2023-10-13","arxiv_id":"2310.09107","n_code_links":1,"syntology":null},{"paper":"/paper/human-in-the-loop-machine-translation-with","slug":"human-in-the-loop-machine-translation-with","title":"Human-in-the-loop Machine Translation with Large Language Model","date":"2023-10-13","arxiv_id":"2310.08908","n_code_links":1,"syntology":null},{"paper":"/paper/pali-3-vision-language-models-smaller-faster","slug":"pali-3-vision-language-models-smaller-faster","title":"PaLI-3 Vision Language Models: Smaller, Faster, Stronger","date":"2023-10-13","arxiv_id":"2310.09199","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 3 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":"/paper/retrieval-generation-alignment-for-end-to-end","slug":"retrieval-generation-alignment-for-end-to-end","title":"Retrieval-Generation Alignment for End-to-End Task-Oriented Dialogue System","date":"2023-10-13","arxiv_id":"2310.08877","n_code_links":1,"syntology":null},{"paper":null,"slug":"table-gpt-table-tuned-gpt-for-diverse-table","title":"Table-GPT: Table-tuned GPT for Diverse Table Tasks","date":"2023-10-13","arxiv_id":"2310.09263","n_code_links":0,"syntology":null},{"paper":"/paper/towards-end-to-end-4-bit-inference-on","slug":"towards-end-to-end-4-bit-inference-on","title":"QUIK: Towards End-to-End 4-Bit Inference on Generative Large Language Models","date":"2023-10-13","arxiv_id":"2310.09259","n_code_links":1,"syntology":{"ran":9,"of":13,"n_ran_checked":9,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["ist-daslab/quik"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/towards-example-based-nmt-with-multi","slug":"towards-example-based-nmt-with-multi","title":"Towards Example-Based NMT with Multi-Levenshtein Transformers","date":"2023-10-13","arxiv_id":"2310.08967","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["maxwell1447/fairseq"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"age-estimation-based-on-graph-convolutional","title":"Age Estimation Based on Graph Convolutional Networks and Multi-head Attention Mechanisms","date":"2023-10-12","arxiv_id":"2310.08064","n_code_links":0,"syntology":null},{"paper":null,"slug":"analyzing-textual-data-for-fatality","title":"Analyzing Textual Data for Fatality Classification in Afghanistan's Armed Conflicts: A BERT Approach","date":"2023-10-12","arxiv_id":"2310.08653","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-gpt-models-be-financial-analysts-an","title":"Can GPT models be Financial Analysts? An Evaluation of ChatGPT and GPT-4 on mock CFA Exams","date":"2023-10-12","arxiv_id":"2310.08678","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-large-language-models-really-improve-by","title":"Can Large Language Models Really Improve by Self-critiquing Their Own Plans?","date":"2023-10-12","arxiv_id":"2310.08118","n_code_links":0,"syntology":null},{"paper":"/paper/covid-19-detection-using-swin-transformer","slug":"covid-19-detection-using-swin-transformer","title":"COVID-19 detection using ViT transformer-based approach from Computed Tomography Images","date":"2023-10-12","arxiv_id":"2310.08165","n_code_links":1,"syntology":null},{"paper":null,"slug":"cross-episodic-curriculum-for-transformer","title":"Cross-Episodic Curriculum for Transformer Agents","date":"2023-10-12","arxiv_id":"2310.08549","n_code_links":0,"syntology":null},{"paper":null,"slug":"detection-and-prediction-of-clopidogrel","title":"Detection and prediction of clopidogrel treatment failures using longitudinal structured electronic health records","date":"2023-10-12","arxiv_id":"2310.08757","n_code_links":0,"syntology":null},{"paper":null,"slug":"do-pretrained-transformers-really-learn-in","title":"Do pretrained Transformers Learn In-Context by Gradient Descent?","date":"2023-10-12","arxiv_id":"2310.08540","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-the-effectiveness-of-capsule","slug":"evaluating-the-effectiveness-of-capsule","title":"Evaluating The Effectiveness of Capsule Neural Network in Toxic Comment Classification using Pre-trained BERT Embeddings","date":"2023-10-12","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/expanding-the-vocabulary-of-bert-for","slug":"expanding-the-vocabulary-of-bert-for","title":"Expanding the Vocabulary of BERT for Knowledge Base Construction","date":"2023-10-12","arxiv_id":"2310.08291","n_code_links":1,"syntology":null},{"paper":"/paper/impact-of-time-and-note-duration","slug":"impact-of-time-and-note-duration","title":"Impact of time and note duration tokenizations on deep learning symbolic music modeling","date":"2023-10-12","arxiv_id":"2310.08497","n_code_links":1,"syntology":null},{"paper":"/paper/interpreting-reward-models-in-rlhf-tuned","slug":"interpreting-reward-models-in-rlhf-tuned","title":"Interpreting Learned Feedback Patterns in Large Language Models","date":"2023-10-12","arxiv_id":"2310.08164","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":0,"n_instrument":1,"unverified":1,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["apartresearch/interpreting-reward-models"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"investigating-the-robustness-and-properties","title":"Investigating the Robustness and Properties of Detection Transformers (DETR) Toward Difficult Images","date":"2023-10-12","arxiv_id":"2310.08772","n_code_links":0,"syntology":null},{"paper":"/paper/jailbreaking-black-box-large-language-models","slug":"jailbreaking-black-box-large-language-models","title":"Jailbreaking Black Box Large Language Models in Twenty Queries","date":"2023-10-12","arxiv_id":"2310.08419","n_code_links":1,"syntology":{"ran":2,"of":7,"n_ran_checked":1,"n_instrument":1,"unverified":5,"pointer_only":0,"phrase":"2 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 5 unverified","official":{"repos":["patrickrchao/jailbreakingllms"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"large-language-models-can-replicate-cross","title":"Large language models can replicate cross-cultural differences in personality","date":"2023-10-12","arxiv_id":"2310.10679","n_code_links":0,"syntology":null},{"paper":"/paper/lemon-lossless-model-expansion","slug":"lemon-lossless-model-expansion","title":"LEMON: Lossless model expansion","date":"2023-10-12","arxiv_id":"2310.07999","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["YiteWang/lemon-pytorch"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"llm-augmented-preference-learning-from","title":"LLM-augmented Preference Learning from Natural Language","date":"2023-10-12","arxiv_id":"2310.08523","n_code_links":0,"syntology":null},{"paper":null,"slug":"multiclass-classification-of-policy-documents","title":"Multiclass Classification of Policy Documents with Large Language Models","date":"2023-10-12","arxiv_id":"2310.08167","n_code_links":0,"syntology":null},{"paper":"/paper/octopus-embodied-vision-language-programmer","slug":"octopus-embodied-vision-language-programmer","title":"Octopus: Embodied Vision-Language Programmer from Environmental Feedback","date":"2023-10-12","arxiv_id":"2310.08588","n_code_links":1,"syntology":{"ran":12,"of":12,"n_ran_checked":6,"n_instrument":6,"unverified":0,"pointer_only":12,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 1 honoured, 1 violated, 4 with no contract checked; 6 where Syntology's instrument failed) · 0 unverified","official":{"repos":["dongyh20/octopus"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/prometheus-inducing-fine-grained-evaluation","slug":"prometheus-inducing-fine-grained-evaluation","title":"Prometheus: Inducing Fine-grained Evaluation Capability in Language Models","date":"2023-10-12","arxiv_id":"2310.08491","n_code_links":3,"syntology":{"ran":4,"of":9,"n_ran_checked":3,"n_instrument":1,"unverified":5,"pointer_only":5,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 1 where Syntology's instrument failed) · 5 unverified","official":{"repos":["kaistAI/Prometheus"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"promptor-a-conversational-and-autonomous","title":"Promptor: A Conversational and Autonomous Prompt Generation Agent for Intelligent Text Entry Techniques","date":"2023-10-12","arxiv_id":"2310.08101","n_code_links":0,"syntology":null},{"paper":"/paper/qasina-religious-domain-question-answering","slug":"qasina-religious-domain-question-answering","title":"QASiNa: Religious Domain Question Answering using Sirah Nabawiyah","date":"2023-10-12","arxiv_id":"2310.08102","n_code_links":1,"syntology":null},{"paper":null,"slug":"revisiting-data-augmentation-for-rotational","title":"Revisiting Data Augmentation for Rotational Invariance in Convolutional Neural Networks","date":"2023-10-12","arxiv_id":"2310.08429","n_code_links":0,"syntology":null},{"paper":null,"slug":"self-supervised-visual-learning-for-analyzing","title":"Self-supervised visual learning for analyzing firearms trafficking activities on the Web","date":"2023-10-12","arxiv_id":"2310.07975","n_code_links":0,"syntology":null},{"paper":"/paper/the-uncertainty-based-retrieval-framework-for-1","slug":"the-uncertainty-based-retrieval-framework-for-1","title":"The Uncertainty-based Retrieval Framework for Ancient Chinese CWS and POS","date":"2023-10-12","arxiv_id":"2310.08496","n_code_links":1,"syntology":null},{"paper":null,"slug":"training-generative-question-answering-on","title":"Training Generative Question-Answering on Synthetic Data Obtained from an Instruct-tuned Model","date":"2023-10-12","arxiv_id":"2310.08072","n_code_links":0,"syntology":null},{"paper":null,"slug":"transformer-choice-net-a-transformer-neural","title":"Transformer Choice Net: A Transformer Neural Network for Choice Prediction","date":"2023-10-12","arxiv_id":"2310.08716","n_code_links":0,"syntology":null},{"paper":"/paper/transport-hub-aware-spatial-temporal-adaptive","slug":"transport-hub-aware-spatial-temporal-adaptive","title":"Transport-Hub-Aware Spatial-Temporal Adaptive Graph Transformer for Traffic Flow Prediction","date":"2023-10-12","arxiv_id":"2310.08328","n_code_links":1,"syntology":null},{"paper":null,"slug":"ziya-vl-bilingual-large-vision-language-model","title":"Ziya-Visual: Bilingual Large Vision-Language Model via Multi-Task Instruction Tuning","date":"2023-10-12","arxiv_id":"2310.08166","n_code_links":0,"syntology":null},{"paper":"/paper/3d-transunet-advancing-medical-image","slug":"3d-transunet-advancing-medical-image","title":"3D TransUNet: Advancing Medical Image Segmentation through Vision Transformers","date":"2023-10-11","arxiv_id":"2310.07781","n_code_links":3,"syntology":null},{"paper":null,"slug":"accelerating-vision-transformers-based-on","title":"Accelerating Vision Transformers Based on Heterogeneous Attention Patterns","date":"2023-10-11","arxiv_id":"2310.07664","n_code_links":0,"syntology":null},{"paper":null,"slug":"atom-motif-contrastive-transformer-for","title":"Atom-Motif Contrastive Transformer for Molecular Property Prediction","date":"2023-10-11","arxiv_id":"2310.07351","n_code_links":0,"syntology":null},{"paper":"/paper/cognate-transformer-for-automated","slug":"cognate-transformer-for-automated","title":"Cognate Transformer for Automated Phonological Reconstruction and Cognate Reflex Prediction","date":"2023-10-11","arxiv_id":"2310.07487","n_code_links":1,"syntology":null},{"paper":"/paper/daspeech-directed-acyclic-transformer-for-1","slug":"daspeech-directed-acyclic-transformer-for-1","title":"DASpeech: Directed Acyclic Transformer for Fast and High-quality Speech-to-Speech Translation","date":"2023-10-11","arxiv_id":"2310.07403","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":5,"n_instrument":0,"unverified":2,"pointer_only":7,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["ictnlp/daspeech"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"distance-based-weighted-transformer-network","title":"Distance Weighted Trans Network for Image Completion","date":"2023-10-11","arxiv_id":"2310.07440","n_code_links":0,"syntology":null},{"paper":null,"slug":"distilling-efficient-vision-transformers-from","title":"Distilling Efficient Vision Transformers from CNNs for Semantic Segmentation","date":"2023-10-11","arxiv_id":"2310.07265","n_code_links":0,"syntology":null},{"paper":null,"slug":"diversity-of-thought-improves-reasoning","title":"Diversity of Thought Improves Reasoning Abilities of LLMs","date":"2023-10-11","arxiv_id":"2310.07088","n_code_links":0,"syntology":null},{"paper":null,"slug":"does-resistance-to-style-transfer-equal-shape","title":"Does resistance to style-transfer equal Global Shape Bias? Measuring network sensitivity to global shape configuration","date":"2023-10-11","arxiv_id":"2310.07555","n_code_links":0,"syntology":null},{"paper":null,"slug":"ethical-reasoning-over-moral-alignment-a-case","title":"Ethical Reasoning over Moral Alignment: A Case and Framework for In-Context Ethical Policies in LLMs","date":"2023-10-11","arxiv_id":"2310.07251","n_code_links":0,"syntology":null}],"record_sha256":"42119a4f9d1665eced768db78604772b9708d97762707d70b8b0e9d1d824b6f1","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}