{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/attention/papers/176","list_of":"/method/attention","method":"Attention","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":176,"pages_in_order":316,"rows_per_page":100,"rows":[17501,17600],"of":31583,"counts":{"archive_papers_tagged":31583,"with_a_code_link":13473,"where_syntology_ran_a_sample":3998,"not_listed_spam_title":0,"listed":31583,"listed_where_code_ran":3998,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3366,"every_run_a_failure_of_syntologys_instrument":632,"listed_with_a_run_with_no_instrument_failure":3366,"listed_every_run_a_failure_of_syntologys_instrument":632,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/attention","prev":"/method/attention/papers/175","next":"/method/attention/papers/177","papers":[{"paper":"/paper/once-upon-a-textit-time-in-textit-graph","slug":"once-upon-a-textit-time-in-textit-graph","title":"Once Upon a $\\textit{Time}$ in $\\textit{Graph}$: Relative-Time Pretraining for Complex Temporal Reasoning","date":"2023-10-23","arxiv_id":"2310.14709","n_code_links":1,"syntology":{"ran":7,"of":11,"n_ran_checked":7,"n_instrument":0,"unverified":4,"pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["damo-nlp-sg/rememo"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/p2at-pyramid-pooling-axial-transformer-for","slug":"p2at-pyramid-pooling-axial-transformer-for","title":"P2AT: Pyramid Pooling Axial Transformer for Real-time Semantic Segmentation","date":"2023-10-23","arxiv_id":"2310.15025","n_code_links":1,"syntology":null},{"paper":"/paper/partialformer-modeling-part-instead-of-whole","slug":"partialformer-modeling-part-instead-of-whole","title":"PartialFormer: Modeling Part Instead of Whole for Machine Translation","date":"2023-10-23","arxiv_id":"2310.14921","n_code_links":1,"syntology":null},{"paper":null,"slug":"prefix-tuning-based-unsupervised-text-style","title":"Prefix-Tuning Based Unsupervised Text Style Transfer","date":"2023-10-23","arxiv_id":"2310.14599","n_code_links":0,"syntology":null},{"paper":null,"slug":"samclr-contrastive-pre-training-on-complex","title":"SAMCLR: Contrastive pre-training on complex scenes using SAM for view sampling","date":"2023-10-23","arxiv_id":"2310.14736","n_code_links":0,"syntology":null},{"paper":"/paper/slog-a-structural-generalization-benchmark","slug":"slog-a-structural-generalization-benchmark","title":"SLOG: A Structural Generalization Benchmark for Semantic Parsing","date":"2023-10-23","arxiv_id":"2310.15040","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["bingzhilee/slog"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/teleqna-a-benchmark-dataset-to-assess-large","slug":"teleqna-a-benchmark-dataset-to-assess-large","title":"TeleQnA: A Benchmark Dataset to Assess Large Language Models Telecommunications Knowledge","date":"2023-10-23","arxiv_id":"2310.15051","n_code_links":1,"syntology":null},{"paper":"/paper/text2topic-multi-label-text-classification","slug":"text2topic-multi-label-text-classification","title":"Text2Topic: Multi-Label Text Classification System for Efficient Topic Detection in User Generated Content with Zero-Shot Capabilities","date":"2023-10-23","arxiv_id":"2310.14817","n_code_links":1,"syntology":null},{"paper":"/paper/the-continued-usefulness-of-vocabulary-tests","slug":"the-continued-usefulness-of-vocabulary-tests","title":"Establishing Vocabulary Tests as a Benchmark for Evaluating Large Language Models","date":"2023-10-23","arxiv_id":"2310.14703","n_code_links":1,"syntology":null},{"paper":"/paper/towards-a-mechanistic-interpretation-of-multi","slug":"towards-a-mechanistic-interpretation-of-multi","title":"Towards a Mechanistic Interpretation of Multi-Step Reasoning Capabilities of Language Models","date":"2023-10-23","arxiv_id":"2310.14491","n_code_links":2,"syntology":{"ran":6,"of":7,"n_ran_checked":6,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["yifan-h/mechanisticprobe"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"unleashing-the-potential-of-prompt","title":"Unleashing the potential of prompt engineering for large language models","date":"2023-10-23","arxiv_id":"2310.14735","n_code_links":0,"syntology":null},{"paper":"/paper/a-pytorch-reproduction-of-masked-generative","slug":"a-pytorch-reproduction-of-masked-generative","title":"A Pytorch Reproduction of Masked Generative Image Transformer","date":"2023-10-22","arxiv_id":"2310.14400","n_code_links":1,"syntology":null},{"paper":"/paper/affine-consistent-transformer-for-multi-class-1","slug":"affine-consistent-transformer-for-multi-class-1","title":"Affine-Consistent Transformer for Multi-Class Cell Nuclei Detection","date":"2023-10-22","arxiv_id":"2310.14154","n_code_links":1,"syntology":null},{"paper":null,"slug":"an-overview-of-text-to-speech-systems-and","title":"An overview of text-to-speech systems and media applications","date":"2023-10-22","arxiv_id":"2310.14301","n_code_links":0,"syntology":null},{"paper":"/paper/can-language-models-laugh-at-youtube-short","slug":"can-language-models-laugh-at-youtube-short","title":"Can Language Models Laugh at YouTube Short-form Videos?","date":"2023-10-22","arxiv_id":"2310.14159","n_code_links":1,"syntology":null},{"paper":null,"slug":"convivit-a-deep-neural-network-combining","title":"ConViViT -- A Deep Neural Network Combining Convolutions and Factorized Self-Attention for Human Activity Recognition","date":"2023-10-22","arxiv_id":"2310.14416","n_code_links":0,"syntology":null},{"paper":"/paper/cxr-llava-multimodal-large-language-model-for","slug":"cxr-llava-multimodal-large-language-model-for","title":"CXR-LLAVA: a multimodal large language model for interpreting chest X-ray images","date":"2023-10-22","arxiv_id":"2310.18341","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["ecofri/cxr_llava"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/hierarchical-vector-quantized-transformer-for-1","slug":"hierarchical-vector-quantized-transformer-for-1","title":"Hierarchical Vector Quantized Transformer for Multi-class Unsupervised Anomaly Detection","date":"2023-10-22","arxiv_id":"2310.14228","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ruiyinglu/hvq-trans"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/is-chatgpt-a-game-changer-for-geocoding-a","slug":"is-chatgpt-a-game-changer-for-geocoding-a","title":"Is ChatGPT a game changer for geocoding -- a benchmark for geocoding address parsing techniques","date":"2023-10-22","arxiv_id":"2310.14360","n_code_links":1,"syntology":null},{"paper":null,"slug":"item-unsupervised-image-text-embedding","title":"ITEm: Unsupervised Image-Text Embedding Learning for eCommerce","date":"2023-10-22","arxiv_id":"2311.02084","n_code_links":0,"syntology":null},{"paper":"/paper/large-language-models-are-biased-to","slug":"large-language-models-are-biased-to","title":"Large Language Models are biased to overestimate profoundness","date":"2023-10-22","arxiv_id":"2310.14422","n_code_links":1,"syntology":null},{"paper":"/paper/manifold-preserving-transformers-are","slug":"manifold-preserving-transformers-are","title":"Manifold-Preserving Transformers are Effective for Short-Long Range Encoding","date":"2023-10-22","arxiv_id":"2310.14206","n_code_links":1,"syntology":null},{"paper":null,"slug":"mmtf-des-a-fusion-of-multimodal-transformer","title":"MMTF-DES: A Fusion of Multimodal Transformer Models for Desire, Emotion, and Sentiment Analysis of Social Media Data","date":"2023-10-22","arxiv_id":"2310.14143","n_code_links":0,"syntology":null},{"paper":"/paper/text-generation-for-dataset-augmentation-in","slug":"text-generation-for-dataset-augmentation-in","title":"Text generation for dataset augmentation in security classification tasks","date":"2023-10-22","arxiv_id":"2310.14429","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-harmful-erotic-content-detection","title":"Towards Harmful Erotic Content Detection through Coreference-Driven Contextual Analysis","date":"2023-10-22","arxiv_id":"2310.14325","n_code_links":0,"syntology":null},{"paper":"/paper/unimap-universal-smiles-graph-representation","slug":"unimap-universal-smiles-graph-representation","title":"UniMAP: Universal SMILES-Graph Representation Learning","date":"2023-10-22","arxiv_id":"2310.14216","n_code_links":1,"syntology":null},{"paper":"/paper/visual-attribute-prompt-learning-for","slug":"visual-attribute-prompt-learning-for","title":"Visual-Attribute Prompt Learning for Progressive Mild Cognitive Impairment Prediction","date":"2023-10-22","arxiv_id":"2310.14158","n_code_links":1,"syntology":null},{"paper":null,"slug":"asbart-accelerated-soft-bayes-additive","title":"ASBART:Accelerated Soft Bayes Additive Regression Trees","date":"2023-10-21","arxiv_id":"2310.13975","n_code_links":0,"syntology":null},{"paper":null,"slug":"concept-based-anomaly-detection-in-retail","title":"Concept-based Anomaly Detection in Retail Stores for Automatic Correction using Mobile Robots","date":"2023-10-21","arxiv_id":"2310.14063","n_code_links":0,"syntology":null},{"paper":null,"slug":"covidfakeexplainer-an-explainable-machine","title":"COVIDFakeExplainer: An Explainable Machine Learning based Web Application for Detecting COVID-19 Fake News","date":"2023-10-21","arxiv_id":"2310.13890","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-driving-behavior-for-autonomous","title":"Exploring Driving Behavior for Autonomous Vehicles Based on Gramian Angular Field Vision Transformer","date":"2023-10-21","arxiv_id":"2310.13906","n_code_links":0,"syntology":null},{"paper":"/paper/gemba-mqm-detecting-translation-quality-error","slug":"gemba-mqm-detecting-translation-quality-error","title":"GEMBA-MQM: Detecting Translation Quality Error Spans with GPT-4","date":"2023-10-21","arxiv_id":"2310.13988","n_code_links":1,"syntology":{"ran":8,"of":8,"n_ran_checked":8,"n_instrument":0,"unverified":0,"pointer_only":8,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"haterephrase-zero-and-few-shot-reduction-of","title":"HateRephrase: Zero- and Few-Shot Reduction of Hate Intensity in Online Posts using Large Language Models","date":"2023-10-21","arxiv_id":"2310.13985","n_code_links":0,"syntology":null},{"paper":"/paper/llm-prop-predicting-physical-and-electronic","slug":"llm-prop-predicting-physical-and-electronic","title":"LLM-Prop: Predicting Physical And Electronic Properties Of Crystalline Solids From Their Text Descriptions","date":"2023-10-21","arxiv_id":"2310.14029","n_code_links":1,"syntology":{"ran":0,"of":2,"n_ran_checked":0,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"0 ran · 2 unverified","official":{"repos":["vertaix/llm-prop"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":[]}}},{"paper":"/paper/multimodal-transformer-using-cross-channel","slug":"multimodal-transformer-using-cross-channel","title":"Multimodal Transformer Using Cross-Channel attention for Object Detection in Remote Sensing Images","date":"2023-10-21","arxiv_id":"2310.13876","n_code_links":1,"syntology":null},{"paper":"/paper/small-language-models-fine-tuned-to","slug":"small-language-models-fine-tuned-to","title":"Small Language Models Fine-tuned to Coordinate Larger Language Models improve Complex Reasoning","date":"2023-10-21","arxiv_id":"2310.18338","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":5,"n_instrument":0,"unverified":2,"pointer_only":7,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["lcs2-iiitd/daslam"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/a-simple-baseline-for-knowledge-based-visual","slug":"a-simple-baseline-for-knowledge-based-visual","title":"A Simple Baseline for Knowledge-Based Visual Question Answering","date":"2023-10-20","arxiv_id":"2310.13570","n_code_links":0,"syntology":{"ran":2,"of":3,"n_ran_checked":0,"n_instrument":2,"unverified":1,"pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":null}},{"paper":null,"slug":"alltogether-investigating-the-efficacy-of","title":"AllTogether: Investigating the Efficacy of Spliced Prompt for Web Navigation using Large Language Models","date":"2023-10-20","arxiv_id":"2310.18331","n_code_links":0,"syntology":null},{"paper":null,"slug":"anomaly-detection-of-command-shell-sessions","title":"Anomaly Detection of Command Shell Sessions based on DistilBERT: Unsupervised and Supervised Approaches","date":"2023-10-20","arxiv_id":"2310.13247","n_code_links":0,"syntology":null},{"paper":null,"slug":"ask-language-model-to-clean-your-noisy","title":"Ask Language Model to Clean Your Noisy Translation Data","date":"2023-10-20","arxiv_id":"2310.13469","n_code_links":0,"syntology":null},{"paper":null,"slug":"auxiliary-features-guided-super-resolution","title":"Auxiliary Features-Guided Super Resolution for Monte Carlo Rendering","date":"2023-10-20","arxiv_id":"2310.13235","n_code_links":0,"syntology":null},{"paper":"/paper/botchat-evaluating-llms-capabilities-of","slug":"botchat-evaluating-llms-capabilities-of","title":"BotChat: Evaluating LLMs' Capabilities of Having Multi-Turn Dialogues","date":"2023-10-20","arxiv_id":"2310.13650","n_code_links":1,"syntology":{"ran":1,"of":3,"n_ran_checked":0,"n_instrument":1,"unverified":2,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["open-compass/botchat"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/bridging-the-gap-between-synthetic-and","slug":"bridging-the-gap-between-synthetic-and","title":"Bridging the Gap between Synthetic and Authentic Images for Multimodal Machine Translation","date":"2023-10-20","arxiv_id":"2310.13361","n_code_links":1,"syntology":null},{"paper":"/paper/cache-me-if-you-can-an-online-cost-aware","slug":"cache-me-if-you-can-an-online-cost-aware","title":"Cache me if you Can: an Online Cost-aware Teacher-Student framework to Reduce the Calls to Large Language Models","date":"2023-10-20","arxiv_id":"2310.13395","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["stoyian/OCaTS"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"challenges-and-contributing-factors-in-the","title":"Challenges and Contributing Factors in the Utilization of Large Language Models (LLMs)","date":"2023-10-20","arxiv_id":"2310.13343","n_code_links":0,"syntology":null},{"paper":null,"slug":"design-inclusive-language-models-for","title":"She had Cobalt Blue Eyes: Prompt Testing to Create Aligned and Sustainable Language Models","date":"2023-10-20","arxiv_id":"2310.18333","n_code_links":0,"syntology":null},{"paper":null,"slug":"equivariant-transformer-is-all-you-need","title":"Equivariant Transformer is all you need","date":"2023-10-20","arxiv_id":"2310.13222","n_code_links":0,"syntology":null},{"paper":"/paper/evaluation-metrics-in-the-era-of-gpt-4","slug":"evaluation-metrics-in-the-era-of-gpt-4","title":"Evaluation Metrics in the Era of GPT-4: Reliably Evaluating Large Language Models on Sequence to Sequence Tasks","date":"2023-10-20","arxiv_id":"2310.13800","n_code_links":1,"syntology":{"ran":4,"of":5,"n_ran_checked":4,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["protagolabs/seq2seq_llm_evaluation"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/exploring-the-impact-of-corpus-diversity-on","slug":"exploring-the-impact-of-corpus-diversity-on","title":"Exploring the Impact of Corpus Diversity on Financial Pretrained Language Models","date":"2023-10-20","arxiv_id":"2310.13312","n_code_links":1,"syntology":null},{"paper":null,"slug":"fabula-intelligence-report-generation-using","title":"FABULA: Intelligence Report Generation Using Retrieval-Augmented Narrative Construction","date":"2023-10-20","arxiv_id":"2310.13848","n_code_links":0,"syntology":null},{"paper":null,"slug":"fmrt-learning-accurate-feature-matching-with","title":"FMRT: Learning Accurate Feature Matching with Reconciliatory Transformer","date":"2023-10-20","arxiv_id":"2310.13605","n_code_links":0,"syntology":null},{"paper":null,"slug":"foundation-model-s-embedded-representations","title":"Foundation Model's Embedded Representations May Detect Distribution Shift","date":"2023-10-20","arxiv_id":"2310.13836","n_code_links":0,"syntology":null},{"paper":"/paper/improving-cross-lingual-transfer-through","slug":"improving-cross-lingual-transfer-through","title":"Improving Cross-Lingual Transfer through Subtree-Aware Word Reordering","date":"2023-10-20","arxiv_id":"2310.13583","n_code_links":1,"syntology":null},{"paper":"/paper/moqagpt-zero-shot-multi-modal-open-domain","slug":"moqagpt-zero-shot-multi-modal-open-domain","title":"MoqaGPT : Zero-Shot Multi-modal Open-domain Question Answering with Large Language Model","date":"2023-10-20","arxiv_id":"2310.13265","n_code_links":1,"syntology":null},{"paper":"/paper/multi-level-contrastive-learning-for-script","slug":"multi-level-contrastive-learning-for-script","title":"Multi-level Contrastive Learning for Script-based Character Understanding","date":"2023-10-20","arxiv_id":"2310.13231","n_code_links":1,"syntology":null},{"paper":"/paper/plausibility-processing-in-transformer","slug":"plausibility-processing-in-transformer","title":"Plausibility Processing in Transformer Language Models: Focusing on the Role of Attention Heads in GPT","date":"2023-10-20","arxiv_id":"2310.13824","n_code_links":1,"syntology":null},{"paper":null,"slug":"potloc-pseudo-label-oriented-transformer-for","title":"POTLoc: Pseudo-Label Oriented Transformer for Point-Supervised Temporal Action Localization","date":"2023-10-20","arxiv_id":"2310.13585","n_code_links":0,"syntology":null},{"paper":null,"slug":"robust-training-for-conversational-question","title":"Robust Training for Conversational Question Answering Models with Reinforced Reformulation Generation","date":"2023-10-20","arxiv_id":"2310.13505","n_code_links":0,"syntology":null},{"paper":"/paper/skin-lesion-segmentation-improved-by","slug":"skin-lesion-segmentation-improved-by","title":"Skin Lesion Segmentation Improved by Transformer-based Networks with Inter-scale Dependency Modeling","date":"2023-10-20","arxiv_id":"2310.13604","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-impact-of-performance-expectancy-workload","title":"The Impact of Performance Expectancy, Workload, Risk, and Satisfaction on Trust in ChatGPT: Cross-sectional Survey Analysis","date":"2023-10-20","arxiv_id":"2311.05632","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-perils-promises-of-fact-checking-with","title":"The Perils & Promises of Fact-checking with Large Language Models","date":"2023-10-20","arxiv_id":"2310.13549","n_code_links":0,"syntology":null},{"paper":"/paper/tuna-instruction-tuning-using-feedback-from","slug":"tuna-instruction-tuning-using-feedback-from","title":"Tuna: Instruction Tuning using Feedback from Large Language Models","date":"2023-10-20","arxiv_id":"2310.13385","n_code_links":1,"syntology":null},{"paper":null,"slug":"wordart-designer-user-driven-artistic","title":"WordArt Designer: User-Driven Artistic Typography Synthesis using Large Language Models","date":"2023-10-20","arxiv_id":"2310.18332","n_code_links":0,"syntology":null},{"paper":null,"slug":"2d-3d-interlaced-transformer-for-point-cloud-1","title":"2D-3D Interlaced Transformer for Point Cloud Segmentation with Scene-Level Supervision","date":"2023-10-19","arxiv_id":"2310.12817","n_code_links":0,"syntology":null},{"paper":"/paper/agenttuning-enabling-generalized-agent","slug":"agenttuning-enabling-generalized-agent","title":"AgentTuning: Enabling Generalized Agent Abilities for LLMs","date":"2023-10-19","arxiv_id":"2310.12823","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["thudm/agenttuning"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"an-exploration-of-in-context-learning-for","title":"Exploring In-Context Learning of Textless Speech Language Model for Speech Classification Tasks","date":"2023-10-19","arxiv_id":"2310.12477","n_code_links":0,"syntology":null},{"paper":null,"slug":"automatic-hallucination-assessment-for","title":"ReEval: Automatic Hallucination Evaluation for Retrieval-Augmented Large Language Models via Transferable Adversarial Attacks","date":"2023-10-19","arxiv_id":"2310.12516","n_code_links":0,"syntology":null},{"paper":"/paper/automix-automatically-mixing-language-models","slug":"automix-automatically-mixing-language-models","title":"AutoMix: Automatically Mixing Language Models","date":"2023-10-19","arxiv_id":"2310.12963","n_code_links":1,"syntology":{"ran":2,"of":4,"n_ran_checked":0,"n_instrument":2,"unverified":2,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["automix-llm/automix"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/avtenet-audio-visual-transformer-based","slug":"avtenet-audio-visual-transformer-based","title":"AVTENet: Audio-Visual Transformer-based Ensemble Network Exploiting Multiple Experts for Video Deepfake Detection","date":"2023-10-19","arxiv_id":"2310.13103","n_code_links":0,"syntology":null},{"paper":"/paper/character-level-chinese-backpack-language","slug":"character-level-chinese-backpack-language","title":"Character-level Chinese Backpack Language Models","date":"2023-10-19","arxiv_id":"2310.12751","n_code_links":1,"syntology":null},{"paper":"/paper/da-transunet-integrating-spatial-and-channel","slug":"da-transunet-integrating-spatial-and-channel","title":"DA-TransUNet: Integrating Spatial and Channel Dual Attention with Transformer U-Net for Medical Image Segmentation","date":"2023-10-19","arxiv_id":"2310.12570","n_code_links":1,"syntology":null},{"paper":null,"slug":"energy-based-models-for-speech-synthesis","title":"Energy-Based Models For Speech Synthesis","date":"2023-10-19","arxiv_id":"2310.12765","n_code_links":0,"syntology":null},{"paper":"/paper/eureka-human-level-reward-design-via-coding","slug":"eureka-human-level-reward-design-via-coding","title":"Eureka: Human-Level Reward Design via Coding Large Language Models","date":"2023-10-19","arxiv_id":"2310.12931","n_code_links":1,"syntology":{"ran":11,"of":12,"n_ran_checked":9,"n_instrument":2,"unverified":1,"pointer_only":1,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 1 violated, 8 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["eureka-research/Eureka"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":1,"ran_from_kinds":["community","official"]}}},{"paper":null,"slug":"experimental-narratives-a-comparison-of-human","title":"Experimental Narratives: A Comparison of Human Crowdsourced Storytelling and AI Storytelling","date":"2023-10-19","arxiv_id":"2310.12902","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-graph-neural-networks-for-indian","title":"Exploring Graph Neural Networks for Indian Legal Judgment Prediction","date":"2023-10-19","arxiv_id":"2310.12800","n_code_links":0,"syntology":null},{"paper":null,"slug":"heart-disease-detection-using-vision-based","title":"Heart Disease Detection using Vision-Based Transformer Models from ECG Images","date":"2023-10-19","arxiv_id":"2310.12630","n_code_links":0,"syntology":null},{"paper":"/paper/identifying-and-adapting-transformer","slug":"identifying-and-adapting-transformer","title":"Identifying and Adapting Transformer-Components Responsible for Gender Bias in an English Language Model","date":"2023-10-19","arxiv_id":"2310.12611","n_code_links":1,"syntology":{"ran":4,"of":5,"n_ran_checked":4,"n_instrument":0,"unverified":1,"pointer_only":5,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["iabhijith/bias-causal-analysis"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"laser-linear-compression-in-wireless","title":"LASER: Linear Compression in Wireless Distributed Optimization","date":"2023-10-19","arxiv_id":"2310.13033","n_code_links":0,"syntology":null},{"paper":"/paper/letfuser-light-weight-end-to-end-transformer","slug":"letfuser-light-weight-end-to-end-transformer","title":"LeTFuser: Light-weight End-to-end Transformer-Based Sensor Fusion for Autonomous Driving with Multi-Task Learning","date":"2023-10-19","arxiv_id":"2310.13135","n_code_links":1,"syntology":null},{"paper":null,"slug":"medai-dialog-corpus-medic-zero-shot","title":"MedAI Dialog Corpus (MEDIC): Zero-Shot Classification of Doctor and AI Responses in Health Consultations","date":"2023-10-19","arxiv_id":"2310.12489","n_code_links":0,"syntology":null},{"paper":"/paper/minimalist-and-high-performance-semantic","slug":"minimalist-and-high-performance-semantic","title":"Minimalist and High-Performance Semantic Segmentation with Plain Vision Transformers","date":"2023-10-19","arxiv_id":"2310.12755","n_code_links":1,"syntology":null},{"paper":null,"slug":"multi-granularity-backprojection-transformer","title":"Multi-granularity Backprojection Transformer for Remote Sensing Image Super-Resolution","date":"2023-10-19","arxiv_id":"2310.12507","n_code_links":0,"syntology":null},{"paper":"/paper/non-autoregressive-sentence-ordering","slug":"non-autoregressive-sentence-ordering","title":"Non-Autoregressive Sentence Ordering","date":"2023-10-19","arxiv_id":"2310.12640","n_code_links":1,"syntology":null},{"paper":null,"slug":"not-all-countries-celebrate-thanksgiving-on","title":"Not All Countries Celebrate Thanksgiving: On the Cultural Dominance in Large Language Models","date":"2023-10-19","arxiv_id":"2310.12481","n_code_links":0,"syntology":null},{"paper":"/paper/predicting-ovarian-cancer-treatment-response","slug":"predicting-ovarian-cancer-treatment-response","title":"Predicting Ovarian Cancer Treatment Response in Histopathology using Hierarchical Vision Transformers and Multiple Instance Learning","date":"2023-10-19","arxiv_id":"2310.12866","n_code_links":1,"syntology":null},{"paper":"/paper/product-attribute-value-extraction-using","slug":"product-attribute-value-extraction-using","title":"ExtractGPT: Exploring the Potential of Large Language Models for Product Attribute Value Extraction","date":"2023-10-19","arxiv_id":"2310.12537","n_code_links":1,"syntology":null},{"paper":"/paper/real-time-motion-prediction-via-heterogeneous-1","slug":"real-time-motion-prediction-via-heterogeneous-1","title":"Real-Time Motion Prediction via Heterogeneous Polyline Transformer with Relative Pose Encoding","date":"2023-10-19","arxiv_id":"2310.12970","n_code_links":2,"syntology":{"ran":3,"of":5,"n_ran_checked":3,"n_instrument":0,"unverified":2,"pointer_only":5,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["zhejz/hptr"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/sequence-length-independent-norm-based","slug":"sequence-length-independent-norm-based","title":"Sequence Length Independent Norm-Based Generalization Bounds for Transformers","date":"2023-10-19","arxiv_id":"2310.13088","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["traugerjacob/transformer-gen-bounds"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/the-foundation-model-transparency-index","slug":"the-foundation-model-transparency-index","title":"The Foundation Model Transparency Index","date":"2023-10-19","arxiv_id":"2310.12941","n_code_links":1,"syntology":null},{"paper":"/paper/the-shifted-and-the-overlooked-a-task","slug":"the-shifted-and-the-overlooked-a-task","title":"The Shifted and The Overlooked: A Task-oriented Investigation of User-GPT Interactions","date":"2023-10-19","arxiv_id":"2310.12418","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ozyyshr/sharegpt_investigation"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/to-grok-or-not-to-grok-disentangling","slug":"to-grok-or-not-to-grok-disentangling","title":"To grok or not to grok: Disentangling generalization and memorization on corrupted algorithmic datasets","date":"2023-10-19","arxiv_id":"2310.13061","n_code_links":1,"syntology":{"ran":10,"of":10,"n_ran_checked":10,"n_instrument":0,"unverified":0,"pointer_only":10,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["d-doshi/Grokking"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"towards-robust-pruning-an-adaptive-knowledge","title":"Towards Robust Pruning: An Adaptive Knowledge-Retention Pruning Strategy for Language Models","date":"2023-10-19","arxiv_id":"2310.13191","n_code_links":0,"syntology":null},{"paper":"/paper/transformer-based-entity-legal-form","slug":"transformer-based-entity-legal-form","title":"Transformer-based Entity Legal Form Classification","date":"2023-10-19","arxiv_id":"2310.12766","n_code_links":1,"syntology":null},{"paper":"/paper/understanding-addition-in-transformers","slug":"understanding-addition-in-transformers","title":"Understanding Addition in Transformers","date":"2023-10-19","arxiv_id":"2310.13121","n_code_links":4,"syntology":{"ran":12,"of":16,"n_ran_checked":12,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["apartresearch/conceptual-interp","apartresearch/integer_addition","apartresearch/interger_addition"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/amr-parsing-with-causal-hierarchical","slug":"amr-parsing-with-causal-hierarchical","title":"AMR Parsing with Causal Hierarchical Attention and Pointers","date":"2023-10-18","arxiv_id":"2310.11964","n_code_links":1,"syntology":{"ran":5,"of":5,"n_ran_checked":3,"n_instrument":2,"unverified":0,"pointer_only":2,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["louchao98/chap_amr_parser"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"deep-learning-based-on-transformer","title":"Deep learning based on Transformer architecture for power system short-term voltage stability assessment with class imbalance","date":"2023-10-18","arxiv_id":"2310.11690","n_code_links":0,"syntology":null},{"paper":null,"slug":"direct-neural-machine-translation-with-task","title":"Direct Neural Machine Translation with Task-level Mixture of Experts models","date":"2023-10-18","arxiv_id":"2310.12236","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-the-symbol-binding-ability-of","title":"Evaluating the Symbol Binding Ability of Large Language Models for Multiple-Choice Questions in Vietnamese General Education","date":"2023-10-18","arxiv_id":"2310.12059","n_code_links":0,"syntology":null},{"paper":"/paper/fast-multipole-attention-a-divide-and-conquer","slug":"fast-multipole-attention-a-divide-and-conquer","title":"Fast Multipole Attention: A Divide-and-Conquer Attention Mechanism for Long Sequences","date":"2023-10-18","arxiv_id":"2310.11960","n_code_links":1,"syntology":null},{"paper":null,"slug":"field-testing-items-using-artificial","title":"Field-testing items using artificial intelligence: Natural language processing with transformers","date":"2023-10-18","arxiv_id":"2310.11655","n_code_links":0,"syntology":null}],"record_sha256":"d435e7c573b20286bc60fecf718208a08a35c06062133a618672cc61682ad42d","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}