{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/multi-head-attention/papers/123","list_of":"/method/multi-head-attention","method":"Multi-Head Attention","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":123,"pages_in_order":249,"rows_per_page":100,"rows":[12201,12300],"of":24855,"counts":{"archive_papers_tagged":24855,"with_a_code_link":11214,"where_syntology_ran_a_sample":3454,"not_listed_spam_title":0,"listed":24855,"listed_where_code_ran":3454,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":2916,"every_run_a_failure_of_syntologys_instrument":538,"listed_with_a_run_with_no_instrument_failure":2916,"listed_every_run_a_failure_of_syntologys_instrument":538,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/multi-head-attention","prev":"/method/multi-head-attention/papers/122","next":"/method/multi-head-attention/papers/124","papers":[{"paper":null,"slug":"the-paradigm-shifts-in-artificial","title":"The Paradigm Shifts in Artificial Intelligence","date":"2023-08-02","arxiv_id":"2308.02558","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-better-query-classification-with","title":"Towards Better Query Classification with Multi-Expert Knowledge Condensation in JD Ads Search","date":"2023-08-02","arxiv_id":"2308.01098","n_code_links":0,"syntology":null},{"paper":null,"slug":"chatmof-an-autonomous-ai-system-for","title":"ChatMOF: An Autonomous AI System for Predicting and Generating Metal-Organic Frameworks","date":"2023-08-01","arxiv_id":"2308.01423","n_code_links":0,"syntology":null},{"paper":null,"slug":"choir-transformer-generating-polyphonic-music","title":"Choir Transformer: Generating Polyphonic Music with Relative Attention on Transformer","date":"2023-08-01","arxiv_id":"2308.02531","n_code_links":0,"syntology":null},{"paper":null,"slug":"counterfactual-graph-transformer-for-traffic","title":"Counterfactual Graph Transformer for Traffic Flow Prediction","date":"2023-08-01","arxiv_id":"2308.00391","n_code_links":0,"syntology":null},{"paper":"/paper/dino-cxr-a-self-supervised-method-based-on","slug":"dino-cxr-a-self-supervised-method-based-on","title":"DINO-CXR: A self supervised method based on vision transformer for chest X-ray classification","date":"2023-08-01","arxiv_id":"2308.00475","n_code_links":0,"syntology":null},{"paper":"/paper/flatten-transformer-vision-transformer-using","slug":"flatten-transformer-vision-transformer-using","title":"FLatten Transformer: Vision Transformer using Focused Linear Attention","date":"2023-08-01","arxiv_id":"2308.00442","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["leaplabthu/flatten-transformer"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/improving-pixel-based-mim-by-reducing-wasted","slug":"improving-pixel-based-mim-by-reducing-wasted","title":"Improving Pixel-based MIM by Reducing Wasted Modeling Capability","date":"2023-08-01","arxiv_id":"2308.00261","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["open-mmlab/mmpretrain"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/instructed-to-bias-instruction-tuned-language","slug":"instructed-to-bias-instruction-tuned-language","title":"Instructed to Bias: Instruction-Tuned Language Models Exhibit Emergent Cognitive Bias","date":"2023-08-01","arxiv_id":"2308.00225","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["itay1itzhak/instructedtobias"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/partitioned-saliency-ranking-with-dense","slug":"partitioned-saliency-ranking-with-dense","title":"Partitioned Saliency Ranking with Dense Pyramid Transformers","date":"2023-08-01","arxiv_id":"2308.00236","n_code_links":1,"syntology":null},{"paper":"/paper/retrieval-augmented-generation-and","slug":"retrieval-augmented-generation-and","title":"Retrieval Augmented Generation and Representative Vector Summarization for large unstructured textual data in Medical Education","date":"2023-08-01","arxiv_id":"2308.00479","n_code_links":2,"syntology":null},{"paper":"/paper/self-supervised-contrastive-bert-fine-tuning","slug":"self-supervised-contrastive-bert-fine-tuning","title":"Self-Supervised Contrastive BERT Fine-tuning for Fusion-based Reviewed-Item Retrieval","date":"2023-08-01","arxiv_id":"2308.00762","n_code_links":2,"syntology":null},{"paper":"/paper/strip-attention-for-image-restoration","slug":"strip-attention-for-image-restoration","title":"Strip Attention for Image Restoration","date":"2023-08-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/towards-effective-ancient-chinese-translation","slug":"towards-effective-ancient-chinese-translation","title":"Towards Effective Ancient Chinese Translation: Dataset, Model, and Evaluation","date":"2023-08-01","arxiv_id":"2308.00240","n_code_links":1,"syntology":null},{"paper":"/paper/vit2eeg-leveraging-hybrid-pretrained-vision","slug":"vit2eeg-leveraging-hybrid-pretrained-vision","title":"ViT2EEG: Leveraging Hybrid Pretrained Vision Transformers for EEG Data","date":"2023-08-01","arxiv_id":"2308.00454","n_code_links":2,"syntology":null},{"paper":null,"slug":"a-pre-trained-data-deduplication-model-based","title":"A Pre-trained Data Deduplication Model based on Active Learning","date":"2023-07-31","arxiv_id":"2308.00721","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-effective-data-creation-pipeline-to","title":"An Effective Data Creation Pipeline to Generate High-quality Financial Instruction Data for Large Language Model","date":"2023-07-31","arxiv_id":"2308.01415","n_code_links":0,"syntology":null},{"paper":null,"slug":"capturing-co-existing-distortions-in-user","title":"Capturing Co-existing Distortions in User-Generated Content for No-reference Video Quality Assessment","date":"2023-07-31","arxiv_id":"2307.16813","n_code_links":0,"syntology":null},{"paper":"/paper/classifying-multilingual-party-manifestos","slug":"classifying-multilingual-party-manifestos","title":"Classifying multilingual party manifestos: Domain transfer across country, time, and genre","date":"2023-07-31","arxiv_id":"2307.16511","n_code_links":1,"syntology":null},{"paper":"/paper/conformal-pid-control-for-time-series-1","slug":"conformal-pid-control-for-time-series-1","title":"Conformal PID Control for Time Series Prediction","date":"2023-07-31","arxiv_id":"2307.16895","n_code_links":1,"syntology":{"ran":6,"of":10,"n_ran_checked":6,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["aangelopoulos/conformal-time-series"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"dctm-dilated-convolutional-transformer-model","title":"DCTM: Dilated Convolutional Transformer Model for Multimodal Engagement Estimation in Conversation","date":"2023-07-31","arxiv_id":"2308.01966","n_code_links":0,"syntology":null},{"paper":null,"slug":"deception-abilities-emerged-in-large-language","title":"Deception Abilities Emerged in Large Language Models","date":"2023-07-31","arxiv_id":"2307.16513","n_code_links":0,"syntology":null},{"paper":"/paper/does-fine-tuning-gpt-3-with-the-openai-api","slug":"does-fine-tuning-gpt-3-with-the-openai-api","title":"Does fine-tuning GPT-3 with the OpenAI API leak personally-identifiable information?","date":"2023-07-31","arxiv_id":"2307.16382","n_code_links":1,"syntology":null},{"paper":"/paper/hagrid-a-human-llm-collaborative-dataset-for","slug":"hagrid-a-human-llm-collaborative-dataset-for","title":"HAGRID: A Human-LLM Collaborative Dataset for Generative Information-Seeking with Attribution","date":"2023-07-31","arxiv_id":"2307.16883","n_code_links":1,"syntology":null},{"paper":"/paper/no-that-s-not-what-i-meant-handling-third","slug":"no-that-s-not-what-i-meant-handling-third","title":"No that's not what I meant: Handling Third Position Repair in Conversational Question Answering","date":"2023-07-31","arxiv_id":"2307.16689","n_code_links":1,"syntology":null},{"paper":"/paper/noisy-self-training-with-data-augmentations","slug":"noisy-self-training-with-data-augmentations","title":"Noisy Self-Training with Data Augmentations for Offensive and Hate Speech Detection Tasks","date":"2023-07-31","arxiv_id":"2307.16609","n_code_links":1,"syntology":null},{"paper":null,"slug":"ontology-engineering-with-large-language","title":"Ontology engineering with Large Language Models","date":"2023-07-31","arxiv_id":"2307.16699","n_code_links":0,"syntology":null},{"paper":null,"slug":"performance-evaluation-of-swin-vision","title":"Performance Evaluation of Swin Vision Transformer Model using Gradient Accumulation Optimization Technique","date":"2023-07-31","arxiv_id":"2308.00197","n_code_links":0,"syntology":null},{"paper":"/paper/vacancysbert-the-approach-for-representation","slug":"vacancysbert-the-approach-for-representation","title":"VacancySBERT: the approach for representation of titles and skills for semantic similarity search in the recruitment domain","date":"2023-07-31","arxiv_id":"2307.16638","n_code_links":1,"syntology":null},{"paper":"/paper/clgt-a-graph-transformer-for-student","slug":"clgt-a-graph-transformer-for-student","title":"CLGT: A Graph Transformer for Student Performance Prediction in Collaborative Learning","date":"2023-07-30","arxiv_id":"2308.02038","n_code_links":1,"syntology":null},{"paper":null,"slug":"evaluating-chatgpt-and-gpt-4-for-visual","title":"Evaluating ChatGPT and GPT-4 for Visual Programming","date":"2023-07-30","arxiv_id":"2308.02522","n_code_links":0,"syntology":null},{"paper":null,"slug":"laficmil-rethinking-large-file-classification","title":"LaFiCMIL: Rethinking Large File Classification from the Perspective of Correlated Multiple Instance Learning","date":"2023-07-30","arxiv_id":"2308.01413","n_code_links":0,"syntology":null},{"paper":null,"slug":"pre-training-end-to-end-asr-models-with","title":"Pre-training End-to-end ASR Models with Augmented Speech Samples Queried by Text","date":"2023-07-30","arxiv_id":"2307.16332","n_code_links":0,"syntology":null},{"paper":"/paper/scribblevc-scribble-supervised-medical-image","slug":"scribblevc-scribble-supervised-medical-image","title":"ScribbleVC: Scribble-supervised Medical Image Segmentation with Vision-Class Embedding","date":"2023-07-30","arxiv_id":"2307.16226","n_code_links":1,"syntology":null},{"paper":"/paper/seed-bench-benchmarking-multimodal-llms-with","slug":"seed-bench-benchmarking-multimodal-llms-with","title":"SEED-Bench: Benchmarking Multimodal LLMs with Generative Comprehension","date":"2023-07-30","arxiv_id":"2307.16125","n_code_links":3,"syntology":null},{"paper":"/paper/styleprompter-all-styles-need-is-attention","slug":"styleprompter-all-styles-need-is-attention","title":"StylePrompter: All Styles Need Is Attention","date":"2023-07-30","arxiv_id":"2307.16151","n_code_links":1,"syntology":null},{"paper":"/paper/transfusion-a-practical-and-effective","slug":"transfusion-a-practical-and-effective","title":"TransFusion: A Practical and Effective Transformer-based Diffusion Model for 3D Human Motion Prediction","date":"2023-07-30","arxiv_id":"2307.16106","n_code_links":1,"syntology":null},{"paper":null,"slug":"video-frame-interpolation-with-flow","title":"Video Frame Interpolation with Flow Transformer","date":"2023-07-30","arxiv_id":"2307.16144","n_code_links":0,"syntology":null},{"paper":null,"slug":"covid-19-detection-leveraging-vision","title":"CoVid-19 Detection leveraging Vision Transformers and Explainable AI","date":"2023-07-29","arxiv_id":"2307.16033","n_code_links":0,"syntology":null},{"paper":null,"slug":"handmim-pose-aware-self-supervised-learning","title":"HandMIM: Pose-Aware Self-Supervised Learning for 3D Hand Mesh Estimation","date":"2023-07-29","arxiv_id":"2307.16061","n_code_links":0,"syntology":null},{"paper":null,"slug":"monaural-multi-speaker-speech-separation","title":"Monaural Multi-Speaker Speech Separation Using Efficient Transformer Model","date":"2023-07-29","arxiv_id":"2308.00010","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-critical-review-of-large-language-models","title":"A Critical Review of Large Language Models: Sensitivity, Bias, and the Path Toward Specialized AI","date":"2023-07-28","arxiv_id":"2307.15425","n_code_links":0,"syntology":null},{"paper":null,"slug":"beyond-reality-the-pivotal-role-of-generative","title":"Beyond Reality: The Pivotal Role of Generative AI in the Metaverse","date":"2023-07-28","arxiv_id":"2308.06272","n_code_links":0,"syntology":null},{"paper":"/paper/chathome-development-and-evaluation-of-a","slug":"chathome-development-and-evaluation-of-a","title":"ChatHome: Development and Evaluation of a Domain-Specific Language Model for Home Renovation","date":"2023-07-28","arxiv_id":"2307.15290","n_code_links":1,"syntology":{"ran":7,"of":17,"n_ran_checked":4,"n_instrument":3,"unverified":10,"pointer_only":1,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 3 where Syntology's instrument failed) · 10 unverified","official":{"repos":["lianjiatech/belle"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":10,"ran_from_kinds":["official"]}}},{"paper":"/paper/differential-evolution-algorithm-based-hyper","slug":"differential-evolution-algorithm-based-hyper","title":"Differential Evolution Algorithm based Hyper-Parameters Selection of Transformer Neural Network Model for Load Forecasting","date":"2023-07-28","arxiv_id":"2307.15299","n_code_links":1,"syntology":null},{"paper":null,"slug":"docdeshadower-frequency-aware-transformer-for","title":"DocDeshadower: Frequency-Aware Transformer for Document Shadow Removal","date":"2023-07-28","arxiv_id":"2307.15318","n_code_links":0,"syntology":null},{"paper":"/paper/med-halt-medical-domain-hallucination-test","slug":"med-halt-medical-domain-hallucination-test","title":"Med-HALT: Medical Domain Hallucination Test for Large Language Models","date":"2023-07-28","arxiv_id":"2307.15343","n_code_links":1,"syntology":null},{"paper":"/paper/memotr-long-term-memory-augmented-transformer","slug":"memotr-long-term-memory-augmented-transformer","title":"MeMOTR: Long-Term Memory-Augmented Transformer for Multi-Object Tracking","date":"2023-07-28","arxiv_id":"2307.15700","n_code_links":1,"syntology":{"ran":10,"of":14,"n_ran_checked":7,"n_instrument":3,"unverified":4,"pointer_only":0,"phrase":"10 ran (of which 5 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 3 where Syntology's instrument failed) · 4 unverified","official":{"repos":["mcg-nju/memotr"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":5,"n_ran_no_instrument_failure":7,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"pcnn-a-lightweight-parallel-conformer-neural","title":"PCNN: A Lightweight Parallel Conformer Neural Network for Efficient Monaural Speech Enhancement","date":"2023-07-28","arxiv_id":"2307.15251","n_code_links":0,"syntology":null},{"paper":null,"slug":"point-clouds-are-specialized-images-a","title":"Point Clouds Are Specialized Images: A Knowledge Transfer Approach for 3D Understanding","date":"2023-07-28","arxiv_id":"2307.15569","n_code_links":0,"syntology":null},{"paper":"/paper/prompt-guided-transformer-for-multi-task","slug":"prompt-guided-transformer-for-multi-task","title":"Prompt Guided Transformer for Multi-Task Dense Prediction","date":"2023-07-28","arxiv_id":"2307.15362","n_code_links":1,"syntology":null},{"paper":"/paper/rsgpt-a-remote-sensing-vision-language-model","slug":"rsgpt-a-remote-sensing-vision-language-model","title":"RSGPT: A Remote Sensing Vision Language Model and Benchmark","date":"2023-07-28","arxiv_id":"2307.15266","n_code_links":2,"syntology":null},{"paper":null,"slug":"simdetr-simplifying-self-supervised","title":"Aligned Unsupervised Pretraining of Object Detectors with Self-training","date":"2023-07-28","arxiv_id":"2307.15697","n_code_links":0,"syntology":null},{"paper":null,"slug":"tutorials-on-stance-detection-using-pre","title":"Tutorials on Stance Detection using Pre-trained Language Models: Fine-tuning BERT and Prompting Large Language Models","date":"2023-07-28","arxiv_id":"2307.15331","n_code_links":0,"syntology":null},{"paper":null,"slug":"verigen-a-large-language-model-for-verilog","title":"VeriGen: A Large Language Model for Verilog Code Generation","date":"2023-07-28","arxiv_id":"2308.00708","n_code_links":0,"syntology":null},{"paper":"/paper/vpp-efficient-conditional-3d-generation-via-1","slug":"vpp-efficient-conditional-3d-generation-via-1","title":"VPP: Efficient Conditional 3D Generation via Voxel-Point Progressive Representation","date":"2023-07-28","arxiv_id":"2307.16605","n_code_links":2,"syntology":{"ran":6,"of":7,"n_ran_checked":4,"n_instrument":2,"unverified":1,"pointer_only":4,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 2 honoured, 0 violated, 2 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["qizekun/vpp"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"a-transformer-based-approach-for-arabic","title":"A Transformer-based Approach for Arabic Offline Handwritten Text Recognition","date":"2023-07-27","arxiv_id":"2307.15045","n_code_links":0,"syntology":null},{"paper":null,"slug":"arc-nlp-at-pan-2023-hierarchical-long-text","title":"ARC-NLP at PAN 2023: Hierarchical Long Text Classification for Trigger Detection","date":"2023-07-27","arxiv_id":"2307.14912","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-generative-models-for-graph-to","slug":"evaluating-generative-models-for-graph-to","title":"Evaluating Generative Models for Graph-to-Text Generation","date":"2023-07-27","arxiv_id":"2307.14712","n_code_links":1,"syntology":null},{"paper":"/paper/htnet-for-micro-expression-recognition","slug":"htnet-for-micro-expression-recognition","title":"HTNet for micro-expression recognition","date":"2023-07-27","arxiv_id":"2307.14637","n_code_links":1,"syntology":null},{"paper":"/paper/iml-vit-image-manipulation-localization-by","slug":"iml-vit-image-manipulation-localization-by","title":"IML-ViT: Benchmarking Image Manipulation Localization by Vision Transformer","date":"2023-07-27","arxiv_id":"2307.14863","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["sunnyhaze/iml-vit"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/improving-aspect-based-sentiment-with-end-to","slug":"improving-aspect-based-sentiment-with-end-to","title":"Improving Aspect-Based Sentiment with End-to-End Semantic Role Labeling Model","date":"2023-07-27","arxiv_id":"2307.14785","n_code_links":1,"syntology":null},{"paper":null,"slug":"llmediator-gpt-4-assisted-online-dispute","title":"LLMediator: GPT-4 Assisted Online Dispute Resolution","date":"2023-07-27","arxiv_id":"2307.16732","n_code_links":0,"syntology":null},{"paper":"/paper/mcpa-multi-scale-cross-perceptron-attention","slug":"mcpa-multi-scale-cross-perceptron-attention","title":"MCPA: Multi-scale Cross Perceptron Attention Network for 2D Medical Image Segmentation","date":"2023-07-27","arxiv_id":"2307.14588","n_code_links":1,"syntology":null},{"paper":"/paper/metric-based-in-context-learning-a-case-study","slug":"metric-based-in-context-learning-a-case-study","title":"Metric-Based In-context Learning: A Case Study in Text Simplification","date":"2023-07-27","arxiv_id":"2307.14632","n_code_links":1,"syntology":null},{"paper":"/paper/new-interaction-paradigm-for-complex-eda","slug":"new-interaction-paradigm-for-complex-eda","title":"New Interaction Paradigm for Complex EDA Software Leveraging GPT","date":"2023-07-27","arxiv_id":"2307.14740","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["smarton-empower/smarton-ai"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/scaling-session-based-transformer","slug":"scaling-session-based-transformer","title":"Scaling Session-Based Transformer Recommendations using Optimized Negative Sampling and Loss Functions","date":"2023-07-27","arxiv_id":"2307.14906","n_code_links":1,"syntology":null},{"paper":"/paper/scaling-transnormer-to-175-billion-parameters","slug":"scaling-transnormer-to-175-billion-parameters","title":"TransNormerLLM: A Faster and Better Large Language Model with Improved TransNormer","date":"2023-07-27","arxiv_id":"2307.14995","n_code_links":2,"syntology":{"ran":3,"of":5,"n_ran_checked":1,"n_instrument":2,"unverified":2,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["opennlplab/transnormerllm"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"self-supervised-graph-transformer-for","title":"Self-Supervised Graph Transformer for Deepfake Detection","date":"2023-07-27","arxiv_id":"2307.15019","n_code_links":0,"syntology":null},{"paper":"/paper/superclue-a-comprehensive-chinese-large","slug":"superclue-a-comprehensive-chinese-large","title":"SuperCLUE: A Comprehensive Chinese Large Language Model Benchmark","date":"2023-07-27","arxiv_id":"2307.15020","n_code_links":0,"syntology":null},{"paper":null,"slug":"textmania-enriching-visual-feature-by-text","title":"TextManiA: Enriching Visual Feature by Text-driven Manifold Augmentation","date":"2023-07-27","arxiv_id":"2307.14611","n_code_links":0,"syntology":null},{"paper":null,"slug":"visu-at-wassa-2023-shared-task-detecting","title":"VISU at WASSA 2023 Shared Task: Detecting Emotions in Reaction to News Stories Leveraging BERT and Stacked Embeddings","date":"2023-07-27","arxiv_id":"2307.15164","n_code_links":0,"syntology":null},{"paper":null,"slug":"affective-natural-language-generation-of","title":"Affective Natural Language Generation of Event Descriptions through Fine-grained Appraisal Conditions","date":"2023-07-26","arxiv_id":"2307.14004","n_code_links":0,"syntology":null},{"paper":null,"slug":"are-transformers-with-one-layer-self","title":"Are Transformers with One Layer Self-Attention Using Low-Rank Weight Matrices Universal Approximators?","date":"2023-07-26","arxiv_id":"2307.14023","n_code_links":0,"syntology":null},{"paper":null,"slug":"clinidigest-a-case-study-in-large-language","title":"CliniDigest: A Case Study in Large Language Model Based Large-Scale Summarization of Clinical Trial Descriptions","date":"2023-07-26","arxiv_id":"2307.14522","n_code_links":0,"syntology":null},{"paper":null,"slug":"comparative-analysis-of-libraries-for-the","title":"Comparative Analysis of Libraries for the Sentimental Analysis","date":"2023-07-26","arxiv_id":"2307.14311","n_code_links":0,"syntology":null},{"paper":null,"slug":"decoding-chatgpt-a-taxonomy-of-existing","title":"Decoding ChatGPT: A Taxonomy of Existing Research, Current Challenges, and Possible Future Directions","date":"2023-07-26","arxiv_id":"2307.14107","n_code_links":0,"syntology":null},{"paper":null,"slug":"developing-and-evaluating-tiny-to-medium","title":"Developing and Evaluating Tiny to Medium-Sized Turkish BERT Models","date":"2023-07-26","arxiv_id":"2307.14134","n_code_links":0,"syntology":null},{"paper":null,"slug":"dpbert-efficient-inference-for-bert-based-on","title":"DPBERT: Efficient Inference for BERT based on Dynamic Planning","date":"2023-07-26","arxiv_id":"2308.00108","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhanced-security-against-adversarial","title":"Enhanced Security against Adversarial Examples Using a Random Ensemble of Encrypted Vision Transformer Models","date":"2023-07-26","arxiv_id":"2307.13985","n_code_links":0,"syntology":null},{"paper":"/paper/essaformer-efficient-transformer-for","slug":"essaformer-efficient-transformer-for","title":"ESSAformer: Efficient Transformer for Hyperspectral Image Super-resolution","date":"2023-07-26","arxiv_id":"2307.14010","n_code_links":1,"syntology":{"ran":7,"of":9,"n_ran_checked":7,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"7 ran (of which 4 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["rexzhan/essaformer"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":4,"n_ran_no_instrument_failure":7,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/event-based-vision-for-early-prediction-of","slug":"event-based-vision-for-early-prediction-of","title":"Event-based Vision for Early Prediction of Manipulation Actions","date":"2023-07-26","arxiv_id":"2307.14332","n_code_links":1,"syntology":null},{"paper":"/paper/exedec-execution-decomposition-for","slug":"exedec-execution-decomposition-for","title":"ExeDec: Execution Decomposition for Compositional Generalization in Neural Program Synthesis","date":"2023-07-26","arxiv_id":"2307.13883","n_code_links":0,"syntology":{"ran":13,"of":16,"n_ran_checked":13,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 0 honoured, 0 violated, 13 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":null}},{"paper":null,"slug":"fintree-financial-dataset-pretrain","title":"FinTree: Financial Dataset Pretrain Transformer Encoder for Relation Extraction","date":"2023-07-26","arxiv_id":"2307.13900","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-user-language-affects-conflict-fatality","title":"How User Language Affects Conflict Fatality Estimates in ChatGPT","date":"2023-07-26","arxiv_id":"2308.00072","n_code_links":0,"syntology":null},{"paper":"/paper/improving-existing-segmentators-performance","slug":"improving-existing-segmentators-performance","title":"Improving existing segmentators performance with zero-shot segmentators","date":"2023-07-26","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/leveraging-large-language-models-for-mental","slug":"leveraging-large-language-models-for-mental","title":"Mental-LLM: Leveraging Large Language Models for Mental Health Prediction via Online Text Data","date":"2023-07-26","arxiv_id":"2307.14385","n_code_links":1,"syntology":null},{"paper":"/paper/midas-v3-1-a-model-zoo-for-robust-monocular","slug":"midas-v3-1-a-model-zoo-for-robust-monocular","title":"MiDaS v3.1 -- A Model Zoo for Robust Monocular Relative Depth Estimation","date":"2023-07-26","arxiv_id":"2307.14460","n_code_links":2,"syntology":{"ran":1,"of":3,"n_ran_checked":0,"n_instrument":1,"unverified":2,"pointer_only":3,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["isl-org/MiDaS"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"speed-reading-tool-powered-by-artificial","title":"Speed Reading Tool Powered by Artificial Intelligence for Students with ADHD, Dyslexia, or Short Attention Span","date":"2023-07-26","arxiv_id":"2307.14544","n_code_links":0,"syntology":null},{"paper":null,"slug":"understanding-deep-neural-networks-via-linear","title":"Understanding Deep Neural Networks via Linear Separability of Hidden Layers","date":"2023-07-26","arxiv_id":"2307.13962","n_code_links":0,"syntology":null},{"paper":null,"slug":"unveiling-security-privacy-and-ethical","title":"Unveiling Security, Privacy, and Ethical Concerns of ChatGPT","date":"2023-07-26","arxiv_id":"2307.14192","n_code_links":0,"syntology":null},{"paper":null,"slug":"visual-prompt-flexible-modal-face-anti","title":"Visual Prompt Flexible-Modal Face Anti-Spoofing","date":"2023-07-26","arxiv_id":"2307.13958","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-end-to-end-workflow-using-topic","title":"An End-to-End Workflow using Topic Segmentation and Text Summarisation Methods for Improved Podcast Comprehension","date":"2023-07-25","arxiv_id":"2307.13394","n_code_links":0,"syntology":null},{"paper":null,"slug":"arb-advanced-reasoning-benchmark-for-large","title":"ARB: Advanced Reasoning Benchmark for Large Language Models","date":"2023-07-25","arxiv_id":"2307.13692","n_code_links":0,"syntology":null},{"paper":null,"slug":"conditional-cross-attention-network-for-multi","title":"Conditional Cross Attention Network for Multi-Space Embedding without Entanglement in Only a SINGLE Network","date":"2023-07-25","arxiv_id":"2307.13254","n_code_links":0,"syntology":null},{"paper":null,"slug":"curvature-based-transformer-for-molecular","title":"CTAGE: Curvature-Based Topology-Aware Graph Embedding for Learning Molecular Representations","date":"2023-07-25","arxiv_id":"2307.13275","n_code_links":0,"syntology":null},{"paper":"/paper/geotransformer-fast-and-robust-point-cloud","slug":"geotransformer-fast-and-robust-point-cloud","title":"GeoTransformer: Fast and Robust Point Cloud Registration with Geometric Transformer","date":"2023-07-25","arxiv_id":"2308.03768","n_code_links":1,"syntology":null},{"paper":null,"slug":"gpt-3-models-are-few-shot-financial-reasoners","title":"GPT-3 Models are Few-Shot Financial Reasoners","date":"2023-07-25","arxiv_id":"2307.13617","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-can-large-language-models-help-humans-in","title":"How Can Large Language Models Help Humans in Design and Manufacturing?","date":"2023-07-25","arxiv_id":"2307.14377","n_code_links":0,"syntology":null},{"paper":null,"slug":"is-gpt-a-computational-model-of-emotion","title":"Is GPT a Computational Model of Emotion? Detailed Analysis","date":"2023-07-25","arxiv_id":"2307.13779","n_code_links":0,"syntology":null}],"record_sha256":"60bd915f471bcd1d7eef6fe2f003b9ef872b1949ab1ec63a15632c8fb5c7814f","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}