{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/linear-layer/papers/64","list_of":"/method/linear-layer","method":"Linear Layer","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":64,"pages_in_order":255,"rows_per_page":100,"rows":[6301,6400],"of":25421,"counts":{"archive_papers_tagged":25421,"with_a_code_link":11479,"where_syntology_ran_a_sample":3523,"not_listed_spam_title":0,"listed":25421,"listed_where_code_ran":3523,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":2976,"every_run_a_failure_of_syntologys_instrument":547,"listed_with_a_run_with_no_instrument_failure":2976,"listed_every_run_a_failure_of_syntologys_instrument":547,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/linear-layer","prev":"/method/linear-layer/papers/63","next":"/method/linear-layer/papers/65","papers":[{"paper":null,"slug":"enhancing-the-llm-based-robot-manipulation","title":"Enhancing the LLM-Based Robot Manipulation Through Human-Robot Collaboration","date":"2024-06-20","arxiv_id":"2406.14097","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-implicit-bias-in-large-language","slug":"evaluating-implicit-bias-in-large-language","title":"Evaluating Implicit Bias in Large Language Models by Attacking From a Psychometric Perspective","date":"2024-06-20","arxiv_id":"2406.14023","n_code_links":1,"syntology":null},{"paper":"/paper/evaluating-rag-fusion-with-ragelo-an","slug":"evaluating-rag-fusion-with-ragelo-an","title":"Evaluating RAG-Fusion with RAGElo: an Automated Elo-based Framework","date":"2024-06-20","arxiv_id":"2406.14783","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":3,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["zetaalphavector/ragelo"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"factual-dialogue-summarization-via-learning","title":"Factual Dialogue Summarization via Learning from Large Language Models","date":"2024-06-20","arxiv_id":"2406.14709","n_code_links":0,"syntology":null},{"paper":null,"slug":"generative-ai-for-enhancing-active-learning","title":"Generative AI for Enhancing Active Learning in Education: A Comparative Study of GPT-3.5 and GPT-4 in Crafting Customized Test Questions","date":"2024-06-20","arxiv_id":"2406.13903","n_code_links":0,"syntology":null},{"paper":null,"slug":"healing-powers-of-bert-how-task-specific-fine","title":"Healing Powers of BERT: How Task-Specific Fine-Tuning Recovers Corrupted Language Models","date":"2024-06-20","arxiv_id":"2406.14459","n_code_links":0,"syntology":null},{"paper":"/paper/how-to-compute-the-probability-of-a-word","slug":"how-to-compute-the-probability-of-a-word","title":"How to Compute the Probability of a Word","date":"2024-06-20","arxiv_id":"2406.14561","n_code_links":2,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["tpimentelms/probability-of-a-word"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"identifying-user-goals-from-ui-trajectories","title":"Identifying User Goals from UI Trajectories","date":"2024-06-20","arxiv_id":"2406.14314","n_code_links":0,"syntology":null},{"paper":"/paper/learning-to-plan-for-retrieval-augmented","slug":"learning-to-plan-for-retrieval-augmented","title":"Learning to Plan for Retrieval-Augmented Large Language Models from Knowledge Graphs","date":"2024-06-20","arxiv_id":"2406.14282","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["zjukg/lpkg"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/llasa-large-multimodal-agent-for-human","slug":"llasa-large-multimodal-agent-for-human","title":"LLaSA: A Multimodal LLM for Human Activity Analysis Through Wearable and Smartphone Sensors","date":"2024-06-20","arxiv_id":"2406.14498","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":3,"n_instrument":3,"unverified":1,"pointer_only":7,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 1 violated, 2 with no contract checked; 3 where Syntology's instrument failed) · 1 unverified","official":{"repos":["bashlab/llasa"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/mm-gtunets-unified-multi-modal-graph-deep","slug":"mm-gtunets-unified-multi-modal-graph-deep","title":"MM-GTUNets: Unified Multi-Modal Graph Deep Learning for Brain Disorders Prediction","date":"2024-06-20","arxiv_id":"2406.14455","n_code_links":1,"syntology":null},{"paper":"/paper/mmbench-video-a-long-form-multi-shot","slug":"mmbench-video-a-long-form-multi-shot","title":"MMBench-Video: A Long-Form Multi-Shot Benchmark for Holistic Video Understanding","date":"2024-06-20","arxiv_id":"2406.14515","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":1,"n_instrument":4,"unverified":1,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","official":{"repos":["open-compass/vlmevalkit"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"mr-ben-a-comprehensive-meta-reasoning","title":"MR-Ben: A Meta-Reasoning Benchmark for Evaluating System-2 Thinking in LLMs","date":"2024-06-20","arxiv_id":"2406.13975","n_code_links":0,"syntology":null},{"paper":"/paper/persuasiveness-of-generated-free-text","slug":"persuasiveness-of-generated-free-text","title":"Persuasiveness of Generated Free-Text Rationales in Subjective Decisions: A Case Study on Pairwise Argument Ranking","date":"2024-06-20","arxiv_id":"2406.13905","n_code_links":1,"syntology":null},{"paper":"/paper/prism-a-framework-for-decoupling-and","slug":"prism-a-framework-for-decoupling-and","title":"Prism: A Framework for Decoupling and Assessing the Capabilities of VLMs","date":"2024-06-20","arxiv_id":"2406.14544","n_code_links":1,"syntology":{"ran":9,"of":13,"n_ran_checked":9,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["sparksjoe/prism"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"relation-extraction-with-fine-tuned-large","title":"Relation Extraction with Fine-Tuned Large Language Models in Retrieval Augmented Generation Frameworks","date":"2024-06-20","arxiv_id":"2406.14745","n_code_links":0,"syntology":null},{"paper":null,"slug":"rtformer-re-parameter-tsbn-spiking","title":"RTFormer: Re-parameter TSBN Spiking Transformer","date":"2024-06-20","arxiv_id":"2406.14180","n_code_links":0,"syntology":null},{"paper":null,"slug":"scidmt-a-large-scale-corpus-for-detecting","title":"SciDMT: A Large-Scale Corpus for Detecting Scientific Mentions","date":"2024-06-20","arxiv_id":"2406.14756","n_code_links":0,"syntology":null},{"paper":"/paper/seg-lstm-performance-of-xlstm-for-semantic","slug":"seg-lstm-performance-of-xlstm-for-semantic","title":"Seg-LSTM: Performance of xLSTM for Semantic Segmentation of Remotely Sensed Images","date":"2024-06-20","arxiv_id":"2406.14086","n_code_links":1,"syntology":null},{"paper":"/paper/sorry-bench-systematically-evaluating-large","slug":"sorry-bench-systematically-evaluating-large","title":"SORRY-Bench: Systematically Evaluating Large Language Model Safety Refusal Behaviors","date":"2024-06-20","arxiv_id":"2406.14598","n_code_links":1,"syntology":null},{"paper":null,"slug":"spl-a-socratic-playground-for-learning","title":"SPL: A Socratic Playground for Learning Powered by Large Language Model","date":"2024-06-20","arxiv_id":"2406.13919","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-use-of-multimodal-large-language-models","title":"The Use of Multimodal Large Language Models to Detect Objects from Thermal Images: Transportation Applications","date":"2024-06-20","arxiv_id":"2406.13898","n_code_links":0,"syntology":null},{"paper":null,"slug":"ttqa-rs-a-break-down-prompting-approach-for","title":"TTQA-RS- A break-down prompting approach for Multi-hop Table-Text Question Answering with Reasoning and Summarization","date":"2024-06-20","arxiv_id":"2406.14732","n_code_links":0,"syntology":null},{"paper":null,"slug":"unmasking-database-vulnerabilities-zero","title":"Unmasking Database Vulnerabilities: Zero-Knowledge Schema Inference Attacks in Text-to-SQL Systems","date":"2024-06-20","arxiv_id":"2406.14545","n_code_links":0,"syntology":null},{"paper":"/paper/a-pure-transformer-pretraining-framework-on","slug":"a-pure-transformer-pretraining-framework-on","title":"A Pure Transformer Pretraining Framework on Text-attributed Graphs","date":"2024-06-19","arxiv_id":"2406.13873","n_code_links":1,"syntology":null},{"paper":"/paper/alanavlm-a-multimodal-embodied-ai-foundation","slug":"alanavlm-a-multimodal-embodied-ai-foundation","title":"AlanaVLM: A Multimodal Embodied AI Foundation Model for Egocentric Video Understanding","date":"2024-06-19","arxiv_id":"2406.13807","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["alanaai/evud"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/can-long-context-language-models-subsume","slug":"can-long-context-language-models-subsume","title":"Can Long-Context Language Models Subsume Retrieval, RAG, SQL, and More?","date":"2024-06-19","arxiv_id":"2406.13121","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["google-deepmind/loft"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/do-multimodal-foundation-models-understand","slug":"do-multimodal-foundation-models-understand","title":"WONDERBREAD: A Benchmark for Evaluating Multimodal Foundation Models on Business Process Management Tasks","date":"2024-06-19","arxiv_id":"2406.13264","n_code_links":1,"syntology":{"ran":11,"of":12,"n_ran_checked":11,"n_instrument":0,"unverified":1,"pointer_only":12,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 0 violated, 11 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["hazyresearch/wonderbread"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/efficient-sharpness-aware-minimization-for-2","slug":"efficient-sharpness-aware-minimization-for-2","title":"Efficient Sharpness-Aware Minimization for Molecular Graph Transformer Models","date":"2024-06-19","arxiv_id":"2406.13137","n_code_links":1,"syntology":{"ran":7,"of":7,"n_ran_checked":7,"n_instrument":0,"unverified":0,"pointer_only":7,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["YL-wang/GraphSAM"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/enhancing-language-model-factuality-via","slug":"enhancing-language-model-factuality-via","title":"Enhancing Language Model Factuality via Activation-Based Confidence Calibration and Guided Decoding","date":"2024-06-19","arxiv_id":"2406.13230","n_code_links":1,"syntology":null},{"paper":null,"slug":"fine-tuning-berts-for-definition-extraction","title":"Fine-Tuning BERTs for Definition Extraction from Mathematical Text","date":"2024-06-19","arxiv_id":"2406.13827","n_code_links":0,"syntology":null},{"paper":null,"slug":"forag-factuality-optimized-retrieval","title":"FoRAG: Factuality-optimized Retrieval Augmented Generation for Web-enhanced Long-form Question Answering","date":"2024-06-19","arxiv_id":"2406.13779","n_code_links":0,"syntology":null},{"paper":null,"slug":"guided-context-gating-learning-to-leverage","title":"Guided Context Gating: Learning to leverage salient lesions in retinal fundus images","date":"2024-06-19","arxiv_id":"2406.13126","n_code_links":0,"syntology":null},{"paper":"/paper/instructrag-instructing-retrieval-augmented","slug":"instructrag-instructing-retrieval-augmented","title":"InstructRAG: Instructing Retrieval-Augmented Generation via Self-Synthesized Rationales","date":"2024-06-19","arxiv_id":"2406.13629","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":0,"n_instrument":2,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["weizhepei/instructrag"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"is-gpt-4-conscious","title":"Is GPT-4 conscious?","date":"2024-06-19","arxiv_id":"2407.09517","n_code_links":0,"syntology":null},{"paper":null,"slug":"liveness-detection-in-computer-vision","title":"Liveness Detection in Computer Vision: Transformer-based Self-Supervised Learning for Face Anti-Spoofing","date":"2024-06-19","arxiv_id":"2406.13860","n_code_links":0,"syntology":null},{"paper":null,"slug":"m3t-multi-modal-medical-transformer-to-bridge","title":"M3T: Multi-Modal Medical Transformer to bridge Clinical Context with Visual Insights for Retinal Image Medical Description Generation","date":"2024-06-19","arxiv_id":"2406.13129","n_code_links":0,"syntology":null},{"paper":"/paper/model-internals-based-answer-attribution-for","slug":"model-internals-based-answer-attribution-for","title":"Model Internals-based Answer Attribution for Trustworthy Retrieval-Augmented Generation","date":"2024-06-19","arxiv_id":"2406.13663","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["betswish/mirage"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/morehopqa-more-than-multi-hop-reasoning","slug":"morehopqa-more-than-multi-hop-reasoning","title":"MoreHopQA: More Than Multi-hop Reasoning","date":"2024-06-19","arxiv_id":"2406.13397","n_code_links":1,"syntology":null},{"paper":"/paper/multi-meta-rag-improving-rag-for-multi-hop","slug":"multi-meta-rag-improving-rag-for-multi-hop","title":"Multi-Meta-RAG: Improving RAG for Multi-Hop Queries using Database Filtering with LLM-Extracted Metadata","date":"2024-06-19","arxiv_id":"2406.13213","n_code_links":1,"syntology":null},{"paper":"/paper/on-ai-inspired-ui-design","slug":"on-ai-inspired-ui-design","title":"On AI-Inspired UI-Design","date":"2024-06-19","arxiv_id":"2406.13631","n_code_links":1,"syntology":null},{"paper":null,"slug":"open-generative-large-language-models-for","title":"Open Generative Large Language Models for Galician","date":"2024-06-19","arxiv_id":"2406.13893","n_code_links":0,"syntology":null},{"paper":"/paper/part-aware-unified-representation-of-language-1","slug":"part-aware-unified-representation-of-language-1","title":"Part-aware Unified Representation of Language and Skeleton for Zero-shot Action Recognition","date":"2024-06-19","arxiv_id":"2406.13327","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 1 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["azzh1/purls"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/patholm-identifying-pathogenicity-from-the","slug":"patholm-identifying-pathogenicity-from-the","title":"PathoLM: Identifying pathogenicity from the DNA sequence through the Genome Foundation Model","date":"2024-06-19","arxiv_id":"2406.13133","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["Sajib-006/Patho-LM"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/r-2ag-incorporating-retrieval-information","slug":"r-2ag-incorporating-retrieval-information","title":"R^2AG: Incorporating Retrieval Information into Retrieval Augmented Generation","date":"2024-06-19","arxiv_id":"2406.13249","n_code_links":1,"syntology":{"ran":2,"of":4,"n_ran_checked":0,"n_instrument":2,"unverified":2,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["yefd/RRAG"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/sali-short-term-alignment-and-long-term-1","slug":"sali-short-term-alignment-and-long-term-1","title":"SALI: Short-term Alignment and Long-term Interaction Network for Colonoscopy Video Polyp Segmentation","date":"2024-06-19","arxiv_id":"2406.13532","n_code_links":1,"syntology":null},{"paper":null,"slug":"swinstyleformer-is-a-favorable-choice-for","title":"SwinStyleformer is a favorable choice for image inversion","date":"2024-06-19","arxiv_id":"2406.13153","n_code_links":0,"syntology":null},{"paper":null,"slug":"wikicontradict-a-benchmark-for-evaluating","title":"WikiContradict: A Benchmark for Evaluating LLMs on Real-World Knowledge Conflicts from Wikipedia","date":"2024-06-19","arxiv_id":"2406.13805","n_code_links":0,"syntology":null},{"paper":"/paper/an-investigation-of-neuron-activation-as-a","slug":"an-investigation-of-neuron-activation-as-a","title":"An Investigation of Neuron Activation as a Unified Lens to Explain Chain-of-Thought Eliciting Arithmetic Reasoning of LLMs","date":"2024-06-18","arxiv_id":"2406.12288","n_code_links":2,"syntology":{"ran":34,"of":36,"n_ran_checked":32,"n_instrument":2,"unverified":2,"pointer_only":12,"phrase":"34 ran (of which 0 constructed an object rather than computing a result; 32 with no instrument failure: 0 honoured, 3 violated, 29 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["dakingrai/neuron-analysis-cot-arithmetic-reasoning"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"assessing-ai-vs-human-authored-spear-phishing","title":"Assessing AI vs Human-Authored Spear Phishing SMS Attacks: An Empirical Study","date":"2024-06-18","arxiv_id":"2406.13049","n_code_links":0,"syntology":null},{"paper":"/paper/can-large-language-models-always-solve-easy","slug":"can-large-language-models-always-solve-easy","title":"Can Large Language Models Always Solve Easy Problems if They Can Solve Harder Ones?","date":"2024-06-18","arxiv_id":"2406.12809","n_code_links":1,"syntology":{"ran":11,"of":14,"n_ran_checked":11,"n_instrument":0,"unverified":3,"pointer_only":4,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 0 violated, 11 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["QwenLM/ConsisEval"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/chatglm-a-family-of-large-language-models","slug":"chatglm-a-family-of-large-language-models","title":"ChatGLM: A Family of Large Language Models from GLM-130B to GLM-4 All Tools","date":"2024-06-18","arxiv_id":"2406.12793","n_code_links":7,"syntology":{"ran":21,"of":29,"n_ran_checked":20,"n_instrument":1,"unverified":8,"pointer_only":1,"phrase":"21 ran (of which 0 constructed an object rather than computing a result; 20 with no instrument failure: 0 honoured, 0 violated, 20 with no contract checked; 1 where Syntology's instrument failed) · 8 unverified","official":{"repos":["thudm/chatglm-6b"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/class-specific-data-augmentation-for-plant","slug":"class-specific-data-augmentation-for-plant","title":"Class-specific Data Augmentation for Plant Stress Classification","date":"2024-06-18","arxiv_id":"2406.13081","n_code_links":1,"syntology":null},{"paper":"/paper/cyclic-2-5d-perceptual-loss-for-cross-modal","slug":"cyclic-2-5d-perceptual-loss-for-cross-modal","title":"Cyclic 2.5D Perceptual Loss for Cross-Modal 3D Medical Image Synthesis: T1w MRI to Tau PET","date":"2024-06-18","arxiv_id":"2406.12632","n_code_links":1,"syntology":null},{"paper":"/paper/dart-math-difficulty-aware-rejection-tuning-1","slug":"dart-math-difficulty-aware-rejection-tuning-1","title":"DART-Math: Difficulty-Aware Rejection Tuning for Mathematical Problem-Solving","date":"2024-06-18","arxiv_id":"2407.13690","n_code_links":1,"syntology":{"ran":10,"of":12,"n_ran_checked":10,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["hkust-nlp/dart-math"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"from-rags-to-rich-parameters-probing-how","title":"From RAGs to rich parameters: Probing how language models utilize external knowledge over parametric information for factual queries","date":"2024-06-18","arxiv_id":"2406.12824","n_code_links":0,"syntology":null},{"paper":null,"slug":"generating-educational-materials-with","title":"Generating Educational Materials with Different Levels of Readability using LLMs","date":"2024-06-18","arxiv_id":"2406.12787","n_code_links":0,"syntology":null},{"paper":null,"slug":"generative-artificial-intelligence-guided","title":"Generative Artificial Intelligence-Guided User Studies: An Application for Air Taxi Services","date":"2024-06-18","arxiv_id":"2406.12296","n_code_links":0,"syntology":null},{"paper":"/paper/hierarchical-associative-memory-parallelized","slug":"hierarchical-associative-memory-parallelized","title":"Hierarchical Associative Memory, Parallelized MLP-Mixer, and Symmetry Breaking","date":"2024-06-18","arxiv_id":"2406.12220","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["Toshihiro-Ota/paramixer"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"intermediate-distillation-data-efficient","title":"Intermediate Distillation: Data-Efficient Distillation from Black-Box LLMs for Information Retrieval","date":"2024-06-18","arxiv_id":"2406.12169","n_code_links":0,"syntology":null},{"paper":"/paper/interpretable-preferences-via-multi-objective","slug":"interpretable-preferences-via-multi-objective","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","date":"2024-06-18","arxiv_id":"2406.12845","n_code_links":2,"syntology":{"ran":9,"of":9,"n_ran_checked":9,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["RLHFlow/RLHF-Reward-Modeling"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/ipeval-a-bilingual-intellectual-property","slug":"ipeval-a-bilingual-intellectual-property","title":"IPEval: A Bilingual Intellectual Property Agency Consultation Evaluation Benchmark for Large Language Models","date":"2024-06-18","arxiv_id":"2406.12386","n_code_links":1,"syntology":null},{"paper":"/paper/judging-the-judges-evaluating-alignment-and","slug":"judging-the-judges-evaluating-alignment-and","title":"Judging the Judges: Evaluating Alignment and Vulnerabilities in LLMs-as-Judges","date":"2024-06-18","arxiv_id":"2406.12624","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["UMass-Meta-LLM-Eval/llm_eval"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/measuring-psychological-depth-in-language","slug":"measuring-psychological-depth-in-language","title":"Measuring Psychological Depth in Language Models","date":"2024-06-18","arxiv_id":"2406.12680","n_code_links":1,"syntology":null},{"paper":"/paper/mixing-natural-and-synthetic-images-for","slug":"mixing-natural-and-synthetic-images-for","title":"MixDiff: Mixing Natural and Synthetic Images for Robust Self-Supervised Representations","date":"2024-06-18","arxiv_id":"2406.12368","n_code_links":1,"syntology":null},{"paper":"/paper/pcie-egohandpose-solution-for-egoexo4d-hand","slug":"pcie-egohandpose-solution-for-egoexo4d-hand","title":"PCIE_EgoHandPose Solution for EgoExo4D Hand Pose Challenge","date":"2024-06-18","arxiv_id":"2406.12219","n_code_links":1,"syntology":null},{"paper":"/paper/planrag-a-plan-then-retrieval-augmented","slug":"planrag-a-plan-then-retrieval-augmented","title":"PlanRAG: A Plan-then-Retrieval Augmented Generation for Generative Large Language Models as Decision Makers","date":"2024-06-18","arxiv_id":"2406.12430","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":5,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["myeon9h/planrag"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"recognition-of-dynamic-hand-gestures-in-long","title":"Recognition of Dynamic Hand Gestures in Long Distance using a Web-Camera for Robot Guidance","date":"2024-06-18","arxiv_id":"2406.12424","n_code_links":0,"syntology":null},{"paper":"/paper/restorer-solving-multiple-image-restoration","slug":"restorer-solving-multiple-image-restoration","title":"Restorer: Removing Multi-Degradation with All-Axis Attention and Prompt Guidance","date":"2024-06-18","arxiv_id":"2406.12587","n_code_links":1,"syntology":null},{"paper":null,"slug":"retrieval-augmented-generation-for-generative","title":"Retrieval-Augmented Generation for Generative Artificial Intelligence in Medicine","date":"2024-06-18","arxiv_id":"2406.12449","n_code_links":0,"syntology":null},{"paper":null,"slug":"richrag-crafting-rich-responses-for-multi","title":"RichRAG: Crafting Rich Responses for Multi-faceted Queries in Retrieval-Augmented Generation","date":"2024-06-18","arxiv_id":"2406.12566","n_code_links":0,"syntology":null},{"paper":null,"slug":"sample-efficient-imitative-multi-token","title":"Physics-informed Imitative Reinforcement Learning for Real-world Driving","date":"2024-06-18","arxiv_id":"2407.02508","n_code_links":0,"syntology":null},{"paper":"/paper/towards-a-client-centered-assessment-of-llm","slug":"towards-a-client-centered-assessment-of-llm","title":"Towards a Client-Centered Assessment of LLM Therapists by Client Simulation","date":"2024-06-18","arxiv_id":"2406.12266","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["wangjs9/clientcast"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/ubench-benchmarking-uncertainty-in-large","slug":"ubench-benchmarking-uncertainty-in-large","title":"UBENCH: Benchmarking Uncertainty in Large Language Models with Multiple Choice Questions","date":"2024-06-18","arxiv_id":"2406.12784","n_code_links":1,"syntology":null},{"paper":"/paper/unified-active-retrieval-for-retrieval","slug":"unified-active-retrieval-for-retrieval","title":"Unified Active Retrieval for Retrieval Augmented Generation","date":"2024-06-18","arxiv_id":"2406.12534","n_code_links":1,"syntology":null},{"paper":null,"slug":"urbanllm-autonomous-urban-activity-planning","title":"UrbanLLM: Autonomous Urban Activity Planning and Management with Large Language Models","date":"2024-06-18","arxiv_id":"2406.12360","n_code_links":0,"syntology":null},{"paper":null,"slug":"vernacular-i-barely-know-her-challenges-with","title":"Vernacular? I Barely Know Her: Challenges with Style Control and Stereotyping","date":"2024-06-18","arxiv_id":"2406.12679","n_code_links":0,"syntology":null},{"paper":null,"slug":"what-makes-two-models-think-alike","title":"What Makes Two Language Models Think Alike?","date":"2024-06-18","arxiv_id":"2406.12620","n_code_links":0,"syntology":null},{"paper":null,"slug":"you-gotta-be-a-doctor-lin-an-investigation-of","title":"\"You Gotta be a Doctor, Lin\": An Investigation of Name-Based Bias of Large Language Models in Employment Recommendations","date":"2024-06-18","arxiv_id":"2406.12232","n_code_links":0,"syntology":null},{"paper":"/paper/a-two-dimensional-zero-shot-dialogue-state","slug":"a-two-dimensional-zero-shot-dialogue-state","title":"A Two-dimensional Zero-shot Dialogue State Tracking Evaluation Method using GPT-4","date":"2024-06-17","arxiv_id":"2406.11651","n_code_links":1,"syntology":null},{"paper":null,"slug":"an-exploration-of-length-generalization-in","title":"An Exploration of Length Generalization in Transformer-Based Speech Enhancement","date":"2024-06-17","arxiv_id":"2406.11401","n_code_links":0,"syntology":null},{"paper":"/paper/are-large-language-models-true-healthcare","slug":"are-large-language-models-true-healthcare","title":"Are Large Language Models True Healthcare Jacks-of-All-Trades? Benchmarking Across Health Professions Beyond Physician Exams","date":"2024-06-17","arxiv_id":"2406.11328","n_code_links":1,"syntology":null},{"paper":null,"slug":"attention-based-deep-reinforcement-learning-1","title":"Attention-Based Deep Reinforcement Learning for Qubit Allocation in Modular Quantum Architectures","date":"2024-06-17","arxiv_id":"2406.11452","n_code_links":0,"syntology":null},{"paper":"/paper/beyond-boundaries-learning-a-universal-entity","slug":"beyond-boundaries-learning-a-universal-entity","title":"Beyond Boundaries: Learning a Universal Entity Taxonomy across Datasets and Languages for Open Named Entity Recognition","date":"2024-06-17","arxiv_id":"2406.11192","n_code_links":1,"syntology":null},{"paper":null,"slug":"breaking-boundaries-investigating-the-effects","title":"Breaking Boundaries: Investigating the Effects of Model Editing on Cross-linguistic Performance","date":"2024-06-17","arxiv_id":"2406.11139","n_code_links":0,"syntology":null},{"paper":"/paper/building-another-spanish-dictionary-this-time","slug":"building-another-spanish-dictionary-this-time","title":"Building another Spanish dictionary, this time with GPT-4","date":"2024-06-17","arxiv_id":"2406.11218","n_code_links":1,"syntology":null},{"paper":"/paper/citrus-chunked-instruction-aware-state","slug":"citrus-chunked-instruction-aware-state","title":"CItruS: Chunked Instruction-aware State Eviction for Long Sequence Modeling","date":"2024-06-17","arxiv_id":"2406.12018","n_code_links":1,"syntology":{"ran":5,"of":10,"n_ran_checked":0,"n_instrument":5,"unverified":5,"pointer_only":10,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 5 where Syntology's instrument failed) · 5 unverified","official":{"repos":["ybai-nlp/CItruS"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":"/paper/cram-credibility-aware-attention-modification","slug":"cram-credibility-aware-attention-modification","title":"CrAM: Credibility-Aware Attention Modification in LLMs for Combating Misinformation in RAG","date":"2024-06-17","arxiv_id":"2406.11497","n_code_links":1,"syntology":null},{"paper":null,"slug":"crossfusor-a-cross-attention-transformer","title":"Crossfusor: A Cross-Attention Transformer Enhanced Conditional Diffusion Model for Car-Following Trajectory Prediction","date":"2024-06-17","arxiv_id":"2406.11941","n_code_links":0,"syntology":null},{"paper":null,"slug":"cultural-conditioning-or-placebo-on-the","title":"Cultural Conditioning or Placebo? On the Effectiveness of Socio-Demographic Prompting","date":"2024-06-17","arxiv_id":"2406.11661","n_code_links":0,"syntology":null},{"paper":null,"slug":"decoding-the-narratives-analyzing-personal","title":"Decoding the Narratives: Analyzing Personal Drug Experiences Shared on Reddit","date":"2024-06-17","arxiv_id":"2406.12117","n_code_links":0,"syntology":null},{"paper":"/paper/diffusion-based-adaptation-for-classification","slug":"diffusion-based-adaptation-for-classification","title":"Diffusion-Based Adaptation for Classification of Unknown Degraded Images","date":"2024-06-17","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/ditto-tts-efficient-and-scalable-zero-shot","slug":"ditto-tts-efficient-and-scalable-zero-shot","title":"DiTTo-TTS: Diffusion Transformers for Scalable Text-to-Speech without Domain-Specific Factors","date":"2024-06-17","arxiv_id":"2406.11427","n_code_links":1,"syntology":null},{"paper":null,"slug":"enabling-robots-to-follow-abstract","title":"Enabling robots to follow abstract instructions and complete complex dynamic tasks","date":"2024-06-17","arxiv_id":"2406.11231","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhancing-biomedical-knowledge-retrieval","title":"SeRTS: Self-Rewarding Tree Search for Biomedical Retrieval-Augmented Generation","date":"2024-06-17","arxiv_id":"2406.11258","n_code_links":0,"syntology":null},{"paper":"/paper/enhancing-text-classification-through-llm","slug":"enhancing-text-classification-through-llm","title":"Enhancing Text Classification through LLM-Driven Active Learning and Human Annotation","date":"2024-06-17","arxiv_id":"2406.12114","n_code_links":1,"syntology":null},{"paper":null,"slug":"estimating-the-increase-in-emissions-caused","title":"Estimating the Increase in Emissions caused by AI-augmented Search","date":"2024-06-17","arxiv_id":"2407.16894","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-open-language-models-across-task","slug":"evaluating-open-language-models-across-task","title":"Are Small Language Models Ready to Compete with Large Language Models for Practical Applications?","date":"2024-06-17","arxiv_id":"2406.11402","n_code_links":1,"syntology":null},{"paper":"/paper/evaluating-the-efficacy-of-open-source-llms","slug":"evaluating-the-efficacy-of-open-source-llms","title":"Evaluating the Efficacy of Open-Source LLMs in Enterprise-Specific RAG Systems: A Comparative Study of Performance and Scalability","date":"2024-06-17","arxiv_id":"2406.11424","n_code_links":1,"syntology":null},{"paper":null,"slug":"exploring-safety-utility-trade-offs-in","title":"Exploring Safety-Utility Trade-Offs in Personalized Language Models","date":"2024-06-17","arxiv_id":"2406.11107","n_code_links":0,"syntology":null}],"record_sha256":"bacc0fa4a0103bebe7d5937704e71f21474f4519af74d932faddf98e70e2881a","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}