{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/multimodal-deep-learning/papers/2","list_of":"/task/multimodal-deep-learning","task":"Multimodal Deep Learning","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":2,"pages_in_order":3,"rows_per_page":100,"rows":[101,200],"of":213,"counts":{"archive_papers_tagged":213,"with_a_code_link":97,"where_syntology_ran_a_sample":21,"not_listed_spam_title":0,"listed":213,"listed_where_code_ran":21,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":16,"every_run_a_failure_of_syntologys_instrument":5,"listed_with_a_run_with_no_instrument_failure":16,"listed_every_run_a_failure_of_syntologys_instrument":5,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/multimodal-deep-learning","prev":"/task/multimodal-deep-learning","next":"/task/multimodal-deep-learning/papers/3","papers":[{"url":null,"slug":"newsnet-sdf-stochastic-discount-factor","title":"NewsNet-SDF: Stochastic Discount Factor Estimation with Pretrained Language Model News Embeddings via Adversarial Networks","date":"2025-05-11","arxiv_id":"2505.06864","repositories_listed":0,"syntology":null},{"url":null,"slug":"bmmdetect-a-multimodal-deep-learning","title":"BMMDetect: A Multimodal Deep Learning Framework for Comprehensive Biomedical Misconduct Detection","date":"2025-05-09","arxiv_id":"2505.05763","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-empowered-beam","title":"Multimodal Deep Learning-Empowered Beam Prediction in Future THz ISAC Systems","date":"2025-05-05","arxiv_id":"2505.02381","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-for-stroke","title":"Multimodal Deep Learning for Stroke Prediction and Detection using Retinal Imaging and Clinical Data","date":"2025-05-05","arxiv_id":"2505.02677","repositories_listed":0,"syntology":null},{"url":null,"slug":"timing-is-everything-finding-the-optimal","title":"Timing Is Everything: Finding the Optimal Fusion Points in Multimodal Medical Imaging","date":"2025-05-05","arxiv_id":"2505.02467","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-doctor-in-the-loop-a-clinically","title":"Multimodal Doctor-in-the-Loop: A Clinically-Guided Explainable Framework for Predicting Pathological Response in Non-Small Cell Lung Cancer","date":"2025-05-02","arxiv_id":"2505.01390","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-multimodal-deep-learning-approach-for-white","title":"A Multimodal Deep Learning Approach for White Matter Shape Prediction in Diffusion MRI Tractography","date":"2025-04-25","arxiv_id":"2504.18400","repositories_listed":0,"syntology":null},{"url":null,"slug":"integrating-vision-and-location-with","title":"Integrating Vision and Location with Transformers: A Multimodal Deep Learning Framework for Medical Wound Analysis","date":"2025-04-14","arxiv_id":"2504.10452","repositories_listed":0,"syntology":null},{"url":null,"slug":"improving-neonatal-care-an-active-dry-contact","title":"Improving Neonatal Care: An Active Dry-Contact Electrode-based Continuous EEG Monitoring System with Seizure Detection","date":"2025-03-30","arxiv_id":"2503.23338","repositories_listed":0,"syntology":null},{"url":null,"slug":"tabulatime-a-novel-multimodal-deep-learning","title":"TabulaTime: A Novel Multimodal Deep Learning Framework for Advancing Acute Coronary Syndrome Prediction through Environmental and Clinical Data Integration","date":"2025-02-24","arxiv_id":"2502.17049","repositories_listed":0,"syntology":null},{"url":null,"slug":"evolution-of-data-driven-single-and-multi","title":"Evolution of Data-driven Single- and Multi-Hazard Susceptibility Mapping and Emergence of Deep Learning Methods","date":"2025-02-13","arxiv_id":"2502.09045","repositories_listed":0,"syntology":null},{"url":"/paper/admn-a-layer-wise-adaptive-multimodal-network","slug":"admn-a-layer-wise-adaptive-multimodal-network","title":"ADMN: A Layer-Wise Adaptive Multimodal Network for Dynamic Input Noise and Compute Resources","date":"2025-02-11","arxiv_id":"2502.07862","repositories_listed":0,"syntology":{"n":4,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/admn-a-layer-wise-adaptive-multimodal-network#ran","syntology_url":"https://syntology.ai/paper/2502.07862","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2502.07862"}},"official":null}},{"url":null,"slug":"a-self-supervised-multimodal-deep-learning","title":"A Self-supervised Multimodal Deep Learning Approach to Differentiate Post-radiotherapy Progression from Pseudoprogression in Glioblastoma","date":"2025-02-06","arxiv_id":"2502.03999","repositories_listed":0,"syntology":null},{"url":null,"slug":"innovative-framework-for-early-estimation-of","title":"Innovative Framework for Early Estimation of Mental Disorder Scores to Enable Timely Interventions","date":"2025-02-06","arxiv_id":"2502.03965","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-prescriptive-deep-learning","title":"Multimodal Prescriptive Deep Learning","date":"2025-01-24","arxiv_id":"2501.14152","repositories_listed":0,"syntology":null},{"url":null,"slug":"reducing-overtreatment-of-indeterminate","title":"Reducing Overtreatment of Indeterminate Thyroid Nodules Using a Multimodal Deep Learning Model","date":"2024-09-27","arxiv_id":"2409.19171","repositories_listed":0,"syntology":null},{"url":null,"slug":"validation-exploration-of-multimodal-deep","title":"Validation & Exploration of Multimodal Deep-Learning Camera-Lidar Calibration models","date":"2024-09-20","arxiv_id":"2409.13402","repositories_listed":0,"syntology":null},{"url":null,"slug":"2408-02686","title":"A Systematic Review of Intermediate Fusion in Multimodal Deep Learning for Biomedical Applications","date":"2024-08-02","arxiv_id":"2408.02686","repositories_listed":0,"syntology":null},{"url":null,"slug":"audio-visual-approach-for-multimodal","title":"Audio-Visual Approach For Multimodal Concurrent Speaker Detection","date":"2024-07-01","arxiv_id":"2407.01774","repositories_listed":0,"syntology":null},{"url":null,"slug":"advanced-multimodal-deep-learning","title":"Advanced Multimodal Deep Learning Architecture for Image-Text Matching","date":"2024-06-13","arxiv_id":"2406.15306","repositories_listed":0,"syntology":null},{"url":null,"slug":"research-on-optimization-of-natural-language","title":"Research on Optimization of Natural Language Processing Model Based on Multimodal Deep Learning","date":"2024-06-13","arxiv_id":"2406.08838","repositories_listed":0,"syntology":null},{"url":null,"slug":"carbonsense-a-multimodal-dataset-and-baseline","title":"CarbonSense: A Multimodal Dataset and Baseline for Carbon Flux Modelling","date":"2024-06-07","arxiv_id":"2406.04940","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-for-low-resource","title":"Multimodal Deep Learning for Low-Resource Settings: A Vector Embedding Alignment Approach for Healthcare Applications","date":"2024-06-02","arxiv_id":"2406.02601","repositories_listed":0,"syntology":null},{"url":null,"slug":"comprehensive-multimodal-deep-learning","title":"Comprehensive Multimodal Deep Learning Survival Prediction Enabled by a Transformer Architecture: A Multicenter Study in Glioblastoma","date":"2024-05-21","arxiv_id":"2405.12963","repositories_listed":0,"syntology":null},{"url":null,"slug":"the-role-of-emotions-in-informational-support","title":"The Role of Emotions in Informational Support Question-Response Pairs in Online Health Communities: A Multimodal Deep Learning Approach","date":"2024-05-21","arxiv_id":"2405.13099","repositories_listed":0,"syntology":null},{"url":null,"slug":"research-on-image-recognition-technology","title":"Research on Image Recognition Technology Based on Multimodal Deep Learning","date":"2024-05-06","arxiv_id":"2405.03091","repositories_listed":0,"syntology":null},{"url":null,"slug":"ecc-analyzer-extract-trading-signal-from","title":"ECC Analyzer: Extract Trading Signal from Earnings Conference Calls using Large Language Model for Stock Performance Prediction","date":"2024-04-29","arxiv_id":"2404.18470","repositories_listed":0,"syntology":null},{"url":null,"slug":"describe-and-dissect-interpreting-neurons-in","title":"Describe-and-Dissect: Interpreting Neurons in Vision Networks with Language Models","date":"2024-03-20","arxiv_id":"2403.13771","repositories_listed":0,"syntology":null},{"url":null,"slug":"integrating-wearable-sensor-data-and-self","title":"Integrating Wearable Sensor Data and Self-reported Diaries for Personalized Affect Forecasting","date":"2024-03-16","arxiv_id":"2403.13841","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-multimodal-intermediate-fusion-network-with","title":"A Multimodal Intermediate Fusion Network with Manifold Learning for Stress Detection","date":"2024-03-12","arxiv_id":"2403.08077","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-approach-to","title":"Multimodal deep learning approach to predicting neurological recovery from coma after cardiac arrest","date":"2024-03-09","arxiv_id":"2403.06027","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-learning-to-improve-cardiac-late","title":"Multimodal Learning To Improve Cardiac Late Mechanical Activation Detection From Cine MR Images","date":"2024-02-28","arxiv_id":"2402.18507","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-of-word-of-mouth","title":"Multimodal Deep Learning of Word-of-Mouth Text and Demographics to Predict Customer Rating: Handling Consumer Heterogeneity in Marketing","date":"2024-01-22","arxiv_id":"2401.11888","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-learning-for-detecting-urban","title":"Multimodal Urban Areas of Interest Generation via Remote Sensing Imagery and Geographical Prior","date":"2024-01-12","arxiv_id":"2401.06550","repositories_listed":0,"syntology":null},{"url":null,"slug":"predicting-the-skies-a-novel-model-for-flight","title":"Predicting the Skies: A Novel Model for Flight-Level Passenger Traffic Forecasting","date":"2024-01-07","arxiv_id":"2401.03397","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-self-supervised-learning-for-1","title":"Multimodal self-supervised learning for lesion localization","date":"2024-01-03","arxiv_id":"2401.01524","repositories_listed":0,"syntology":null},{"url":null,"slug":"integrating-chemical-language-and-molecular","title":"Integrating Chemical Language and Molecular Graph in Multimodal Fused Deep Learning for Drug Property Prediction","date":"2023-12-29","arxiv_id":"2312.17495","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-graph-based-multimodal-framework-to-predict","title":"A graph-based multimodal framework to predict gentrification","date":"2023-12-25","arxiv_id":"2312.15646","repositories_listed":0,"syntology":null},{"url":null,"slug":"synthscribe-deep-multimodal-tools-for","title":"SynthScribe: Deep Multimodal Tools for Synthesizer Sound Retrieval and Exploration","date":"2023-12-07","arxiv_id":"2312.04690","repositories_listed":0,"syntology":null},{"url":null,"slug":"textaug-test-time-text-augmentation-for","title":"TextAug: Test time Text Augmentation for Multimodal Person Re-identification","date":"2023-12-04","arxiv_id":"2312.01605","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-for-mapping-forest","title":"Multimodal deep learning for mapping forest dominant height by fusing GEDI with earth observation data","date":"2023-11-20","arxiv_id":"2311.11777","repositories_listed":0,"syntology":null},{"url":null,"slug":"malfake-a-multimodal-fake-news-identification","title":"MalFake: A Multimodal Fake News Identification for Malayalam using Recurrent Neural Networks and VGG-16","date":"2023-10-27","arxiv_id":"2310.18263","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-for-scientific","title":"Multimodal Deep Learning for Scientific Imaging Interpretation","date":"2023-09-21","arxiv_id":"2309.12460","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-multimodal-deep-learning-architecture-for","title":"A multimodal deep learning architecture for smoking detection with a small data approach","date":"2023-09-19","arxiv_id":"2309.10561","repositories_listed":0,"syntology":null},{"url":null,"slug":"arc-nlp-at-multimodal-hate-speech-event","title":"ARC-NLP at Multimodal Hate Speech Event Detection 2023: Multimodal Methods Boosted by Ensemble Learning, Syntactical and Entity Features","date":"2023-07-25","arxiv_id":"2307.13829","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-scoping-review-on-multimodal-deep-learning","title":"A scoping review on multimodal deep learning in biomedical images and texts","date":"2023-07-14","arxiv_id":"2307.07362","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-for-personalized","title":"Multimodal Deep Learning for Personalized Renal Cell Carcinoma Prognosis: Integrating CT Imaging and Clinical Data","date":"2023-07-07","arxiv_id":"2307.03575","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-novel-site-agnostic-multimodal-deep","title":"A Novel Site-Agnostic Multimodal Deep Learning Model to Identify Pro-Eating Disorder Content on Social Media","date":"2023-07-06","arxiv_id":"2307.06775","repositories_listed":0,"syntology":null},{"url":null,"slug":"performance-optimization-using-multimodal","title":"Performance Optimization using Multimodal Modeling and Heterogeneous GNN","date":"2023-04-25","arxiv_id":"2304.12568","repositories_listed":0,"syntology":null},{"url":null,"slug":"empowering-ai-drug-discovery-with-explicit","title":"Towards Unified AI Drug Discovery with Multiple Knowledge Modalities","date":"2023-04-17","arxiv_id":"2305.01523","repositories_listed":0,"syntology":null},{"url":null,"slug":"medical-intervention-duration-estimation","title":"P-Transformer: A Prompt-based Multimodal Transformer Architecture For Medical Tabular Data","date":"2023-03-30","arxiv_id":"2303.17408","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-comprehensive-and-versatile-multimodal-deep","title":"A Comprehensive and Versatile Multimodal Deep Learning Approach for Predicting Diverse Properties of Advanced Materials","date":"2023-03-29","arxiv_id":"2303.16412","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-to-differentiate","title":"Multimodal Deep Learning to Differentiate Tumor Recurrence from Treatment Effect in Human Glioblastoma","date":"2023-02-27","arxiv_id":"2302.14124","repositories_listed":0,"syntology":null},{"url":null,"slug":"show-me-your-nft-and-i-tell-you-how-it-will","title":"Show me your NFT and I tell you how it will perform: Multimodal representation learning for NFT selling price prediction","date":"2023-02-03","arxiv_id":"2302.01676","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-survey-on-knowledge-enhanced-multimodal","title":"A survey on knowledge-enhanced multimodal learning","date":"2022-11-19","arxiv_id":"2211.12328","repositories_listed":0,"syntology":null},{"url":null,"slug":"problem-behaviors-recognition-in-videos-using","title":"Language-Assisted Deep Learning for Autistic Behaviors Recognition","date":"2022-11-17","arxiv_id":"2211.09310","repositories_listed":0,"syntology":null},{"url":null,"slug":"multicrossvit-multimodal-vision-transformer","title":"MultiCrossViT: Multimodal Vision Transformer for Schizophrenia Prediction using Structural MRI and Functional Network Connectivity Data","date":"2022-11-12","arxiv_id":"2211.06726","repositories_listed":0,"syntology":null},{"url":null,"slug":"identification-of-cognitive-workload-during","title":"Identification of Cognitive Workload during Surgical Tasks with Multimodal Deep Learning","date":"2022-09-12","arxiv_id":"2209.06208","repositories_listed":0,"syntology":null},{"url":null,"slug":"r2d2-at-semeval-2022-task-5-attention-is-only","title":"R2D2 at SemEval-2022 Task 5: Attention is only as good as its Values! A multimodal system for identifying misogynist memes","date":"2022-07-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"vision-aided-frame-capture-based-csi","title":"Vision-Aided Frame-Capture-Based CSI Recomposition for WiFi Sensing: A Multimodal Approach","date":"2022-06-03","arxiv_id":"2206.01414","repositories_listed":0,"syntology":null},{"url":null,"slug":"detection-of-propaganda-techniques-in-visuo","title":"Detection of Propaganda Techniques in Visuo-Lingual Metaphor in Memes","date":"2022-05-03","arxiv_id":"2205.02937","repositories_listed":0,"syntology":null},{"url":null,"slug":"multi-objective-optimization-determines-when","title":"Multi-objective optimization determines when, which and how to fuse deep networks: an application to predict COVID-19 outcomes","date":"2022-04-07","arxiv_id":"2204.03772","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-review-on-methods-and-applications-in","title":"A Review on Methods and Applications in Multimodal Deep Learning","date":"2022-02-18","arxiv_id":"2202.09195","repositories_listed":0,"syntology":null},{"url":null,"slug":"emotion-based-hate-speech-detection-using","title":"Emotion Based Hate Speech Detection using Multimodal Learning","date":"2022-02-13","arxiv_id":"2202.06218","repositories_listed":0,"syntology":null},{"url":null,"slug":"geometric-multimodal-deep-learning-with-multi","title":"Geometric Multimodal Deep Learning with Multi-Scaled Graph Wavelet Convolutional Network","date":"2021-11-26","arxiv_id":"2111.13361","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-approach-for-metadata-extraction","title":"Multimodal Approach for Metadata Extraction from German Scientific Publications","date":"2021-11-10","arxiv_id":"2111.05736","repositories_listed":0,"syntology":null},{"url":null,"slug":"from-multimodal-to-unimodal-attention-in","title":"From Multimodal to Unimodal Attention in Transformers using Knowledge Distillation","date":"2021-10-15","arxiv_id":"2110.08270","repositories_listed":0,"syntology":null},{"url":null,"slug":"deepstroke-an-efficient-stroke-screening","title":"DeepStroke: An Efficient Stroke Screening Framework for Emergency Rooms with Multimodal Adversarial Deep Learning","date":"2021-09-24","arxiv_id":"2109.12065","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-co-learning-challenges","title":"Multimodal Co-learning: Challenges, Applications with Datasets, Recent Advances and Future Directions","date":"2021-07-29","arxiv_id":"2107.13782","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-multimodal-deep-learning-model-for-cardiac","title":"A Multimodal Deep Learning Model for Cardiac Resynchronisation Therapy Response Prediction","date":"2021-07-20","arxiv_id":"2107.10662","repositories_listed":0,"syntology":null},{"url":null,"slug":"deep-learning-for-technical-document","title":"Deep Learning for Technical Document Classification","date":"2021-06-27","arxiv_id":"2106.14269","repositories_listed":0,"syntology":null},{"url":null,"slug":"listen-to-your-favorite-melodies-with","title":"Listen to Your Favorite Melodies with img2Mxml, Producing MusicXML from Sheet Music Image by Measure-based Multimodal Deep Learning-driven Assembly","date":"2021-06-16","arxiv_id":"2106.12037","repositories_listed":0,"syntology":null},{"url":null,"slug":"deepmmsa-a-novel-multimodal-deep-learning","title":"DeepMMSA: A Novel Multimodal Deep Learning Method for Non-small Cell Lung Cancer Survival Analysis","date":"2021-06-12","arxiv_id":"2106.06744","repositories_listed":0,"syntology":null},{"url":null,"slug":"digital-taxonomist-identifying-plant-species","title":"Digital Taxonomist: Identifying Plant Species in Community Scientists' Photographs","date":"2021-06-07","arxiv_id":"2106.03774","repositories_listed":0,"syntology":null},{"url":null,"slug":"how-to-select-and-use-tools-active-perception","title":"How to select and use tools? : Active Perception of Target Objects Using Multimodal Deep Learning","date":"2021-06-04","arxiv_id":"2106.02445","repositories_listed":0,"syntology":null},{"url":null,"slug":"recent-advances-and-trends-in-multimodal-deep","title":"Recent Advances and Trends in Multimodal Deep Learning: A Review","date":"2021-05-24","arxiv_id":"2105.11087","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-framework-for-image","title":"Multimodal Deep Learning Framework for Image Popularity Prediction on Social Media","date":"2021-05-18","arxiv_id":"2105.08809","repositories_listed":0,"syntology":null},{"url":null,"slug":"where-and-when-space-time-attention-for-audio","title":"Where and When: Space-Time Attention for Audio-Visual Explanations","date":"2021-05-04","arxiv_id":"2105.01517","repositories_listed":0,"syntology":null},{"url":null,"slug":"the-influence-of-audio-on-video-memorability","title":"The Influence of Audio on Video Memorability with an Audio Gestalt Regulated Video Memorability System","date":"2021-04-23","arxiv_id":"2104.11568","repositories_listed":0,"syntology":null},{"url":null,"slug":"leveraging-audio-gestalt-to-predict-media","title":"Leveraging Audio Gestalt to Predict Media Memorability","date":"2020-12-31","arxiv_id":"2012.15635","repositories_listed":0,"syntology":null},{"url":null,"slug":"predicting-online-video-advertising-effects","title":"Predicting Online Video Advertising Effects with Multimodal Deep Learning","date":"2020-12-22","arxiv_id":"2012.11851","repositories_listed":0,"syntology":null},{"url":null,"slug":"multi-modal-detection-of-alzheimer-s-disease","title":"Multi-Modal Detection of Alzheimer's Disease from Speech and Text","date":"2020-11-30","arxiv_id":"2012.00096","repositories_listed":0,"syntology":null},{"url":null,"slug":"exploring-multimodal-features-and-fusion","title":"Exploring Multimodal Features and Fusion Strategies for Analyzing Disaster Tweets","date":"2020-11-06","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"m2d-a-multi-modal-framework-for-automatic","title":"M2D: A Multi-modal Framework for Automatic Medical Diagnosis","date":"2020-10-19","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"new-ideas-and-trends-in-deep-multimodal","title":"New Ideas and Trends in Deep Multimodal Content Understanding: A Review","date":"2020-10-16","arxiv_id":"2010.08189","repositories_listed":0,"syntology":null},{"url":null,"slug":"using-neural-architecture-search-for","title":"Using Neural Architecture Search for Improving Software Flaw Detection in Multimodal Deep Learning Models","date":"2020-09-22","arxiv_id":"2009.10644","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-for-flaw-detection","title":"Multimodal Deep Learning for Flaw Detection in Software Programs","date":"2020-09-09","arxiv_id":"2009.04549","repositories_listed":0,"syntology":null},{"url":null,"slug":"data-driven-geophysics-from-dictionary","title":"Data-driven geophysics: from dictionary learning to deep learning","date":"2020-07-13","arxiv_id":"2007.06183","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-unfolding-for-guided-image","title":"Multimodal Deep Unfolding for Guided Image Super-Resolution","date":"2020-01-21","arxiv_id":"2001.07575","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-multimodal-deep-learning-approach-for-named","title":"A multimodal deep learning approach for named entity recognition from social media","date":"2020-01-19","arxiv_id":"2001.06888","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-intelligence-representation","title":"Multimodal Intelligence: Representation Learning, Information Fusion, and Applications","date":"2019-11-10","arxiv_id":"1911.03977","repositories_listed":0,"syntology":null},{"url":null,"slug":"detecting-deception-in-political-debates","title":"Detecting Deception in Political Debates Using Acoustic and Textual Features","date":"2019-10-04","arxiv_id":"1910.01990","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-for-mental-disorders","title":"Multimodal Deep Learning for Mental Disorders Prediction from Audio Speech Samples","date":"2019-09-03","arxiv_id":"1909.01067","repositories_listed":0,"syntology":null},{"url":null,"slug":"toxicity-prediction-by-multimodal-deep","title":"Toxicity Prediction by Multimodal Deep Learning","date":"2019-07-19","arxiv_id":"1907.08333","repositories_listed":0,"syntology":null},{"url":null,"slug":"deep-coupled-representation-learning-for","title":"Deep Coupled-Representation Learning for Sparse Linear Inverse Problems with Side Information","date":"2019-07-04","arxiv_id":"1907.02511","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-for-finance","title":"Multimodal Deep Learning for Finance: Integrating and Forecasting International Stock Markets","date":"2019-03-15","arxiv_id":"1903.06478","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-deep-learning-for-short-term-stock","title":"Multimodal deep learning for short-term stock volatility prediction","date":"2018-12-25","arxiv_id":"1812.10479","repositories_listed":0,"syntology":null},{"url":null,"slug":"hybrid-attention-based-multimodal-network-for","title":"Hybrid Attention based Multimodal Network for Spoken Language Classification","date":"2018-08-01","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"correlation-net-spatio-temporal-multimodal","title":"Correlation Net: Spatiotemporal multimodal deep learning for action recognition","date":"2018-07-22","arxiv_id":"1807.08291","repositories_listed":0,"syntology":null},{"url":null,"slug":"fine-grained-video-attractiveness-prediction","title":"Fine-grained Video Attractiveness Prediction Using Multimodal Deep Learning on a Large Real-world Dataset","date":"2018-04-04","arxiv_id":"1804.01373","repositories_listed":0,"syntology":null}],"record_sha256":"4799bb16e3602a3fb9264dcc4a02fedd020b5e83c689bc3eaff7f557603b9c86","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}