{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/object-detection/papers/63","list_of":"/task/object-detection","task":"Object Detection","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":63,"pages_in_order":110,"rows_per_page":100,"rows":[6201,6300],"of":10957,"counts":{"archive_papers_tagged":10957,"with_a_code_link":4657,"where_syntology_ran_a_sample":1183,"not_listed_spam_title":0,"listed":10957,"listed_where_code_ran":1183,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1038,"every_run_a_failure_of_syntologys_instrument":145,"listed_with_a_run_with_no_instrument_failure":1038,"listed_every_run_a_failure_of_syntologys_instrument":145,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/object-detection","prev":"/task/object-detection/papers/62","next":"/task/object-detection/papers/64","papers":[{"url":null,"slug":"biswift-bandwidth-orchestrator-for-multi","title":"BiSwift: Bandwidth Orchestrator for Multi-Stream Video Analytics on Edge","date":"2023-12-25","arxiv_id":"2312.15740","repositories_listed":0,"syntology":null},{"url":null,"slug":"end-to-end-3d-object-detection-using-lidar","title":"End-to-End 3D Object Detection using LiDAR Point Cloud","date":"2023-12-24","arxiv_id":"2312.15377","repositories_listed":0,"syntology":null},{"url":null,"slug":"idet3d-towards-efficient-interactive-object","title":"iDet3D: Towards Efficient Interactive Object Detection for LiDAR Point Clouds","date":"2023-12-24","arxiv_id":"2312.15449","repositories_listed":0,"syntology":null},{"url":null,"slug":"pre-trained-trojan-attacks-for-visual","title":"Pre-trained Trojan Attacks for Visual Recognition","date":"2023-12-23","arxiv_id":"2312.15172","repositories_listed":0,"syntology":null},{"url":null,"slug":"scale-optimization-using-evolutionary","title":"Scale Optimization Using Evolutionary Reinforcement Learning for Object Detection on Drone Imagery","date":"2023-12-23","arxiv_id":"2312.15219","repositories_listed":0,"syntology":null},{"url":null,"slug":"context-enhanced-transformer-for-single-image","title":"Context Enhanced Transformer for Single Image Object Detection","date":"2023-12-22","arxiv_id":"2312.14492","repositories_listed":0,"syntology":null},{"url":null,"slug":"explainable-multi-camera-3d-object-detection","title":"Explainable Multi-Camera 3D Object Detection with Transformer-Based Saliency Maps","date":"2023-12-22","arxiv_id":"2312.14606","repositories_listed":0,"syntology":null},{"url":null,"slug":"fm-ov3d-foundation-model-based-cross-modal","title":"FM-OV3D: Foundation Model-based Cross-modal Knowledge Blending for Open-Vocabulary 3D Detection","date":"2023-12-22","arxiv_id":"2312.14465","repositories_listed":0,"syntology":null},{"url":null,"slug":"fred-towards-a-full-rotation-equivariance-in","title":"FRED: Towards a Full Rotation-Equivariance in Aerial Image Object Detection","date":"2023-12-22","arxiv_id":"2401.06159","repositories_listed":0,"syntology":null},{"url":null,"slug":"lift-attend-splat-bird-s-eye-view-camera","title":"Lift-Attend-Splat: Bird's-eye-view camera-lidar fusion using transformers","date":"2023-12-22","arxiv_id":"2312.14919","repositories_listed":0,"syntology":null},{"url":null,"slug":"meaod-model-extraction-attack-against-object","title":"MEAOD: Model Extraction Attack against Object Detectors","date":"2023-12-22","arxiv_id":"2312.14677","repositories_listed":0,"syntology":null},{"url":null,"slug":"timepillars-temporally-recurrent-3d-lidar","title":"TimePillars: Temporally-Recurrent 3D LiDAR Object Detection","date":"2023-12-22","arxiv_id":"2312.17260","repositories_listed":0,"syntology":null},{"url":null,"slug":"low-power-event-based-face-detection-with","title":"Low-power event-based face detection with asynchronous neuromorphic hardware","date":"2023-12-21","arxiv_id":"2312.14261","repositories_listed":0,"syntology":null},{"url":null,"slug":"domain-similarity-perceived-label-assignment","title":"Domain Similarity-Perceived Label Assignment for Domain Generalized Underwater Object Detection","date":"2023-12-20","arxiv_id":"2401.05401","repositories_listed":0,"syntology":null},{"url":null,"slug":"integration-and-performance-analysis-of","title":"Integration and Performance Analysis of Artificial Intelligence and Computer Vision Based on Deep Learning Algorithms","date":"2023-12-20","arxiv_id":"2312.12872","repositories_listed":0,"syntology":null},{"url":null,"slug":"pointenet-a-lightweight-framework-for","title":"PointeNet: A Lightweight Framework for Effective and Efficient Point Cloud Analysis","date":"2023-12-20","arxiv_id":"2312.12743","repositories_listed":0,"syntology":null},{"url":null,"slug":"diffusing-more-objects-for-semi-supervised","title":"Diffusing More Objects for Semi-Supervised Domain Adaptation with Less Labeling","date":"2023-12-19","arxiv_id":"2312.12000","repositories_listed":0,"syntology":null},{"url":null,"slug":"first-qualitative-observations-on-deep","title":"First qualitative observations on deep learning vision model YOLO and DETR for automated driving in Austria","date":"2023-12-19","arxiv_id":"2312.12314","repositories_listed":0,"syntology":null},{"url":null,"slug":"lasa-instance-reconstruction-from-real-scans","title":"LASA: Instance Reconstruction from Real Scans using A Large-scale Aligned Shape Annotation Dataset","date":"2023-12-19","arxiv_id":"2312.12418","repositories_listed":0,"syntology":null},{"url":null,"slug":"object-detection-for-automated-coronary","title":"Object Detection for Automated Coronary Artery Using Deep Learning","date":"2023-12-19","arxiv_id":"2312.12135","repositories_listed":0,"syntology":null},{"url":null,"slug":"transformer-network-for-multi-person-tracking","title":"Transformer Network for Multi-Person Tracking and Re-Identification in Unconstrained Environment","date":"2023-12-19","arxiv_id":"2312.11929","repositories_listed":0,"syntology":null},{"url":null,"slug":"uniondet-union-level-detector-towards-real-1","title":"UnionDet: Union-Level Detector Towards Real-Time Human-Object Interaction Detection","date":"2023-12-19","arxiv_id":"2312.12664","repositories_listed":0,"syntology":null},{"url":null,"slug":"unsupervised-segmentation-of-colonoscopy","title":"Unsupervised Segmentation of Colonoscopy Images","date":"2023-12-19","arxiv_id":"2312.12599","repositories_listed":0,"syntology":null},{"url":null,"slug":"unveiling-spaces-architecturally-meaningful","title":"Unveiling Spaces: Architecturally meaningful semantic descriptions from images of interior spaces","date":"2023-12-19","arxiv_id":"2312.12481","repositories_listed":0,"syntology":null},{"url":null,"slug":"vision-based-automatic-groceries-tracking","title":"Vision-Based Automatic Groceries Tracking System -- Smart Homes","date":"2023-12-19","arxiv_id":"2312.12486","repositories_listed":0,"syntology":null},{"url":null,"slug":"weakly-supervised-open-vocabulary-object","title":"Weakly Supervised Open-Vocabulary Object Detection","date":"2023-12-19","arxiv_id":"2312.12437","repositories_listed":0,"syntology":null},{"url":null,"slug":"diffusion-based-particle-detr-for-bev","title":"Diffusion-Based Particle-DETR for BEV Perception","date":"2023-12-18","arxiv_id":"2312.11578","repositories_listed":0,"syntology":null},{"url":null,"slug":"global-feature-pyramid-network","title":"Global Feature Pyramid Network","date":"2023-12-18","arxiv_id":"2312.11231","repositories_listed":0,"syntology":null},{"url":null,"slug":"language-assisted-3d-scene-understanding","title":"Language-Assisted 3D Scene Understanding","date":"2023-12-18","arxiv_id":"2312.11451","repositories_listed":0,"syntology":null},{"url":null,"slug":"liquid-leak-detection-using-thermal-images","title":"Liquid Leak Detection Using Thermal Images","date":"2023-12-18","arxiv_id":"2312.10980","repositories_listed":0,"syntology":null},{"url":null,"slug":"matchdet-a-collaborative-framework-for-image","title":"MatchDet: A Collaborative Framework for Image Matching and Object Detection","date":"2023-12-18","arxiv_id":"2312.10983","repositories_listed":0,"syntology":null},{"url":null,"slug":"satellite-captioning-large-language-models-to","title":"Satellite Captioning: Large Language Models to Augment Labeling","date":"2023-12-18","arxiv_id":"2312.10905","repositories_listed":0,"syntology":null},{"url":null,"slug":"squeezed-edge-yolo-onboard-object-detection","title":"Squeezed Edge YOLO: Onboard Object Detection on Edge Devices","date":"2023-12-18","arxiv_id":"2312.11716","repositories_listed":0,"syntology":null},{"url":null,"slug":"3daxiesprompts-unleashing-the-3d-spatial-task","title":"3DAxiesPrompts: Unleashing the 3D Spatial Task Capabilities of GPT-4V","date":"2023-12-15","arxiv_id":"2312.09738","repositories_listed":0,"syntology":null},{"url":null,"slug":"deep-active-perception-for-object-detection","title":"Deep Active Perception for Object Detection using Navigation Proposals","date":"2023-12-15","arxiv_id":"2312.10200","repositories_listed":0,"syntology":null},{"url":null,"slug":"end-to-end-training-of-neural-networks-for","title":"End-to-End Training of Neural Networks for Automotive Radar Interference Mitigation","date":"2023-12-15","arxiv_id":"2312.09790","repositories_listed":0,"syntology":null},{"url":null,"slug":"hierarchical-uncertainty-aggregation-and","title":"Hierarchical Uncertainty Aggregation and Emphasis Loss for Active Learning in Object Detection","date":"2023-12-15","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"semantic-aware-transformation-invariant-roi","title":"Semantic-Aware Transformation-Invariant RoI Align","date":"2023-12-15","arxiv_id":"2312.09609","repositories_listed":0,"syntology":null},{"url":null,"slug":"slowtrack-increasing-the-latency-of-camera","title":"SlowTrack: Increasing the Latency of Camera-based Perception in Autonomous Driving Using Adversarial Examples","date":"2023-12-15","arxiv_id":"2312.09520","repositories_listed":0,"syntology":null},{"url":null,"slug":"ada-yolo-dynamic-fusion-of-yolov8-and","title":"ADA-YOLO: Dynamic Fusion of YOLOv8 and Adaptive Heads for Precise Image Detection and Diagnosis","date":"2023-12-14","arxiv_id":"2312.10099","repositories_listed":0,"syntology":null},{"url":null,"slug":"class-wise-buffer-management-for-incremental","title":"Class-Wise Buffer Management for Incremental Object Detection: An Effective Buffer Training Strategy","date":"2023-12-14","arxiv_id":"2312.09139","repositories_listed":0,"syntology":null},{"url":null,"slug":"exploration-of-visual-prompt-in-grounded-pre","title":"Exploration of visual prompt in Grounded pre-trained open-set detection","date":"2023-12-14","arxiv_id":"2312.08839","repositories_listed":0,"syntology":null},{"url":null,"slug":"learned-fusion-3d-object-detection-using","title":"Learned Fusion: 3D Object Detection using Calibration-Free Transformer Feature Fusion","date":"2023-12-14","arxiv_id":"2312.09082","repositories_listed":0,"syntology":null},{"url":null,"slug":"advancements-in-content-based-image-retrieval","title":"Advancements in Content-Based Image Retrieval: A Comprehensive Survey of Relevance Feedback Techniques","date":"2023-12-13","arxiv_id":"2312.10089","repositories_listed":0,"syntology":null},{"url":null,"slug":"an-invitation-to-deep-reinforcement-learning","title":"An Invitation to Deep Reinforcement Learning","date":"2023-12-13","arxiv_id":"2312.08365","repositories_listed":0,"syntology":null},{"url":null,"slug":"challenges-of-yolo-series-for-object","title":"Challenges of YOLO Series for Object Detection in Extremely Heavy Rain: CALRA Simulator based Synthetic Evaluation Dataset","date":"2023-12-13","arxiv_id":"2312.07976","repositories_listed":0,"syntology":null},{"url":null,"slug":"instance-aware-multi-camera-3d-object","title":"Instance-aware Multi-Camera 3D Object Detection with Structural Priors Mining and Self-Boosting Learning","date":"2023-12-13","arxiv_id":"2312.08004","repositories_listed":0,"syntology":null},{"url":null,"slug":"edge-wasserstein-distance-loss-for-oriented","title":"Edge Wasserstein Distance Loss for Oriented Object Detection","date":"2023-12-12","arxiv_id":"2312.07048","repositories_listed":0,"syntology":null},{"url":null,"slug":"ia2u-a-transfer-plugin-with-multi-prior-for","title":"IA2U: A Transfer Plugin with Multi-Prior for In-Air Model to Underwater","date":"2023-12-12","arxiv_id":"2312.06955","repositories_listed":0,"syntology":null},{"url":null,"slug":"lightweight-high-resolution-subject-matting","title":"Lightweight high-resolution Subject Matting in the Real World","date":"2023-12-12","arxiv_id":"2312.07100","repositories_listed":0,"syntology":null},{"url":null,"slug":"opensight-a-simple-open-vocabulary-framework","title":"OpenSight: A Simple Open-Vocabulary Framework for LiDAR-Based Object Detection","date":"2023-12-12","arxiv_id":"2312.08876","repositories_listed":0,"syntology":null},{"url":null,"slug":"teaching-unknown-objects-by-leveraging-human","title":"Teaching Unknown Objects by Leveraging Human Gaze and Augmented Reality in Human-Robot Interaction","date":"2023-12-12","arxiv_id":"2312.07638","repositories_listed":0,"syntology":null},{"url":null,"slug":"mitigating-perspective-distortion-induced","title":"Mitigating Perspective Distortion-induced Shape Ambiguity in Image Crops","date":"2023-12-11","arxiv_id":"2312.06594","repositories_listed":0,"syntology":null},{"url":null,"slug":"simmining-3d-altitude-aware-3d-object","title":"SimMining-3D: Altitude-Aware 3D Object Detection in Complex Mining Environments: A Novel Dataset and ROS-Based Automatic Annotation Pipeline","date":"2023-12-11","arxiv_id":"2312.06113","repositories_listed":0,"syntology":null},{"url":null,"slug":"squeezesam-user-friendly-mobile-interactive","title":"SqueezeSAM: User friendly mobile interactive segmentation","date":"2023-12-11","arxiv_id":"2312.06736","repositories_listed":0,"syntology":null},{"url":null,"slug":"user-friendly-and-adaptable-discriminative-ai","title":"User Friendly and Adaptable Discriminative AI: Using the Lessons from the Success of LLMs and Image Generation Models","date":"2023-12-11","arxiv_id":"2312.06826","repositories_listed":0,"syntology":null},{"url":null,"slug":"investigating-yolo-models-towards-outdoor","title":"Investigating YOLO Models Towards Outdoor Obstacle Detection For Visually Impaired People","date":"2023-12-10","arxiv_id":"2312.07571","repositories_listed":0,"syntology":null},{"url":null,"slug":"open-world-object-detection-in-the-era-of","title":"Open World Object Detection in the Era of Foundation Models","date":"2023-12-10","arxiv_id":"2312.05745","repositories_listed":0,"syntology":null},{"url":null,"slug":"immature-green-apple-detection-and-sizing-in","title":"Immature Green Apple Detection and Sizing in Commercial Orchards using YOLOv8 and Shape Fitting Techniques","date":"2023-12-08","arxiv_id":"2401.08629","repositories_listed":0,"syntology":null},{"url":"/paper/lyrics-boosting-fine-grained-language-vision","slug":"lyrics-boosting-fine-grained-language-vision","title":"Lyrics: Boosting Fine-grained Language-Vision Alignment and Comprehension via Semantic-aware Visual Objects","date":"2023-12-08","arxiv_id":"2312.05278","repositories_listed":0,"syntology":null},{"url":null,"slug":"forcing-generative-models-to-degenerate-ones","title":"Forcing Generative Models to Degenerate Ones: The Power of Data Poisoning Attacks","date":"2023-12-07","arxiv_id":"2312.04748","repositories_listed":0,"syntology":null},{"url":null,"slug":"gen2det-generate-to-detect","title":"Gen2Det: Generate to Detect","date":"2023-12-07","arxiv_id":"2312.04566","repositories_listed":0,"syntology":null},{"url":null,"slug":"multiview-aerial-visual-recognition-mavrec","title":"Multiview Aerial Visual Recognition (MAVREC): Can Multi-view Improve Aerial Visual Perception?","date":"2023-12-07","arxiv_id":"2312.04548","repositories_listed":0,"syntology":null},{"url":null,"slug":"stable-diffusion-for-data-augmentation-in","title":"Stable Diffusion for Data Augmentation in COCO and Weed Datasets","date":"2023-12-07","arxiv_id":"2312.03996","repositories_listed":0,"syntology":null},{"url":null,"slug":"automated-multimodal-data-annotation-via","title":"Automated Multimodal Data Annotation via Calibration With Indoor Positioning System","date":"2023-12-06","arxiv_id":"2312.03608","repositories_listed":0,"syntology":null},{"url":null,"slug":"learning-to-holistically-detect-bridges-from","title":"Learning to Holistically Detect Bridges from Large-Size VHR Remote Sensing Imagery","date":"2023-12-05","arxiv_id":"2312.02481","repositories_listed":0,"syntology":null},{"url":null,"slug":"rotatr-detection-transformer-for-dense-and","title":"RotaTR: Detection Transformer for Dense and Rotated Object","date":"2023-12-05","arxiv_id":"2312.02821","repositories_listed":0,"syntology":null},{"url":null,"slug":"uni3dl-unified-model-for-3d-and-language","title":"Uni3DL: Unified Model for 3D and Language Understanding","date":"2023-12-05","arxiv_id":"2312.03026","repositories_listed":0,"syntology":null},{"url":null,"slug":"improv-inpainting-based-multimodal-prompting","title":"IMProv: Inpainting-based Multimodal Prompting for Computer Vision Tasks","date":"2023-12-04","arxiv_id":"2312.01771","repositories_listed":0,"syntology":null},{"url":null,"slug":"learning-pseudo-labeler-beyond-noun-concepts","title":"Learning Pseudo-Labeler beyond Noun Concepts for Open-Vocabulary Object Detection","date":"2023-12-04","arxiv_id":"2312.02103","repositories_listed":0,"syntology":null},{"url":null,"slug":"survey-on-deep-learning-in-multimodal-medical","title":"Survey on deep learning in multimodal medical imaging for cancer detection","date":"2023-12-04","arxiv_id":"2312.01573","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-review-and-a-robust-framework-of-data","title":"A Review and A Robust Framework of Data-Efficient 3D Scene Parsing with Traditional/Learned 3D Descriptors","date":"2023-12-03","arxiv_id":"2312.01262","repositories_listed":0,"syntology":null},{"url":null,"slug":"exploring-adversarial-robustness-of-lidar","title":"Exploring Adversarial Robustness of LiDAR-Camera Fusion Model in Autonomous Driving","date":"2023-12-03","arxiv_id":"2312.01468","repositories_listed":0,"syntology":null},{"url":null,"slug":"scheme-scalable-channer-mixer-for-vision","title":"SCHEME: Scalable Channel Mixer for Vision Transformers","date":"2023-12-01","arxiv_id":"2312.00412","repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-efficient-3d-object-detection-in-bird","title":"Towards Efficient 3D Object Detection in Bird's-Eye-View Space for Autonomous Driving: A Convolutional-Only Approach","date":"2023-12-01","arxiv_id":"2312.00633","repositories_listed":0,"syntology":null},{"url":null,"slug":"cascaded-interaction-with-eroded-deep","title":"Cascaded Interaction with Eroded Deep Supervision for Salient Object Detection","date":"2023-11-30","arxiv_id":"2311.18675","repositories_listed":0,"syntology":null},{"url":null,"slug":"fool-the-hydra-adversarial-attacks-against","title":"Fool the Hydra: Adversarial Attacks against Multi-view Object Detection Systems","date":"2023-11-30","arxiv_id":"2312.00173","repositories_listed":0,"syntology":null},{"url":null,"slug":"hy-tracker-a-novel-framework-for-enhancing","title":"Hy-Tracker: A Novel Framework for Enhancing Efficiency and Accuracy of Object Tracking in Hyperspectral Videos","date":"2023-11-30","arxiv_id":"2311.18199","repositories_listed":0,"syntology":null},{"url":null,"slug":"simulflow-simultaneously-extracting-feature","title":"SimulFlow: Simultaneously Extracting Feature and Identifying Target for Unsupervised Video Object Segmentation","date":"2023-11-30","arxiv_id":"2311.18286","repositories_listed":0,"syntology":null},{"url":null,"slug":"an-efficient-illumination-invariant-tiger","title":"An Efficient Illumination Invariant Tiger Detection Framework for Wildlife Surveillance","date":"2023-11-29","arxiv_id":"2311.17552","repositories_listed":0,"syntology":null},{"url":null,"slug":"pillarnest-embracing-backbone-scaling-and","title":"PillarNeSt: Embracing Backbone Scaling and Pretraining for Pillar-based 3D Object Detection","date":"2023-11-29","arxiv_id":"2311.17770","repositories_listed":0,"syntology":null},{"url":null,"slug":"weakly-semi-supervised-object-detection-in","title":"Weakly-semi-supervised object detection in remotely sensed imagery","date":"2023-11-29","arxiv_id":"2311.17449","repositories_listed":0,"syntology":null},{"url":null,"slug":"cross-level-attention-with-overlapped-windows","title":"Cross-level Attention with Overlapped Windows for Camouflaged Object Detection","date":"2023-11-28","arxiv_id":"2311.16618","repositories_listed":0,"syntology":null},{"url":null,"slug":"feedback-roi-features-improve-aerial-object","title":"Feedback RoI Features Improve Aerial Object Detection","date":"2023-11-28","arxiv_id":"2311.17129","repositories_listed":0,"syntology":null},{"url":null,"slug":"generative-data-augmentation-improves","title":"ScribbleGen: Generative Data Augmentation Improves Scribble-supervised Semantic Segmentation","date":"2023-11-28","arxiv_id":"2311.17121","repositories_listed":0,"syntology":null},{"url":null,"slug":"integration-of-robotics-computer-vision-and","title":"Integration of Robotics, Computer Vision, and Algorithm Design: A Chinese Poker Self-Playing Robot","date":"2023-11-28","arxiv_id":"2312.09455","repositories_listed":0,"syntology":null},{"url":null,"slug":"large-model-based-referring-camouflaged","title":"Large Model Based Referring Camouflaged Object Detection","date":"2023-11-28","arxiv_id":"2311.17122","repositories_listed":0,"syntology":null},{"url":null,"slug":"model-agnostic-body-part-relevance-assessment","title":"Model-agnostic Body Part Relevance Assessment for Pedestrian Detection","date":"2023-11-27","arxiv_id":"2311.15679","repositories_listed":0,"syntology":null},{"url":null,"slug":"seeing-beyond-cancer-multi-institutional","title":"Seeing Beyond Cancer: Multi-Institutional Validation of Object Localization and 3D Semantic Segmentation using Deep Learning for Breast MRI","date":"2023-11-27","arxiv_id":"2311.16213","repositories_listed":0,"syntology":null},{"url":null,"slug":"small-and-dim-target-detection-in-ir-imagery","title":"Small and Dim Target Detection in IR Imagery: A Review","date":"2023-11-27","arxiv_id":"2311.16346","repositories_listed":0,"syntology":null},{"url":null,"slug":"unified-batch-normalization-identifying-and","title":"Unified Batch Normalization: Identifying and Alleviating the Feature Condensation in Batch Normalization and a Unified Framework","date":"2023-11-27","arxiv_id":"2311.15993","repositories_listed":0,"syntology":null},{"url":null,"slug":"an-intelligent-detection-network-for","title":"An Intelligent-Detection Network for Handwritten Mathematical Expression Recognition","date":"2023-11-26","arxiv_id":"2311.15273","repositories_listed":0,"syntology":null},{"url":null,"slug":"gan-based-lidar-intensity-simulation","title":"GAN-Based LiDAR Intensity Simulation","date":"2023-11-26","arxiv_id":"2311.15415","repositories_listed":0,"syntology":null},{"url":null,"slug":"opennet-incremental-learning-for-autonomous","title":"OpenNet: Incremental Learning for Autonomous Driving Object Detection with Balanced Loss","date":"2023-11-25","arxiv_id":"2311.14939","repositories_listed":0,"syntology":null},{"url":null,"slug":"multi-modal-instance-refinement-for-cross","title":"Multi-modal Instance Refinement for Cross-domain Action Recognition","date":"2023-11-24","arxiv_id":"2311.14281","repositories_listed":0,"syntology":null},{"url":null,"slug":"all-in-one-rgb-rgb-d-and-rgb-t-salient-object","title":"All in One: RGB, RGB-D, and RGB-T Salient Object Detection","date":"2023-11-23","arxiv_id":"2311.14746","repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-transferable-multi-modal-perception","title":"Towards Transferable Multi-modal Perception Representation Learning for Autonomy: NeRF-Supervised Masked AutoEncoder","date":"2023-11-23","arxiv_id":"2311.13750","repositories_listed":0,"syntology":null},{"url":null,"slug":"when-side-channel-attacks-break-the-black-box","title":"When Side-Channel Attacks Break the Black-Box Property of Embedded Artificial Intelligence","date":"2023-11-23","arxiv_id":"2311.14005","repositories_listed":0,"syntology":null},{"url":null,"slug":"benthiq-a-transformer-based-benthic","title":"BenthIQ: a Transformer-Based Benthic Classification Model for Coral Restoration","date":"2023-11-22","arxiv_id":"2311.13661","repositories_listed":0,"syntology":null},{"url":null,"slug":"doubleaug-single-domain-generalized-object","title":"DoubleAUG: Single-domain Generalized Object Detector in Urban via Color Perturbation and Dual-style Memory","date":"2023-11-22","arxiv_id":"2311.13198","repositories_listed":0,"syntology":null}],"record_sha256":"9be89b7f4ede49a2e2c580fa0e0eba25e018ebc9b42c39f153bd03f79fac5b19","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}