{"generated_at":"2026-10-10T09:53:13Z","total":1413,"page":0,"pages":15,"page_size":100,"next":"/api/catalog?topic=computer-vision&page=1","documents":[{"id":"doi-c5b64415e6da6f4234fb","title":"Comparative Evaluation of Object Detection and Semantic Segmentation for Diffuse Tunnel Defects Using UAV Imagery","year":2026,"authors":["Leandro Silva de Assis","Jefferson R. de Souza","Renata Virna Lourenco Balbino da Silva"],"author_count":5,"journal":"IEEE Access","doi":"10.1109/access.2026.3738528","source":"doaj","topics":["computer-vision","aerospace-drones"],"keywords":["Tunnel inspection","UAV","DeepLabv3+","ConvNeXt_Tiny","structural pathologies","Electrical engineering. Electronics. Nuclear engineering"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-c5b64415e6da6f4234fb","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-0287f8c25e31f094eda5","title":"HFSC-Net: A Hierarchical-Feature-Guided Spatial-Frequency Dual-Domain Collaborative Network for Remote Sensing Image Semantic Segmentation","year":2026,"authors":["Zizhi Yuan","Quan Zhang","Guotong Geng"],"author_count":13,"journal":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing","doi":"10.1109/jstars.2026.3735023","source":"doaj","topics":["computer-vision","geoscience"],"keywords":["Hierarchical feature extraction framework","remote sensing image semantic segmentation","semantic prior injection","spatial-frequency dual-domain collaborative decoding (SFCD)","structural context modeling","Ocean engineering","Geophysics. Cosmic physics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-0287f8c25e31f094eda5","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-7093479265e8681d4f2d","title":"Comparative analysis of image augmentation methods for improving the quality of computer vision models in animal husbandry","year":2026,"authors":["Irina I. Mikhailenko"],"author_count":1,"journal":"Цифровые модели и решения","doi":"10.29141/2949-477x-2026-5-3-1","source":"doaj","topics":["computer-vision"],"keywords":["augmentation strategies","image transformation","machine learning","neural networks","limited data.","Home economics","Economics as a science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-7093479265e8681d4f2d","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-7fcebb6fc43d71576376","title":"Q-SODA: High-precision small object detection via feature refinement in hilbert space using cross-channel interaction VQCs","year":2026,"authors":["N.Venkatesvara rao","D. Susitra","L. Sharmila"],"author_count":4,"journal":"Results in Engineering","doi":"10.1016/j.rineng.2026.113306","source":"doaj","topics":["computer-vision"],"keywords":["Hybrid quantum-classical learning","Small object detection","Variational quantum circuits (VQC)","Pascal VOC 2012","Hilbert space mapping","Technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-7fcebb6fc43d71576376","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-d09646058e257949c0e2","title":"Automated neutron dosimetry with bubble detectors using CNN-based object detection","year":2026,"authors":["Xin-yan Li","De-bin Zou","Yan-qing Deng"],"author_count":14,"journal":"Nuclear Engineering and Technology","doi":"10.1016/j.net.2026.104493","source":"doaj","topics":["computer-vision"],"keywords":["Neutron yield","Neutron dosimetry","Bubble detector","Convolutional neural network","Object detection","Nuclear engineering. Atomic power"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-d09646058e257949c0e2","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-c0874715c9a0c366f8ab","title":"Multi-Attribute Small Object Detection Based on Feature Fusion","year":2026,"authors":["XIN Zhenwei, ZHANG Lin, WANG Baolong, SHI Shaohua, YANG Yipeng, CHENG Hui, CHEN Tao"],"author_count":1,"journal":"Jisuanji gongcheng","doi":"10.19678/j.issn.1000-3428.0252596","source":"doaj","topics":["computer-vision"],"keywords":["deep learning|computer vision|object detection|multi-attribute|small object","Computer engineering. Computer hardware","Computer software"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-c0874715c9a0c366f8ab","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-804efd8aa22ea96e3efd","title":"Automated Multi-class Classification of Lung Diseases from CT-scans Using Computer Vision","year":2026,"authors":["Abdul Aleem","Farzana Siddique","Ghulam Gilanie"],"author_count":5,"journal":"Sudan Journal of Medical Sciences","doi":"10.18502/sjms.v21i3.19945","source":"doaj","topics":["computer-vision"],"keywords":["Computed Tomography","COVID-19","Fibrosis","Lung Diseases","Pulmonary Nodule","Pure Consolidation","Medicine"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-804efd8aa22ea96e3efd","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-fc694b065eac08eac723","title":"Corrections to “A Review of Computer Vision-Based Monitoring Approaches for Construction Workers’ Work-Related Behaviors”","year":2026,"authors":["Jiaqi Li","Qi Miao","Zheng Zou"],"author_count":7,"journal":"IEEE Access","doi":"10.1109/access.2026.3738164","source":"doaj","topics":["computer-vision"],"keywords":["Electrical engineering. Electronics. Nuclear engineering"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-fc694b065eac08eac723","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-3e601f9a33507dd270ee","title":"Prototype-Guided Structural Refinement for Weakly Supervised Semantic Segmentation of Industrial Aqueduct Images","year":2026,"authors":["Yangjie Wu","Hongqiu Zhu"],"author_count":2,"journal":"IEEE Access","doi":"10.1109/access.2026.3726966","source":"doaj","topics":["computer-vision"],"keywords":["Semantic segmentation of industrial heritage imagery","frequency-calibrated prototypes","weakly supervised semantic segmentation","Electrical engineering. Electronics. Nuclear engineering"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-3e601f9a33507dd270ee","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-95c38495a55fa3e4fbd4","title":"An Integrated Real-Time 3D Perception Pipeline for Object Detection, Depth, and Orientation Estimation","year":2026,"authors":["Naeem Ul Islam","Dharnish Raja","Ehtisham Muhammad Khan"],"author_count":5,"journal":"IEEE Access","doi":"10.1109/access.2026.3731049","source":"doaj","topics":["computer-vision"],"keywords":["3D perception","object detection","monocular depth estimation","orientation prediction","autonomous driving","robotics","Electrical engineering. Electronics. Nuclear engineering"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-95c38495a55fa3e4fbd4","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-cbb21236e2102b5fdc4a","title":"RSO-DETR: Multiscale Feature Integration and Spatial Pixel Calibration for Robust Small Object Detection in UAV Images","year":2026,"authors":["Jie Pan","Yunzhe Hu","Xiaoyu Zou"],"author_count":3,"journal":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing","doi":"10.1109/jstars.2026.3732076","source":"doaj","topics":["computer-vision","aerospace-drones"],"keywords":["Multiscale feature integration","RT-DETR","small object detection","spatial pixel interaction","uncrewed aerial vehicle (UAV) images","Ocean engineering","Geophysics. Cosmic physics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-cbb21236e2102b5fdc4a","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-5fe9bc1bf771f78d156a","title":"RSFF: a reconciliation-sparse fusion framework for semantic segmentation of rocky desertification land in multimodal time series","year":2026,"authors":["Chaokang He","Qinjun Wang","Boqi Yuan"],"author_count":4,"journal":"GIScience & Remote Sensing","doi":"10.1080/15481603.2026.2739002","source":"doaj","topics":["computer-vision"],"keywords":["Rocky desertification land","multimodal","time series","deep learning","fragmented parcels","Mathematical geography. Cartography","Environmental sciences"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-5fe9bc1bf771f78d156a","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-15e81c3369a689d0bf02","title":"Automated Alignment Powered by Computer Vision Streamlines the Two‐Photon Polymerization‐Based Micro 3D Printing of Multiscale and Multimaterial Structures","year":2026,"authors":["Daniel Maher","Valentina Mendoza Martinez","Marcin Piotr Piekarczyk"],"author_count":5,"journal":"Advanced Intelligent Discovery","doi":"10.1002/aidi.202500188","source":"doaj","topics":["computer-vision"],"keywords":["3D printing","4D printing","automated alignment","computer vision","hierarchical structures","Electronic computers. Computer science","Technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-15e81c3369a689d0bf02","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-110d46a08bc4e08787f6","title":"Computer Vision Pipeline for Image Analysis for Freeze‐Fracture Electron Microscopy: Rosette Cellulose Synthase Complexes Case","year":2026,"authors":["Siri Mudunuri","Leala Carbonneau","Eric M. Roberts"],"author_count":7,"journal":"Advanced Intelligent Discovery","doi":"10.1002/aidi.202500116","source":"doaj","topics":["computer-vision"],"keywords":["cellulose synthase","computer vision","deep learning","FF‐TEM","Electronic computers. Computer science","Technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-110d46a08bc4e08787f6","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-28c4b2476246adedbda1","title":"MMLN: Multidirectional and Multiconstraint Learning Network for Remote Sensing Imagery Semantic Segmentation","year":2026,"authors":["Hang Sun","Yongchang Xie","Dong Ren"],"author_count":6,"journal":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing","doi":"10.1109/jstars.2024.3403854","source":"doaj","topics":["computer-vision","geoscience"],"keywords":["Feature discrimination","multiconstraint","multidirectional","remote sensing imagery semantic segmentation","semantic consistency","Ocean engineering","Geophysics. Cosmic physics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-28c4b2476246adedbda1","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-5a23af3a255e1852a2c5","title":"ParaT-UNet: A Parallel-Structured Temporal-Aware UNet for Burned Area Semantic Segmentation From Aerial Video-Derived Frames","year":2026,"authors":["Yuan Feng","Yang Zhang","Xu Huang"],"author_count":4,"journal":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing","doi":"10.1109/jstars.2026.3728674","source":"doaj","topics":["computer-vision"],"keywords":["Deep learning","semantic segmentation","temporal consistency","UAV imagery","UNet","video-based analysis","Ocean engineering","Geophysics. Cosmic physics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-5a23af3a255e1852a2c5","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-1a6bd39c987f48cff0b0","title":"SparseMamba: Efficient Long-Range Context Modeling for LiDAR Semantic Segmentation With State-Space Models","year":2026,"authors":["Zhiyu Liu","Xiangsen Zhu","Fengli Li"],"author_count":6,"journal":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing","doi":"10.1109/jstars.2026.3733097","source":"doaj","topics":["computer-vision"],"keywords":["Autonomous driving","feature alignment","light detection and ranging (LiDAR) semantic segmentation","Mamba","sparse convolution","state-space models (SSMs)","Ocean engineering","Geophysics. Cosmic physics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-1a6bd39c987f48cff0b0","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-4fdd28739480fe276c9e","title":"Multistage Uncertainty-Aware Feature Fusion Framework for Semantic Segmentation of High-Resolution Remote Sensing Imagery","year":2026,"authors":["Aisha Javed","Youkyung Han"],"author_count":2,"journal":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing","doi":"10.1109/jstars.2026.3733099","source":"doaj","topics":["computer-vision","geoscience"],"keywords":["Encoder–decoder networks","feature modulation","high-resolution remote sensing imagery","multistage feature fusion","semantic segmentation","uncertainty-aware learning","Ocean engineering","Geophysics. Cosmic physics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-4fdd28739480fe276c9e","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-82ee48260aab40b35897","title":"A salient object detection model based on feature attention refinement","year":2026,"authors":["Deepak Kumar Khare","Amit Bhagat","Vishnu Priya R."],"author_count":5,"journal":"Connection Science","doi":"10.1080/09540091.2026.2741505","source":"doaj","topics":["computer-vision"],"keywords":["Salient object detection","attention mechanism","contextual feature integration","feature selection","layer-wise feature refinement","Information technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-82ee48260aab40b35897","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-03317546918fd970218e","title":"Improved diffusion model with dedicated pooling and normalized Wasserstein distance for small object detection in UAV images","year":2026,"authors":["Zhu Haijiang","Liu Mengting","Wang Xiao"],"author_count":3,"journal":"Open Computer Science","doi":"10.1515/comp-2025-0077","source":"doaj","topics":["computer-vision","aerospace-drones"],"keywords":["uav image","small object detection","diffusion model","feature pyramid","normalized wasserstein distance (nwd)","Electronic computers. Computer science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-03317546918fd970218e","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-4e9088b3e4b378aa8078","title":"Comparative Analysis of Marker-based vs. Marker-less Computer Vision for a Gamified Lower Limb Rehabilitation","year":2026,"authors":["Shah Nazar Peiman","Jaber Mohamad","Pott Peter P."],"author_count":3,"journal":"Current Directions in Biomedical Engineering","doi":"10.1515/cdbme-2026-0226","source":"doaj","topics":["computer-vision"],"keywords":["marker-based tracking","computer vision","gamification","lower limb rehabilitation","mediapipe","hsv segmentation","Medicine"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-4e9088b3e4b378aa8078","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-bcb095f4e69e440dbef0","title":"Research on fashion object detection based on cross-domain adaptive learning","year":2026,"authors":["Zhu Shiqian","Liu Xiaogang"],"author_count":2,"journal":"AUTEX Research Journal","doi":"10.1515/aut-2025-0098","source":"doaj","topics":["computer-vision"],"keywords":["fashion image analysis","object detection","deep learning","adaptive consistency","domain transfer learning","Textile bleaching, dyeing, printing, etc."],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-bcb095f4e69e440dbef0","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-739c9256f1961b5e420d","title":"A Survey of Computer Vision Applications in Industrial Environments","year":2026,"authors":["Abdullah Albanyan","Reem Abdel-Salam","Abdulrahman Alabduljabbar"],"author_count":3,"journal":"Applied Artificial Intelligence","doi":"10.1080/08839514.2026.2739540","source":"doaj","topics":["computer-vision"],"keywords":["Electronic computers. Computer science","Cybernetics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-739c9256f1961b5e420d","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-236cbdacef1bff381096","title":"LGI-UNet: a robust semantic segmentation framework with local–global interaction for urban green space mapping in mountainous cities","year":2026,"authors":["Fang Zhou","Guangbin Yang","Man Li"],"author_count":6,"journal":"All Earth","doi":"10.1080/27669645.2026.2741485","source":"doaj","topics":["computer-vision"],"keywords":["Mountainous cities","UGS mapping","LGI-UNet","cross-scale modelling","Geology","Physical geography"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-236cbdacef1bff381096","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-305019c328f92139e8f7","title":"AMFNet: adaptive multi-scale fusion network for small object detection","year":2026,"authors":["Bao Tian","MengNan Hu","LingLing Li"],"author_count":6,"journal":"Complex & Intelligent Systems","doi":"10.1007/s40747-026-02369-2","source":"doaj","topics":["computer-vision"],"keywords":["Small object detection","Feature enhancement","Multi-scale feature fusion","Attention mechanism","Deformable convolution","Electronic computers. Computer science","Information technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-305019c328f92139e8f7","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-f00ec7bda7cf5b5da4b9","title":"Foreign object detection of cable pipeline based on stereo vision","year":2026,"authors":["Ye Lu","Qianxiang Meng","Xingming Feng"],"author_count":6,"journal":"Discover Artificial Intelligence","doi":"10.1007/s44163-026-01691-5","source":"doaj","topics":["computer-vision"],"keywords":["Foreign objects detection of cable pipeline","Target size measurement","Second-order semi-global matching","Saliency detection","Feature fusion","Conditional random field","Computational linguistics. Natural language processing","Electronic computers. Computer science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-f00ec7bda7cf5b5da4b9","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-00fb04cff9654fd136c4","title":"A novel hybrid quantum dilated convolutional Kronecker network (QDCKN) for marine object detection and classification","year":2026,"authors":["N Kumaran","Shyam Mohan J S","Thamaraiselvi D"],"author_count":3,"journal":"Scientific Reports","doi":"10.1038/s41598-026-58876-2","source":"doaj","topics":["computer-vision"],"keywords":["Underwater object detection","Quantum dilated convolutional neural network (QDCNN)","Deep Kronecker network (DKN)","Fuzzy logic","YOLO v3","Savitzky–Golay filtering","Medicine","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-00fb04cff9654fd136c4","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-8abe5a7e04a4eef37ff1","title":"Aerial small object detection via dynamic convolution and hierarchical attention fusion for UAV imagery","year":2026,"authors":["Junxia Zhang","Hao Zhong","Gang Du"],"author_count":5,"journal":"Scientific Reports","doi":"10.1038/s41598-026-59033-5","source":"doaj","topics":["computer-vision","aerospace-drones"],"keywords":["UAV","Small object detection","Dynamic convolution","Hierarchical attention","Aerial imagery","RT-DETR","Medicine","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-8abe5a7e04a4eef37ff1","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-64716b295e6609b8bf0d","title":"An attention-guided adaptive multi-scale feature fusion method for remote sensing image object detection","year":2026,"authors":["Xiaopei Liu","Feng Tian"],"author_count":2,"journal":"Scientific Reports","doi":"10.1038/s41598-026-70241-x","source":"doaj","topics":["computer-vision","geoscience"],"keywords":["Remote sensing target detection","Attention mechanism","Multi-scale feature fusion","Medicine","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-64716b295e6609b8bf0d","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-fa4330b99a6215604ef9","title":"DS-PRNet: lightweight camouflaged object detection with dynamic sparse attention and feature-, attention-, and output-level distillation","year":2026,"authors":["Xiaoyi Wang","Bin Li","Mingqiang Zhang"],"author_count":7,"journal":"Scientific Reports","doi":"10.1038/s41598-026-66102-2","source":"doaj","topics":["computer-vision"],"keywords":["Camouflaged object detection","Lightweight neural network","Dynamic sparse attention","Knowledge distillation","Vision transformer","Medicine","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-fa4330b99a6215604ef9","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-cfebbb8a167553424655","title":"Rand transformer net: An efficient network for semantic segmentation of railway engineering entities based on 3D point cloud","year":2026,"authors":["Xi Chen","Liu Yang","Han Bao"],"author_count":7,"journal":"Scientific Reports","doi":"10.1038/s41598-026-58286-4","source":"doaj","topics":["computer-vision"],"keywords":["Semantic segmentation","Deep learning","3D point cloud","Self-attention","Random downsampling","Medicine","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-cfebbb8a167553424655","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-940cff71b4d17880ce9c","title":"SDS-YOLO: drone-based foreign object detection model for power lines using an enhanced YOLOv8n approach","year":2026,"authors":["Jing Sheng","Shuliang Wu","Guoman Liu"],"author_count":6,"journal":"Scientific Reports","doi":"10.1038/s41598-026-59200-8","source":"doaj","topics":["computer-vision","aerospace-drones"],"keywords":["YOLOv8n","Unmanned aerial vehicle","Power line foreign object detection","SPPF","Dynamic detection head","Attention mechanism","Medicine","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-940cff71b4d17880ce9c","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-4cd488ab8368b4d976aa","title":"Industrial Printed Circuit Board Surface Defect Dataset for Object Detection","year":2026,"authors":["Hao Yan","Xiaoguang Yu","Bifang Ma"],"author_count":9,"journal":"Scientific Data","doi":"10.1038/s41597-026-07684-4","source":"doaj","topics":["computer-vision"],"keywords":["Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-4cd488ab8368b4d976aa","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-70f257835d33cb5e06e6","title":"Artificial Intelligence–Based Computer Vision Feedback to Improve Clinical Performance in Surgical Technology Education","year":2026,"authors":["KIARASH KAMBOOZIA","SEDIGHEH HANNANI","NAZANIN SARRAF SHAHRI"],"author_count":5,"journal":"Journal of Advances in Medical Education and Professionalism","doi":"10.30476/jamp.2026.111362.2415","source":"doaj","topics":["computer-vision"],"keywords":["artificial intelligence","surgical instrument","medical education","Education (General)","Medicine (General)"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-70f257835d33cb5e06e6","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-854d8cc04bfb9c16ca0a","title":"Impact of Fλ/d on Object Detection Performance in Infrared Imaging Systems","year":2026,"authors":["Bingbing Ma","Yu Wang","Sheng Liao"],"author_count":4,"journal":"IEEE Photonics Journal","doi":"10.1109/jphot.2026.3723528","source":"doaj","topics":["computer-vision"],"keywords":["Infrared imaging system","system performance evaluation","Fλ/d","system optimization","Applied optics. Photonics","Optics. Light"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-854d8cc04bfb9c16ca0a","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-d0bed4e4355144855c33","title":"Object detection for autonomous driving environmental perception and its applications","year":2026,"authors":["Xu Xu","Yang Liu","Huijun Yao"],"author_count":3,"journal":"Frontiers in Mechanical Engineering","doi":"10.3389/fmech.2026.1940224","source":"doaj","topics":["computer-vision"],"keywords":["autonomous driving","environmental perception","image recognition","object detection","point cloud processing","sensors","Mechanical engineering and machinery"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-d0bed4e4355144855c33","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-ac8bda603f3a7f83ab02","title":"Crop–Weed Multispectral Semantic Segmentation With Foundation Models and Hardness-Aware Learning","year":2026,"authors":["Ilias Papadeas","Georgios Miltiadis Sandalis","Georgios Zamanakos"],"author_count":4,"journal":"IEEE Access","doi":"10.1109/access.2026.3729808","source":"doaj","topics":["computer-vision"],"keywords":["Crop–Weed semantic segmentation","foundation models","hardness-aware learning","multispectral imagery","precision agriculture","semantic segmentation","Electrical engineering. Electronics. Nuclear engineering"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-ac8bda603f3a7f83ab02","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-b1104cd887c362a919b5","title":"Panoptic Segmentation in Computer Vision: An Overview of Current Approaches and Future Trends","year":2026,"authors":["null Ramyashree","Shambhavi Jha","B. N. Anoop"],"author_count":5,"journal":"Journal of Engineering","doi":"10.1155/je/4051869","source":"doaj","topics":["computer-vision"],"keywords":["Engineering (General). Civil engineering (General)"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-b1104cd887c362a919b5","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-885f055071d1f07e992b","title":"Seeing the Room, Not the Student: A Privacy-Preserving Computer Vision Prototype for Anonymous Aggregate Classroom Visual-Engagement Trend Monitoring from Face-Visibility and Head-Orientation Cues","year":2026,"authors":["Hamzah Asyrani  bin Sulaiman","Nazreen bin Abdullasim","Mohamad Lutfi  bin Dolhalit"],"author_count":4,"journal":"Journal of Engineering","doi":"10.31026/j.eng.2026.10.01","source":"doaj","topics":["computer-vision"],"keywords":["Classroom analytics","Computer vision","Educational technology","Head-orientation estimation","Privacy-preserving AI","Data minimization","Engineering (General). Civil engineering (General)"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-885f055071d1f07e992b","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-b83992208e7eebe5b0b8","title":"DCAFD-YOLOv8: Dynamic Class-Aware Focal Feature Distillation for Class-Imbalanced Object Detection","year":2026,"authors":["Qi Yang","Cong Zhang","Lili Li"],"author_count":5,"journal":"Journal of Applied Science and Engineering","doi":"10.6180/jase.202612_35.057","source":"doaj","topics":["computer-vision"],"keywords":["lightweight object detection","class imbalance","focal feature distillation","dynamic class weighting","teacher confidence","Engineering (General). Civil engineering (General)","Chemical engineering","Physics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-b83992208e7eebe5b0b8","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-b567c6002277d3d6d725","title":"A dual-level evaluation framework for anatomical data augmentation: A comparative study on L2 vertebra semantic segmentation","year":2026,"authors":["Luca Di Angelo","Emanuele Guardiani","Tamsir Jobe"],"author_count":6,"journal":"Machine Learning with Applications","doi":"10.1016/j.mlwa.2026.101019","source":"doaj","topics":["computer-vision"],"keywords":["Data augmentation","Semantic segmentation","Medical imaging analysis","Deep learning","Cybernetics","Electronic computers. Computer science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-b567c6002277d3d6d725","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-fab1bff373d0c7710f5e","title":"TinyML-Driven Smart Real-Time Multi-Class Object Detection on Resource-Constrained Standalone Device","year":2026,"authors":["Satishkumar Kataria","Pankaj kumar Prajapati"],"author_count":2,"journal":"ITEGAM-JETIA","doi":"10.5935/jetia.v12i61.4617","source":"doaj","topics":["computer-vision"],"keywords":["Technology","Technology (General)","Science (General)"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-fab1bff373d0c7710f5e","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-84c00109421261ab7d4d","title":"AUTONOMOUS SEARCH AND RESCUE DRONE WITH REAL-TIME YOLO OBJECT DETECTION AND GPS NAVIGATION USING RASPBERRY PI 4","year":2026,"authors":["Ulpan Turmaganbet","Dana Turlykozhayeva","Sayat Akhtanov"],"author_count":6,"journal":"Eurasian Physical Technical Journal","doi":"10.31489/2026n3/116-125","source":"doaj","topics":["computer-vision"],"keywords":["Unmanned aerial vehicles","search and rescue","object detection","computer vision","autonomous navigation","embedded systems","Physics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-84c00109421261ab7d4d","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-d1469173edec62ab10ae","title":"Research on Street Space Age-Friendliness of Lanzhou City Based on Street View Semantic Segmentation","year":2026,"authors":["Hua LIAN","Bitao HU"],"author_count":2,"journal":"兰州交通大学学报","doi":"10.3969/j.issn.2096-9066.2025.074","source":"doaj","topics":["computer-vision"],"keywords":["street view semantic segmentation","poi data","street space","age-friendliness","lanzhou city","Engineering (General). Civil engineering (General)","Electrical engineering. Electronics. Nuclear engineering"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-d1469173edec62ab10ae","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-52ca5bee4386c5b06f8f","title":"Lightweight UAV-based object detection and tracking for intelligent oil and gas field safety monitoring.","year":2026,"authors":["Jianhua Gong","Yifu Wang","Jun Zhang"],"author_count":5,"journal":"PLoS ONE","doi":"10.1371/journal.pone.0358757","source":"doaj","topics":["computer-vision"],"keywords":["Medicine","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-52ca5bee4386c5b06f8f","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-6190eb2b875bcb5c5eab","title":"Computer Vision System and Transformer-Based Multimodal Fusion of Medical Images for Enhanced Diagnostic Decision Making","year":2026,"authors":["soltan Abdullah","Kahlan F . Aljobory"],"author_count":2,"journal":"Wasit Journal of Computer and Mathematics Science","doi":"10.31185/wjcms.543","source":"doaj","topics":["computer-vision"],"keywords":["Deep learning , computer vision , Medical Images , Diagnostic Decision","Electronic computers. Computer science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-6190eb2b875bcb5c5eab","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-cb8bad410bed7d47355d","title":"Semantic Segmentation of Structural Thermal Bridges Via DTEP-Enhanced UNET++: A Comparative Study on V6 And V7 Architectures","year":2026,"authors":["O. H. Özdemir","E. Açmaz","Ö. H. Bettemir"],"author_count":4,"journal":"The International Archives of the Photogrammetry, Remote Sensing and Spatial Information Sciences","doi":"10.5194/isprs-archives-l-4-w3-2026-123-2026","source":"doaj","topics":["computer-vision"],"keywords":["Technology","Engineering (General). Civil engineering (General)","Applied optics. Photonics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-cb8bad410bed7d47355d","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-e4dfc81623fb1fa5d426","title":"MBD-YOLO: A Multi-Path Coordinate Attention and Boundary-Aware Dual-Stream Fusion YOLO for UAV Small Object Detection","year":2026,"authors":["Jiajun Chen","Yongming Li","Shaochen Jiang"],"author_count":4,"journal":"IEEE Access","doi":"10.1109/access.2026.3736695","source":"doaj","topics":["computer-vision"],"keywords":["UAV aerial imagery","small object detection","YOLOv8","dual-stream attention","coordinate attention","boundary-aware supervision","Electrical engineering. Electronics. Nuclear engineering"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-e4dfc81623fb1fa5d426","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-ac3f172c4084b7a7161e","title":"Enhancing semantic segmentation of construction scenes using IFC-derived synthetic point clouds for geometric Digital Twins","year":2026,"authors":["I. Chauhan","D. Abrol","D. Abrol"],"author_count":4,"journal":"ISPRS Annals of the Photogrammetry, Remote Sensing and Spatial Information Sciences","doi":"10.5194/isprs-annals-xii-4-w1-2026-73-2026","source":"doaj","topics":["computer-vision"],"keywords":["Technology","Engineering (General). Civil engineering (General)","Applied optics. Photonics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-ac3f172c4084b7a7161e","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-2dcf2f05ad1d48954f70","title":"Scale-Invariant Object Contour Points (SIOCP): Integrating Instance Segmentation, Computer Vision and Geometric Refinement for Exact Contour and Keypoint Extraction of Façade Elements from Urban Images","year":2026,"authors":["F. Frank","V. Shah","L. Hoegner"],"author_count":4,"journal":"ISPRS Annals of the Photogrammetry, Remote Sensing and Spatial Information Sciences","doi":"10.5194/isprs-annals-xii-4-w1-2026-121-2026","source":"doaj","topics":["computer-vision"],"keywords":["Technology","Engineering (General). Civil engineering (General)","Applied optics. Photonics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-2dcf2f05ad1d48954f70","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-44a82fe6b33a14782516","title":"A Comparative Analysis of Four Semantic Segmentation Models for Classification of Building Structural Elements in Highly Occluded Environments","year":2026,"authors":["M. Akhoundi Khezrabad","D. Shojaei","M. Tomko"],"author_count":3,"journal":"ISPRS Annals of the Photogrammetry, Remote Sensing and Spatial Information Sciences","doi":"10.5194/isprs-annals-xii-4-w1-2026-1-2026","source":"doaj","topics":["computer-vision"],"keywords":["Technology","Engineering (General). Civil engineering (General)","Applied optics. Photonics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-44a82fe6b33a14782516","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-5c23d3166557b307f92f","title":"Integrating Ontologies and Object Compositional Hierarchies for Dynamic Semantic Segmentation of Point Clouds","year":2026,"authors":["M. Codiglione","M. Codiglione","S. Facenda"],"author_count":4,"journal":"The International Archives of the Photogrammetry, Remote Sensing and Spatial Information Sciences","doi":"10.5194/isprs-archives-l-4-w2-2026-33-2026","source":"doaj","topics":["computer-vision"],"keywords":["Technology","Engineering (General). Civil engineering (General)","Applied optics. Photonics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-5c23d3166557b307f92f","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-7c8b43d378586e66170f","title":"ASSESSMENT OF BIOMECHANICAL PARAMETERS OF PHYSICAL EXERCISES BASED ON COMPUTER VISION","year":2026,"authors":["A. K. Kereyev","K. P. Aman","Z. K. Kalkabayeva"],"author_count":5,"journal":"Шәкәрім университетінің хабаршысы: Техника ғылымдар","doi":"10.53360/2788-7995-2026-2(22)-1","source":"doaj","topics":["computer-vision"],"keywords":["artificial intelligence","computer vision","deep learning","yolo architecture","biomechanical analysis","exercise monitoring","personalized training","Technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-7c8b43d378586e66170f","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-0346f2d4a1f7ed8f62a0","title":"Computer vision workflows for physical carrying capacity analysis and monitoring in linear recreation corridors","year":2026,"authors":["Jordan W. Smith","Chase C. Lamborn"],"author_count":2,"journal":"Machine Learning with Applications","doi":"10.1016/j.mlwa.2026.101012","source":"doaj","topics":["computer-vision"],"keywords":["Computer vision","Semantic segmentation","Object detection","Carrying capacity","Recreation corridors","Parking use","Cybernetics","Electronic computers. Computer science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-0346f2d4a1f7ed8f62a0","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-fe47c8d811da5648515d","title":"LDANet: a lightweight depth-aware framework for RGB-D salient object detection","year":2026,"authors":["Baoyu Wang","Pingping Cao","Xiaochun Guo"],"author_count":4,"journal":"Journal of King Saud University: Computer and Information Sciences","doi":"10.1007/s44443-026-01295-0","source":"doaj","topics":["computer-vision"],"keywords":["Salient object detection","Modality fusion module","Deep guided attention","Adaptive edge refine module","Hierarchical feat fusion decoder","Electronic computers. Computer science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-fe47c8d811da5648515d","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-26b90b27f2baef063d04","title":"An improved YOLO11 based object detection algorithm for UAV maritime rescue","year":2026,"authors":["Jing Wang","Lele Niu","Rongliang Shi"],"author_count":4,"journal":"Discover Artificial Intelligence","doi":"10.1007/s44163-026-01620-6","source":"doaj","topics":["computer-vision"],"keywords":["YOLO11","UAV","Maritime rescue","Small object detection","Loss function","Convolution","Computational linguistics. Natural language processing","Electronic computers. Computer science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-26b90b27f2baef063d04","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-70dd9a2987f34485ad9d","title":"Deblurring-aware semantic segmentation of crops and weeds in UAV sorghum imagery via a UNet-ResNet architecture","year":2026,"authors":["Qiyue Li","Shenshen Zou"],"author_count":2,"journal":"Scientific Reports","doi":"10.1038/s41598-026-58738-x","source":"doaj","topics":["computer-vision"],"keywords":["Semantic segmentation","Crop–weed discrimination","UNet-ResNet architecture","Image deblurring","Precision agriculture","Medicine","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-70dd9a2987f34485ad9d","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-9f6059f6454b6b2e9d7d","title":"SFL-YOLO: an improved YOLOv11-based model for underwater object detection","year":2026,"authors":["Xiaokang Wang","Yuanjiang Li","Qingzhi Zu"],"author_count":4,"journal":"Scientific Reports","doi":"10.1038/s41598-026-60347-7","source":"doaj","topics":["computer-vision"],"keywords":["Deep learning","Underwater object detection","YOLOv11","SFL-YOLO","Medicine","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-9f6059f6454b6b2e9d7d","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-45e18d636fda65157cc7","title":"Explainable image aesthetic assessment via semantic segmentation-guided cross-attention fusion","year":2026,"authors":["Hao Fang","Chengcai Cao","Zixi Huang"],"author_count":5,"journal":"Scientific Reports","doi":"10.1038/s41598-026-66430-3","source":"doaj","topics":["computer-vision"],"keywords":["Image aesthetic assessment (IAA)","Semantic segmentation","Cross-attention fusion","Explainable AI (XAI)","SHAP values","Medicine","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-45e18d636fda65157cc7","fulltext_endpoint":null,"references_endpoint":null},{"id":"doaj-ab63dc25f4244f758f3d2104f8b9b463","title":"Traffic object detection for fixed-camera expressway surveillance under degraded imaging conditions","year":2026,"authors":["Wang Liangliang","Xu Shen","Wang Xuexin"],"author_count":6,"journal":"Journal of Measurement Science and Instrumentation","doi":null,"source":"doaj","topics":["computer-vision"],"keywords":["roadside monitoring","weak-target perception","fine-scale prediction","context-aware reassembly of features (CARAFE)","YOLOv5s","real-time inference","Physics","Engineering (General). Civil engineering (General)"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doaj-ab63dc25f4244f758f3d2104f8b9b463","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-52dc99c7564e0c8e13e0","title":"Corrections to “Research on Road Object Detection Model Based on YOLOv4 of Autonomous Vehicle”","year":2026,"authors":["Penghui Wang","Xufei Wang","Yifan Liu"],"author_count":4,"journal":"IEEE Access","doi":"10.1109/access.2026.3735218","source":"doaj","topics":["computer-vision"],"keywords":["Electrical engineering. Electronics. Nuclear engineering"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-52dc99c7564e0c8e13e0","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-0f99c1ad53306dd61bb5","title":"Evaluating Structured Pruning and Quantization Strategies for Edge-Efficient Traffic Object Detection","year":2026,"authors":["Wang Haojing"],"author_count":1,"journal":"ITM Web of Conferences","doi":"10.1051/itmconf/20269103007","source":"doaj","topics":["computer-vision"],"keywords":["Information technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-0f99c1ad53306dd61bb5","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-58ef3f81151eadb31b38","title":"An adaptive density weighting algorithm for semantic segmentation in large-scale oblique photogrammetric point clouds","year":2026,"authors":["Kunyang Ma","Zheng Zhang","Wen Ge"],"author_count":7,"journal":"Geocarto International","doi":"10.1080/10106049.2026.2717473","source":"doaj","topics":["computer-vision"],"keywords":["Density awareness","deep learning","large-scale application","oblique photogrammetric point cloud","point cloud semantic segmentation","Physical geography"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-58ef3f81151eadb31b38","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-8715fe365e2d5ac503ff","title":"Automatic construction productivity monitoring based on object detection using pose-augmented images","year":2026,"authors":["Ahmad Muhtarom","Hung-Ming Chen","Chi-Yu Chiang"],"author_count":4,"journal":"KSCE Journal of Civil Engineering","doi":"10.1016/j.kscej.2026.100752","source":"doaj","topics":["computer-vision"],"keywords":["Pose-augmented RGB images","Construction activity recognition","Object detection","Real-time monitoring","Productivity analysis","Engineering (General). Civil engineering (General)"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-8715fe365e2d5ac503ff","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-86fb01006a45cb385803","title":"WFFM-TA-YOLO: A unified weighted feature fusion module with P2-Aware multi-branch structure for remote sensing small object detection","year":2026,"authors":["Te Qi","Jing Tian","Fengjie Zheng"],"author_count":5,"journal":"Array","doi":"10.1016/j.array.2026.101210","source":"doaj","topics":["computer-vision"],"keywords":["Remote sensing image detection","Small object detection","Weighted feature fusion module (WFFM)","P2-aware multi-branch structure","Triplet attention","Computer engineering. Computer hardware","Electronic computers. Computer science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-86fb01006a45cb385803","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-2606a24ac8e2e90478a4","title":"Modosaic: A multimodal experimentation platform for computer-vision dataset generation and validation","year":2026,"authors":["O. Agost","P. Fraile","J. Rius"],"author_count":7,"journal":"SoftwareX","doi":"10.1016/j.softx.2026.102938","source":"doaj","topics":["computer-vision"],"keywords":["Multimodal datasets","Computer vision","Data validation","Segmentation","Depth estimation","Reproducibility","Computer software"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-2606a24ac8e2e90478a4","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-2cb036345c58798999a7","title":"Pi-Noise-Guided Student–Teacher Distillation Detector for Open-Vocabulary Aerial Object Detection","year":2026,"authors":["Shu Tian","Zhendong Huang","Lin Cao"],"author_count":8,"journal":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing","doi":"10.1109/jstars.2026.3727402","source":"doaj","topics":["computer-vision"],"keywords":["Aerial object detection","data augmentation","open-vocabulary detection","positive-incentive noise (Pi-noise)","semantic alignment","Ocean engineering","Geophysics. Cosmic physics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-2cb036345c58798999a7","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-b6bc42fb791112c140e2","title":"Development and Validation of a Computer Vision-Based Artificial Intelligence System (FenoParasite) for the Microscopic Detection of Ascaris lumbricoides and Giardia lamblia in Stool Samples","year":2026,"authors":["Miguel Hueda-Zavaleta","Exequiel Federico Espeche","Francisco Zea Gamboa"],"author_count":6,"journal":"Tropical Medicine and Infectious Disease","doi":"10.3390/tropicalmed11090246","source":"doaj","topics":["computer-vision"],"keywords":["artificial intelligence","computer vision","stool microscopy","<i>Ascaris lumbricoides</i>","<i>Giardia lamblia</i>","parasitic diseases","Medicine"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-b6bc42fb791112c140e2","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-e1c536bb7e131a29f2c0","title":"A Multi-Generational YOLO Ensemble with Weighted Boxes Fusion for Robust Rescue-Oriented Object Detection in Chaotic Disaster Scenes","year":2026,"authors":["Ming-Hseng Tseng","Yi-Wei Huang"],"author_count":2,"journal":"Technologies","doi":"10.3390/technologies14090530","source":"doaj","topics":["computer-vision"],"keywords":["rescue-oriented object detection","disaster scene analysis","multi-generational YOLO","ensemble learning","weighted boxes fusion","emergency response","Technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-e1c536bb7e131a29f2c0","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-60fab37eab195dbc2c2e","title":"Underwater Computer Vision for Ecologists: A Framework for Curating Image Training Datasets for Object Detection","year":2026,"authors":["Talen Rimmer","Colin Bates","Declan McIntosh"],"author_count":6,"journal":"Sensors","doi":"10.3390/s26185869","source":"doaj","topics":["computer-vision"],"keywords":["ecology","biodiversity","object detection","underwater imagery","image annotation","annotation methods","Chemical technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-60fab37eab195dbc2c2e","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-7e675b028439a32c2b32","title":"RPT-Fusion: A Time-Lag-Aware Quality-Adaptive Radar–Camera Fusion Framework for Water-Surface Object Detection","year":2026,"authors":["Yabin Xu","Sujie Zhan","Junnan Yang"],"author_count":5,"journal":"Sensors","doi":"10.3390/s26185860","source":"doaj","topics":["computer-vision"],"keywords":["water-surface object detection","radar–camera fusion","4D millimeter-wave radar","probabilistic radar prior","temporal reliability modeling","quality-adaptive fusion","Chemical technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-7e675b028439a32c2b32","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-44518a288e0982be9d98","title":"Bridging the Scale Gap: A Multi-Scale Feature Enhancement Framework for UAV Aerial Image Object Detection","year":2026,"authors":["Dan Shan","Zanqi Qiu","Dadi Cai"],"author_count":6,"journal":"Sensors","doi":"10.3390/s26185832","source":"doaj","topics":["computer-vision"],"keywords":["enhanced feature connection module","feature channel attention","Reinforced Attention Feedback","unified query supervision loss","Chemical technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-44518a288e0982be9d98","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-7d728842130e9ae30fc7","title":"MHF-Net: Multi-Modal Residual Fusion Object Detection Network Based on High-Frequency Gated Convolution","year":2026,"authors":["Bailin Chen","Bo Qian"],"author_count":2,"journal":"Sensors","doi":"10.3390/s26185816","source":"doaj","topics":["computer-vision"],"keywords":["autonomous driving","multi-modal","high-frequency gated convolution","dual-path attention mechanism","spatial-channel downsampling module","Chemical technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-7d728842130e9ae30fc7","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-d389f9713b081e9ce6a1","title":"MCL-YOLO: A Multi-Module Collaborative Lightweight Object Detection Method for Bridge Crack Detection","year":2026,"authors":["Bingyu Han","Yang Wu","Wenhao Feng"],"author_count":4,"journal":"Sensors","doi":"10.3390/s26185801","source":"doaj","topics":["computer-vision"],"keywords":["bridge crack detection","MCL-YOLO","lightweight object detection","receptive-field attention","dynamic feature enhancement","Chemical technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-d389f9713b081e9ce6a1","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-343385189778e42c8415","title":"Reliability-Aware Dual-Stream Self-Training for Semi-Supervised Semantic Segmentation of High-Resolution Remote Sensing Imagery","year":2026,"authors":["Wangchenxiao Liu","Yan Lin","Wei Liu"],"author_count":4,"journal":"Remote Sensing","doi":"10.3390/rs18183257","source":"doaj","topics":["computer-vision"],"keywords":["high-resolution remote sensing imagery","semi-supervised semantic segmentation","reliability-aware self-training","dual-stream learning","class-adaptive pseudo-labeling","prototype contrastive learning","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-343385189778e42c8415","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-e6ab3fbf392c6f14bc4d","title":"Development of a Semi-Automated Tool Based on Computer Vision Methods for Safety Assessment of Micromobility Users","year":2026,"authors":["Alejandra Sofía Fonseca-Cabrera","David Llopis-Castelló","Alfredo García"],"author_count":3,"journal":"Sensors","doi":"10.3390/s26185713","source":"doaj","topics":["computer-vision"],"keywords":["micromobility","computer vision","video processing","lateral position","speed estimation","bike lane","Chemical technology"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-e6ab3fbf392c6f14bc4d","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-60223478adaa35452cda","title":"Label-Efficient Semi-Supervised Building Semantic Segmentation for High-Resolution UAV Imagery","year":2026,"authors":["Gahyun Lee","Junseo Baek","Youkyung Han"],"author_count":3,"journal":"Remote Sensing","doi":"10.3390/rs18183250","source":"doaj","topics":["computer-vision"],"keywords":["building segmentation","semi-supervised learning","foundation model","boundary-aware learning","UAV imagery","remote sensing","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-60223478adaa35452cda","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-bc97b01656642b70d044","title":"ProG-Net: Prompted Guidance Network for Visible–Thermal Tiny-Object Detection","year":2026,"authors":["Yan Zhang","Qiang Wang","Hui Li"],"author_count":3,"journal":"Remote Sensing","doi":"10.3390/rs18183239","source":"doaj","topics":["computer-vision"],"keywords":["visible–thermal object detection","tiny-object detection","vision-foundation model","prompted guidance","cross-modality fusion","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-bc97b01656642b70d044","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-c5009e2790c83234a018","title":"SDR-YOLO: Scale-Selective Detail Residual Enhanced YOLO for Visible–Thermal Object Detection","year":2026,"authors":["Lijuan Wang","Zuchao Bao","Baichuan Rong"],"author_count":4,"journal":"Remote Sensing","doi":"10.3390/rs18183216","source":"doaj","topics":["computer-vision"],"keywords":["visible–thermal object detection","small-object detection","cross-modal feature fusion","lightweight object detection","UAV remote sensing","YOLO","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-c5009e2790c83234a018","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-29d2845958642fcb1ee7","title":"TOTMSeg: A Texture-Aware Octree-Based Transformer-Mamba Framework for Large-Scale Urban Mesh Semantic Segmentation","year":2026,"authors":["Ruiming Zhang","Mengyu Ma","Chun Du"],"author_count":6,"journal":"Remote Sensing","doi":"10.3390/rs18183198","source":"doaj","topics":["computer-vision"],"keywords":["3D urban meshes","state space model","semantic segmentation","texture convolution","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-29d2845958642fcb1ee7","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-0e7f473b650943aa1c9e","title":"Boosting Multi-Class SAR Oriented Object Detection via Geo-Topology-Guided Diffusion Synthesis","year":2026,"authors":["Zhen Wang","Gang Wan","Cheng Wang"],"author_count":5,"journal":"Remote Sensing","doi":"10.3390/rs18183128","source":"doaj","topics":["computer-vision"],"keywords":["Synthetic Aperture Radar","oriented object detection","diffusion model","data augmentation","geo-topology","oriented bounding box","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-0e7f473b650943aa1c9e","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-831a0884159fcf8e09e7","title":"Boundary-Guided Dual-Perspective Cross-Modal Fusion Network for RGB-IR Object Detection","year":2026,"authors":["Huachen Lin","Zhiwei Fu","Xiumei Chen"],"author_count":4,"journal":"Remote Sensing","doi":"10.3390/rs18183175","source":"doaj","topics":["computer-vision"],"keywords":["RGB-IR object detection","boundary-guided fusion","dual-perspective perception","cross-modal feature fusion","local reliability rectification","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-831a0884159fcf8e09e7","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-d44ec4a107614a96fc9f","title":"ZOS-Net: A Lightweight RGB-T Object Detection Network with Cross-Modal Relation Enhancement and Selective Target Awareness","year":2026,"authors":["Yudan Dai","Rui Guo","Zixiang Yue"],"author_count":3,"journal":"Remote Sensing","doi":"10.3390/rs18183112","source":"doaj","topics":["computer-vision"],"keywords":["RGB-T object detection","multimodal fusion","lightweight object detection","cross-modal relation enhancement","selective target awareness","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-d44ec4a107614a96fc9f","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-9223862972514fd2ca77","title":"A Hybrid Super-Resolution and Object Detection Framework for Small Ship Recognition in Optical Remote Sensing Imagery","year":2026,"authors":["Muhammad Abubakar Saleem","Waseemullah Nazir","Muhammad Umar Farooq"],"author_count":7,"journal":"Remote Sensing","doi":"10.3390/rs18183092","source":"doaj","topics":["computer-vision"],"keywords":["super-resolution","ship detection","optical remote sensing","vision transformers","generative adversarial networks","YOLO","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-9223862972514fd2ca77","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-d7060f9bf328161d4d01","title":"DGCR-Net: Dynamic Graph Contextual Reasoning Network for Semantic Segmentation of Remote Sensing Imagery","year":2026,"authors":["Hong Wang","Kun Gao","Xiaodian Zhang"],"author_count":7,"journal":"Remote Sensing","doi":"10.3390/rs18183069","source":"doaj","topics":["computer-vision"],"keywords":["remote sensing image semantic segmentation","dynamic graph","contextual reasoning","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-d7060f9bf328161d4d01","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-2dc441c21a49662a030f","title":"OPMS-Seg: A UAV-Based High-Resolution Image Dataset for Semantic Segmentation of Open-Pit Coal Mine Slopes","year":2026,"authors":["Lingkai Shi","Yide Geng","Haoran Wang"],"author_count":8,"journal":"Remote Sensing","doi":"10.3390/rs18183094","source":"doaj","topics":["computer-vision"],"keywords":["open-pit coal mine","slope monitoring","semantic segmentation","UAV remote sensing","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-2dc441c21a49662a030f","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-3a249e2751e88e7d99d2","title":"Confusion-Resistant Learning for Few-Shot Oriented Object Detection in Aerial Images","year":2026,"authors":["Yang Yang","Lingjun Li","Yuanhang Wang"],"author_count":6,"journal":"Remote Sensing","doi":"10.3390/rs18183062","source":"doaj","topics":["computer-vision"],"keywords":["few-shot oriented object detection","aerial remote sensing images","confusion-resistant learning","classification reweighting","edge-vector cosine similarity loss","rotation angle ambiguity","Science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-3a249e2751e88e7d99d2","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-f33e669a03523885a0ea","title":"Augmented Reality, Mixed Reality, Computer Vision, and Three-Dimensional Modeling in Fertility-Preserving Minimally Invasive Gynecologic Surgery: A Scoping Review with Narrative Synthesis","year":2026,"authors":["Eleni Karatrasoglou","Alexandros Rodolakis","Themistoklis Grigoriadis"],"author_count":4,"journal":"Medicina","doi":"10.3390/medicina62091790","source":"doaj","topics":["computer-vision"],"keywords":["augmented reality","mixed reality","computer vision","artificial intelligence","three-dimensional modeling","minimally invasive gynecologic surgery","Medicine (General)"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-f33e669a03523885a0ea","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-c907647fe0b3991f6be4","title":"RIF-YOLO-N: Lightweight P3 Residual Identity Fusion with Resolution-Guided Fine-Tuning for Tiny-Object Detection","year":2026,"authors":["Mansoor Iqbal","Balaj Khalid","Syed Zarak Shah"],"author_count":6,"journal":"Journal of Imaging","doi":"10.3390/jimaging12090450","source":"doaj","topics":["computer-vision"],"keywords":["tiny-object detection","UAV imagery","RIF-YOLO-N","P3 feature enhancement","resolution-guided fine-tuning","lightweight object detection","Photography","Computer applications to medicine. Medical informatics"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-c907647fe0b3991f6be4","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-648255ce93ca5febcbab","title":"An Object Detection Method Based on Frequency-Band Enhancement and Multi-Scale Fusion","year":2026,"authors":["Zhenzhao Dai","Yongsheng Qiu","Yuanyao Lu"],"author_count":3,"journal":"Journal of Imaging","doi":"10.3390/jimaging12090447","source":"doaj","topics":["computer-vision"],"keywords":["object detection","RT-DETR","wavelet transform","feature fusion","autonomous driving","Photography","Computer applications to medicine. Medical informatics","Electronic computers. Computer science"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-648255ce93ca5febcbab","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-98527d913cf1d79b7bd2","title":"Hyperspectral Image Dataset for Benchmarking on Salient Object Detection","year":2026,"authors":["Nevrez Imamoglu","Yu Oishi","Xiaoqiang Zhang"],"author_count":7,"journal":null,"doi":"10.1109/qomex.2018.8463428","source":"arxiv","topics":["computer-vision"],"keywords":["cs.CV"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-98527d913cf1d79b7bd2","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-123c3a49664ad3337daf","title":"AROID: Improving Adversarial Robustness Through Online Instance-Wise Data Augmentation","year":2024,"authors":["Lin Li","Jianing Qiu","Michael Spratling"],"author_count":3,"journal":"International Journal of Computer Vision (IJCV), 133:929-50, 2024","doi":"10.1007/s11263-024-02206-4","source":"arxiv","topics":["computer-vision"],"keywords":["cs.CV","cs.AI","cs.LG"],"license":"CC-BY-4.0","has_fulltext":true,"has_references":false,"rights_class":"licensed_content","access":"paid","version":2,"content_endpoint":"/api/document/doi-123c3a49664ad3337daf","fulltext_endpoint":"/api/document/doi-123c3a49664ad3337daf/fulltext","references_endpoint":null},{"id":"doi-eb489cac34e0acb86128","title":"A Comprehensive Assessment Benchmark for Rigorously Evaluating Deep Learning Image Classifiers","year":2025,"authors":["Michael W. Spratling"],"author_count":1,"journal":"Neural Networks, 192(107801), 2025","doi":"10.1016/j.neunet.2025.107801","source":"arxiv","topics":["computer-vision"],"keywords":["cs.LG","cs.CV"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-eb489cac34e0acb86128","fulltext_endpoint":null,"references_endpoint":null},{"id":"arxiv-2312.10986","title":"Long-Tailed 3D Detection via Multi-Modal Fusion","year":2026,"authors":["Yechi Ma","Neehar Peri","Achal Dave"],"author_count":6,"journal":null,"doi":null,"source":"arxiv","topics":["computer-vision"],"keywords":["cs.CV","cs.RO"],"license":"CC-BY-4.0","has_fulltext":true,"has_references":false,"rights_class":"licensed_content","access":"paid","version":2,"content_endpoint":"/api/document/arxiv-2312.10986","fulltext_endpoint":"/api/document/arxiv-2312.10986/fulltext","references_endpoint":null},{"id":"arxiv-2401.13937","title":"DiDA: Video Object Segmentation with Distillation Learning of Deformable Attention","year":2026,"authors":["Quang-Trung Truong","Duc Thanh Nguyen","Binh-Son Hua"],"author_count":4,"journal":null,"doi":null,"source":"arxiv","topics":["computer-vision"],"keywords":["cs.CV"],"license":"CC-BY-4.0","has_fulltext":true,"has_references":false,"rights_class":"licensed_content","access":"paid","version":2,"content_endpoint":"/api/document/arxiv-2401.13937","fulltext_endpoint":"/api/document/arxiv-2401.13937/fulltext","references_endpoint":null},{"id":"arxiv-2407.01295","title":"Towards Formal Verification of Deep Neural Networks for Object Detection","year":2026,"authors":["Avraham Raviv","Yizhak Y. Elboher","Omri Isac"],"author_count":8,"journal":null,"doi":null,"source":"arxiv","topics":["computer-vision"],"keywords":["cs.CV"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/arxiv-2407.01295","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-efefde34f55fe41de63f","title":"Biased Heritage: How Datasets Shape Models in Facial Expression Recognition","year":2025,"authors":["Iris Dominguez-Catena","Daniel Paternain","Mikel Galar"],"author_count":6,"journal":null,"doi":"10.1109/taffc.2026.3693403","source":"arxiv","topics":["computer-vision"],"keywords":["cs.CV","cs.CY"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-efefde34f55fe41de63f","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-5f882fee31dd2323e60c","title":"Latent Diffusion Autoencoders: Toward Efficient and Meaningful Unsupervised Representation Learning in Medical Imaging","year":2025,"authors":["Gabriele Lozupone","Alessandro Bria","Francesco Fontanella"],"author_count":6,"journal":"Medical Image Analysis, 109 (2026), 103932","doi":"10.1016/j.media.2026.103932","source":"arxiv","topics":["computer-vision"],"keywords":["cs.CV"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/doi-5f882fee31dd2323e60c","fulltext_endpoint":null,"references_endpoint":null},{"id":"doi-f018390e19f08ccbbed2","title":"Decoupling Multi-Contrast Super-Resolution: Self-Supervised Implicit Re-Representation for Unpaired Cross-Modal Synthesis","year":2026,"authors":["Yinzhe Wu","Hongyu Rui","Fanwen Wang"],"author_count":8,"journal":null,"doi":"10.1109/tmi.2026.3739662","source":"arxiv","topics":["computer-vision"],"keywords":["cs.CV"],"license":"CC-BY-4.0","has_fulltext":true,"has_references":false,"rights_class":"licensed_content","access":"paid","version":2,"content_endpoint":"/api/document/doi-f018390e19f08ccbbed2","fulltext_endpoint":"/api/document/doi-f018390e19f08ccbbed2/fulltext","references_endpoint":null},{"id":"arxiv-2507.18625","title":"3D Software Synthesis Driven by Constraint-Expressive Intermediate Representation","year":2026,"authors":["Shuqing Li","Anson Y. Lam","Yun Peng"],"author_count":5,"journal":null,"doi":null,"source":"arxiv","topics":["computer-vision","software-engineering"],"keywords":["cs.CV","cs.AI","cs.MM","cs.SE"],"license":"CC0-1.0","has_fulltext":false,"has_references":false,"rights_class":"metadata_only","access":"paid","version":1,"content_endpoint":"/api/document/arxiv-2507.18625","fulltext_endpoint":null,"references_endpoint":null}]}