@article{bb117700,
AUTHOR = "Lv, T. and Ji, C.M. and Jiang, H. and Liu, Y.",
TITLE = "HF2TNet: A Hierarchical Fusion Two-Stage Training Network for
Infrared and Visible Image Fusion",
JOURNAL = SPLetters,
VOLUME = "31",
YEAR = "2024",
PAGES = "3164-3168",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114285"}
@article{bb117701,
AUTHOR = "Meng, X.C. and Chen, C.Q. and Liu, Q. and Shao, F.",
TITLE = "Multi-domain pseudo-reference quality evaluation for infrared and
visible image fusion",
JOURNAL = IET-IPR,
VOLUME = "18",
YEAR = "2024",
NUMBER = "13",
PAGES = "4095-4113",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114286"}
@article{bb117702,
AUTHOR = "Bai, Y. and Gao, M. and Li, S.Y. and Wang, P. and Guan, N. and Yin, H.Z. and Yan, Y.H.",
TITLE = "IBFusion: An Infrared and Visible Image Fusion Method Based on
Infrared Target Mask and Bimodal Feature Extraction Strategy",
JOURNAL = MultMed,
VOLUME = "26",
YEAR = "2024",
PAGES = "10610-10622",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114287"}
@article{bb117703,
AUTHOR = "Wang, X.X. and Fang, L.X. and Zhao, J.L. and Pan, Z.K. and Li, H. and Li, Y.",
TITLE = "UUD-Fusion: An unsupervised universal image fusion approach via
generative diffusion model",
JOURNAL = CVIU,
VOLUME = "249",
YEAR = "2024",
PAGES = "104218",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114288"}
@article{bb117704,
AUTHOR = "Wu, X. and Cao, Z.H. and Huang, T.Z. and Deng, L.J. and Chanussot, J. and Vivone, G.",
TITLE = "Fully-Connected Transformer for Multi-Source Image Fusion",
JOURNAL = PAMI,
VOLUME = "47",
YEAR = "2025",
NUMBER = "3",
MONTH = "March",
PAGES = "2071-2088",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114289"}
@article{bb117705,
AUTHOR = "Hussain, I. and Tan, S.Q. and Huang, J.W.",
TITLE = "Few-Shot Based Learning Recaptured Image Detection with Multi-Scale
Feature Fusion and Attention",
JOURNAL = PR,
VOLUME = "161",
YEAR = "2025",
PAGES = "111248",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114290"}
@article{bb117706,
AUTHOR = "Tang, H. and Liu, D.W. and Shen, C.C.",
TITLE = "Data-efficient multi-scale fusion vision transformer",
JOURNAL = PR,
VOLUME = "161",
YEAR = "2025",
PAGES = "111305",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114291"}
@article{bb117707,
AUTHOR = "Liu, T.F. and Zhang, M.Y. and Gong, M. and Zhang, Q.F. and Jiang, F.L. and Zheng, H.H. and Lu, D.",
TITLE = "Commonality Feature Representation Learning for Unsupervised
Multimodal Change Detection",
JOURNAL = IP,
VOLUME = "34",
YEAR = "2025",
PAGES = "1219-1233",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114292"}
@article{bb117708,
AUTHOR = "Xu, J.J. and Liu, T.F. and Lei, T. and Chen, H.R.X. and Yokoya, N. and Lv, Z.Y. and Gong, M.",
TITLE = "CGSL: Commonality graph structure learning for unsupervised
multimodal change detection",
JOURNAL = PandRS,
VOLUME = "229",
YEAR = "2025",
PAGES = "92-106",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114293"}
@article{bb117709,
AUTHOR = "Dong, C. and Wang, L.Z. and Zhang, F. and Hua, Q.",
TITLE = "Multi-modal Few-shot Image Recognition with enhanced semantic and
visual integration",
JOURNAL = IVC,
VOLUME = "157",
YEAR = "2025",
PAGES = "105490",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114294"}
@article{bb117710,
AUTHOR = "Tang, L. and Liu, Y. and Tian, Y.J. and Pardalos, P.M.",
TITLE = "Complementary label learning with multi-view data and a
semi-supervised labeling mechanism",
JOURNAL = PR,
VOLUME = "165",
YEAR = "2025",
PAGES = "111651",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114295"}
@article{bb117711,
AUTHOR = "Zhou, M. and Huang, J. and Yan, K.Y. and Hong, D.F. and Jia, X.P. and Chanussot, J. and Li, C.Y.",
TITLE = "A General Spatial-Frequency Learning Framework for Multimodal Image
Fusion",
JOURNAL = PAMI,
VOLUME = "47",
YEAR = "2025",
NUMBER = "7",
MONTH = "July",
PAGES = "5281-5298",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114296"}
@article{bb117712,
AUTHOR = "Wang, Z. and Zhao, L. and Zhang, J.Z. and Song, R. and Song, H.Y. and Meng, J. and Wang, S.D.",
TITLE = "Multi-Text Guidance Is Important: Multi-Modality Image Fusion via Large
Generative Vision-Language Model",
JOURNAL = IJCV,
VOLUME = "133",
YEAR = "2025",
NUMBER = "7",
MONTH = "July",
PAGES = "4646-4668",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114297"}
@article{bb117713,
AUTHOR = "Liu, Y. and Li, C.X. and Xu, S.K. and Han, J.G.",
TITLE = "Part-Whole Relational Fusion Towards Multi-Modal Scene Understanding",
JOURNAL = IJCV,
VOLUME = "133",
YEAR = "2025",
NUMBER = "7",
MONTH = "July",
PAGES = "4483-4503",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114298"}
@article{bb117714,
AUTHOR = "Ravi, J. and Narmadha, R.",
TITLE = "A Systematic Literature Review on Multimodal Image Fusion Models with
Challenges and Future Research Trends",
JOURNAL = IJIG,
VOLUME = "25",
YEAR = "2025",
NUMBER = "4",
MONTH = "July",
PAGES = "2550039",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114299"}
@article{bb117715,
AUTHOR = "Liu, Y.P. and Sun, Z.C. and Yu, B.S. and Zhao, Y.T. and Du, B. and Xu, Y.C. and Cheng, J.",
TITLE = "MIFNet: Learning Modality-Invariant Features for Generalizable
Multimodal Image Matching",
JOURNAL = IP,
VOLUME = "34",
YEAR = "2025",
PAGES = "3593-3608",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114300"}
@article{bb117716,
AUTHOR = "Lu, M. and Jiang, M. and Tao, X.F. and Kong, J.",
TITLE = "AU-Net: Adaptive Unified Network for Joint Multi-Modal Image
Registration and Fusion",
JOURNAL = IP,
VOLUME = "34",
YEAR = "2025",
PAGES = "4721-4735",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114301"}
@article{bb117717,
AUTHOR = "Wang, Q.H. and Li, Z.W. and Zhang, S.Q. and Chi, N. and Dai, Q.H.",
TITLE = "WaveFusion: A Novel Wavelet Vision Transformer With Saliency-Guided
Enhancement for Multimodal Image Fusion",
JOURNAL = CirSysVideo,
VOLUME = "35",
YEAR = "2025",
NUMBER = "8",
MONTH = "August",
PAGES = "7526-7542",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114302"}
@article{bb117718,
AUTHOR = "Liang, P.W. and Jiang, J.J. and Ma, Q. and Wang, C.Y. and Liu, X.M. and Ma, J.Y.",
TITLE = "FusionINV: A Diffusion-Based Approach for Multimodal Image Fusion",
JOURNAL = IP,
VOLUME = "34",
YEAR = "2025",
PAGES = "5355-5368",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114303"}
@article{bb117719,
AUTHOR = "Shi, L.T. and Zhong, B. and Liang, Q.H. and Hu, X.T. and Mo, Z.Y. and Song, S.X.",
TITLE = "Mamba Adapter: Efficient Multi-Modal Fusion for Vision-Language
Tracking",
JOURNAL = CirSysVideo,
VOLUME = "35",
YEAR = "2025",
NUMBER = "9",
MONTH = "September",
PAGES = "9300-9311",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114304"}
@article{bb117720,
AUTHOR = "Liu, X.Y. and Ming, R. and Du, S.L. and He, L.H. and Luo, H.B. and Xiao, G.B.",
TITLE = "HSENet: Hierarchical Semantic-Enriched Network for Multi-Modal Image
Fusion",
JOURNAL = PR,
VOLUME = "170",
YEAR = "2026",
PAGES = "112043",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114305"}
@article{bb117721,
AUTHOR = "Zavras, A. and Michail, D. and Demir, B. and Papoutsis, I.",
TITLE = "Mind the modality gap: Towards a remote sensing vision-language model
via cross-modal alignment",
JOURNAL = PandRS,
VOLUME = "228",
YEAR = "2025",
PAGES = "270-287",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114306"}
@article{bb117722,
AUTHOR = "Cheng, T. and Chen, H. and Zhang, X.H. and Gao, X.W. and Yin, L. and Jiao, J.B.",
TITLE = "Multi-Channel Spatio-Temporal Data Fusion of 'Big' and 'Small'
Network Data Using Transformer Networks",
JOURNAL = IJGI,
VOLUME = "14",
YEAR = "2025",
NUMBER = "8",
PAGES = "286",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114307"}
@article{bb117723,
AUTHOR = "Hu, J.J. and Fan, C. and Ozay, M. and Gao, Q. and Guo, Y.L. and Lam, T.L.",
TITLE = "Robust Depth Estimation Under Sensor Degradations:
A Multi-Sensor Fusion Perspective",
JOURNAL = PAMI,
VOLUME = "47",
YEAR = "2025",
NUMBER = "10",
MONTH = "October",
PAGES = "8691-8707",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114308"}
@article{bb117724,
AUTHOR = "Xin, J.W. and Shi, B. and Wang, N.N. and Li, J. and Gao, X.B.",
TITLE = "MVFusion: Generative Representation Learning With Masked Variational
Autoencoders for Multi-Modality Image Fusion",
JOURNAL = IP,
VOLUME = "34",
YEAR = "2025",
PAGES = "6418-6431",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114309"}
@article{bb117725,
AUTHOR = "Zheng, T.H. and Dong, G.L. and Zhang, P.P. and He, X. and Ren, C.",
TITLE = "Plug-and-Play General Image Registration for Misaligned Multi-Modal
Image Fusion",
JOURNAL = CirSysVideo,
VOLUME = "35",
YEAR = "2025",
NUMBER = "10",
MONTH = "October",
PAGES = "10017-10031",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114310"}
@article{bb117726,
AUTHOR = "Jiao, S.C. and Long, L. and Kuang, L.Q. and Xiong, F.G. and Han, X.",
TITLE = "Multi-modal semantic embedding network for 3D shape recognition and
retrieval",
JOURNAL = JVCIR,
VOLUME = "112",
YEAR = "2025",
PAGES = "104559",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114311"}
@article{bb117727,
AUTHOR = "Sun, H. and Lv, L. and Zhang, P.P. and Tang, T.D. and Tian, F. and Sun, W.B. and Lu, H.C.",
TITLE = "Spatial-Frequency Enhanced Mamba for Multi-Modal Image Fusion",
JOURNAL = IP,
VOLUME = "34",
YEAR = "2025",
PAGES = "7684-7696",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114312"}
@article{bb117728,
AUTHOR = "Zhu, Y.X. and Lv, L. and Zhang, P.P. and Liu, X.H. and Tang, T.D. and Tian, F. and Sun, W.B. and Lu, H.C.",
TITLE = "Interactive Spatial-Frequency Fusion Mamba for Multi-Modal Image
Fusion",
JOURNAL = IP,
VOLUME = "35",
YEAR = "2026",
PAGES = "2380-2392",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114313"}
@article{bb117729,
AUTHOR = "Sun, Y.J. and Dong, W.S. and Wang, S. and Wu, P. and Feng, M.T. and Li, X. and Shi, G.M.",
TITLE = "Distilling Hierarchical Knowledge From Multimodal Fusion for Unimodal
Image Segmentation",
JOURNAL = CirSysVideo,
VOLUME = "35",
YEAR = "2025",
NUMBER = "12",
MONTH = "December",
PAGES = "11797-11809",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114314"}
@article{bb117730,
AUTHOR = "Yu, C.B. and Pei, Z.H. and Wang, X.R. and Zhou, H.B.",
TITLE = "CrossGlue: Cross-Modal Image matching via potential message
investigation and visual-gradient message integration",
JOURNAL = JVCIR,
VOLUME = "114",
YEAR = "2026",
PAGES = "104620",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114315"}
@article{bb117731,
AUTHOR = "Zhou, D.D. and Xu, L. and Wu, K. and Liu, H.Z. and Jiang, M.T.",
TITLE = "DSEPGAN: A Dual-Stream Enhanced Pyramid Based on Generative
Adversarial Network for Spatiotemporal Image Fusion",
JOURNAL = RS,
VOLUME = "17",
YEAR = "2025",
NUMBER = "24",
PAGES = "4050",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114316"}
@article{bb117732,
AUTHOR = "Jiang, J.L. and Hu, G. and Sheng, G.L. and Wei, G.",
TITLE = "PSG-MCANet: Multi-order cross-attention modeling for multimodal
fusion based on punning semantic guidance",
JOURNAL = PR,
VOLUME = "172",
YEAR = "2026",
PAGES = "112723",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114317"}
@article{bb117733,
AUTHOR = "Li, M.Y. and Meng, C. and Fan, X.D.",
TITLE = "Iterative optimal transport for multimodal image registration",
JOURNAL = PR,
VOLUME = "172",
YEAR = "2026",
PAGES = "112736",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114318"}
@article{bb117734,
AUTHOR = "Wang, Y.X. and Shen, Z.W. and Li, H. and Zhang, Y.N. and Xia, Z.P.",
TITLE = "SGCNet: Silhouette Guided Cascaded Network for Multi-Modal Image
Fusion",
JOURNAL = CVIU,
VOLUME = "263",
YEAR = "2026",
PAGES = "104603",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114319"}
@article{bb117735,
AUTHOR = "He, D. and Wang, G.F. and Li, W.S. and Shu, Y.C. and Li, W.B. and Yang, L.J. and Huang, Y.P. and Li, F.Y.",
TITLE = "Rethinking normalization strategies and convolutional kernels for
multimodal image fusion",
JOURNAL = PR,
VOLUME = "173",
YEAR = "2026",
PAGES = "112903",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114320"}
@article{bb117736,
AUTHOR = "Li, S.T. and Tang, H.",
TITLE = "Multimodal Alignment and Fusion: A Survey",
JOURNAL = IJCV,
VOLUME = "134",
YEAR = "2026",
NUMBER = "1",
MONTH = "January",
PAGES = "103",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114321"}
@article{bb117737,
AUTHOR = "Fu, Y. and Ye, X. and Kong, X.Y.",
TITLE = "KPTFusion: Knowledge Prior-based Task-Driven Multimodal Image Fusion",
JOURNAL = IVC,
VOLUME = "167",
YEAR = "2026",
PAGES = "105886",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114322"}
@article{bb117738,
AUTHOR = "Qin, X.R. and Cui, Y.N. and Sun, S.Q. and Chen, R. and Ren, W.Q. and Knoll, A. and Cao, X.C.",
TITLE = "Disentangle to Fuse: Toward Content Preservation and Cross-Modality
Consistency for Multi-Modality Image Fusion",
JOURNAL = IP,
VOLUME = "35",
YEAR = "2026",
PAGES = "1756-1770",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114323"}
@article{bb117739,
AUTHOR = "Chen, H. and Zhou, H.R. and Zhang, Y. and Lin, Z. and Deng, Y.J.",
TITLE = "Dissecting RGB-D Learning for Improved Multi-Modal Fusion",
JOURNAL = IP,
VOLUME = "35",
YEAR = "2026",
PAGES = "1846-1857",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114324"}
@article{bb117740,
AUTHOR = "Zhang, J.J. and Zhao, F. and Liu, H.Q. and Yu, J.",
TITLE = "Generative Information-Guided Heterogeneous Cross-Fusion Network With
Contrastive Learning for Multimodal Remote Sensing Image
Classification",
JOURNAL = CirSysVideo,
VOLUME = "36",
YEAR = "2026",
NUMBER = "2",
MONTH = "February",
PAGES = "1876-1892",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114325"}
@article{bb117741,
AUTHOR = "Mutakabbir, A. and Lung, C.H. and Zaman, M. and Upadhyay, D. and Naik, K. and Millard, K. and Ravichandran, T. and Purcell, R.",
TITLE = "NOAH: A Multi-Modal and Sensor Fusion Dataset for Generative Modeling
in Remote Sensing",
JOURNAL = RS,
VOLUME = "18",
YEAR = "2026",
NUMBER = "3",
PAGES = "466",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114326"}
@article{bb117742,
AUTHOR = "Rao, J.H. and Liu, R. and Guan, J.J. and Tian, X.",
TITLE = "AMS-Former: Adaptive multi-scale transformer for multi-modal image
matching",
JOURNAL = PandRS,
VOLUME = "232",
YEAR = "2026",
PAGES = "957-973",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114327"}
@article{bb117743,
AUTHOR = "Cao, J.Z. and Chen, J.S. and Wang, X.X. and Huang, W.M. and Chen, D.S. and Zhao, T.H. and Tu, W. and Li, Q.Q.",
TITLE = "UrbanMMCL: Urban region representations via multi-modal and
multi-graph self-supervised contrastive learning",
JOURNAL = PandRS,
VOLUME = "232",
YEAR = "2026",
PAGES = "75-93",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114328"}
@article{bb117744,
AUTHOR = "Ying, Z.H. and Guo, J. and Li, Y.S. and Gao, Y. and Li, C.Y.",
TITLE = "Diff-Transformer: Heterogeneous Feature Fusion Network for
Multisource Remote Sensing Classification",
JOURNAL = CirSysVideo,
VOLUME = "36",
YEAR = "2026",
NUMBER = "2",
MONTH = "February",
PAGES = "1501-1516",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114329"}
@article{bb117745,
AUTHOR = "Li, J.Y. and Jiang, C.J. and Jiang, J.J. and Liang, P.W. and Ma, J.Y. and Nie, L.Q.",
TITLE = "Towards Unified Semantic and Controllable Image Fusion: A Diffusion
Transformer Approach",
JOURNAL = PAMI,
VOLUME = "48",
YEAR = "2026",
NUMBER = "4",
MONTH = "April",
PAGES = "3970-3987",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114330"}
@article{bb117746,
AUTHOR = "Panda, G. and Kundu, S. and Bhattacharya, S. and Routray, A.",
TITLE = "L_0-Regularized Sparse Coding-Based Interpretable Network for
Multi-Modal Image Fusion",
JOURNAL = PAMI,
VOLUME = "48",
YEAR = "2026",
NUMBER = "4",
MONTH = "April",
PAGES = "4081-4097",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114331"}
@article{bb117747,
AUTHOR = "Kamara, A.A. and He, S. and Fofanah, A.J.",
TITLE = "FAMAFuse: Functional-Anatomical Multiscale Attention for Multimodal
Image Fusion",
JOURNAL = CirSysVideo,
VOLUME = "36",
YEAR = "2026",
NUMBER = "3",
MONTH = "March",
PAGES = "3215-3230",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114332"}
@article{bb117748,
AUTHOR = "Fan, J. and Bocus, M.J. and Shu, S.L.",
TITLE = "Embodied multi-modal data fusion via geometry anchoring for
continuous perception in ground robots",
JOURNAL = PRL,
VOLUME = "203",
YEAR = "2026",
PAGES = "162-169",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114333"}
@article{bb117749,
AUTHOR = "Zhang, L. and Yang, Y.G. and He, Z.S. and Li, G.L. and Zhao, F. and Hua, W.Q. and Xiao, G.W. and Zhang, J.Y.",
TITLE = "Multimodal Remote Sensing Image Classification Based on Dynamic Group
Convolution and Bidirectional Guided Cross-Attention Fusion",
JOURNAL = RS,
VOLUME = "18",
YEAR = "2026",
NUMBER = "7",
PAGES = "1066",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114334"}
@article{bb117750,
AUTHOR = "Yu, M. and Lu, X. and Yang, Z. and Gao, D. and Zhong, G.Q.",
TITLE = "DAMFusion: Multi-Spectral Image Segmentation via Competitive Query
and Boundary Region Attention",
JOURNAL = RS,
VOLUME = "18",
YEAR = "2026",
NUMBER = "7",
PAGES = "1064",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114335"}
@article{bb117751,
AUTHOR = "Yang, J. and Chung, H. and Jang, I.",
TITLE = "Hierarchical mutual distillation for multi-view fusion: Learning from
all possible view combinations",
JOURNAL = PR,
VOLUME = "178",
YEAR = "2026",
PAGES = "113432",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114336"}
@article{bb117752,
AUTHOR = "Pan, Y.J. and Shi, Y.C. and Yu, C. and Kong, X.Z. and Zhang, Y. and Xiao, N.",
TITLE = "Beyond a single perspective: A multi-agent debate framework for
affective computing",
JOURNAL = PR,
VOLUME = "178",
YEAR = "2026",
PAGES = "113445",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114337"}
@article{bb117753,
AUTHOR = "Yang, A. and Liu, B.Q. and Liu, M.Z. and Ding, H.H. and Mo, P.J. and Zhao, C.Q. and Liu, X.H. and Ye, T.",
TITLE = "RIF-Fuse: Invertible Frequency Decomposition with Residual
Enhancement for Robust Multimodal Fusion",
JOURNAL = RS,
VOLUME = "18",
YEAR = "2026",
NUMBER = "10",
PAGES = "1520",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114338"}
@article{bb117754,
AUTHOR = "Zhang, Y.C. and Chen, R.S. and Zhang, S. and Leng, B.",
TITLE = "Focus-then-fusion: Learning discriminative cross-modal prototypes for
few-shot classification",
JOURNAL = PR,
VOLUME = "179",
YEAR = "2026",
PAGES = "113527",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114339"}
@article{bb117755,
AUTHOR = "Zhou, Z.C. and Yu, T.Y. and Chen, J.F. and Liang, J.",
TITLE = "DGNNF: Dynamic Graph Neural Network Fusion for 3D object detection",
JOURNAL = PR,
VOLUME = "179",
YEAR = "2026",
PAGES = "113707",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114340"}
@article{bb117756,
AUTHOR = "Li, Y. and Xing, Y.F. and Lan, X.Y. and Li, X. and Chen, H.F. and Jiang, D.M.",
TITLE = "AlignMamba-2: Enhancing multimodal fusion and sentiment analysis with
modality-aware Mamba",
JOURNAL = PR,
VOLUME = "179",
YEAR = "2026",
PAGES = "113517",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114341"}
@article{bb117757,
AUTHOR = "Chen, T. and Wang, C. and Zhang, Y.D. and Xia, K.J. and Qian, P.J.",
TITLE = "DMFusion: Degradation-Customized Mixture-of-Experts With Adaptive
Discrimination for Multi-Modal Image Fusion",
JOURNAL = CirSysVideo,
VOLUME = "36",
YEAR = "2026",
NUMBER = "6",
MONTH = "June",
PAGES = "8506-8521",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114342"}
@article{bb117758,
AUTHOR = "Li, Z.P. and Hu, J. and Guan, W. and Miao, J.W. and Wu, K. and Xia, Z.G. and Yang, B. and Wu, J.S.",
TITLE = "FTransMamba: A multi-stage fusion transformer and mamba modeling for
multimodal remote sensing scene understanding",
JOURNAL = PR,
VOLUME = "179",
YEAR = "2026",
PAGES = "113625",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114343"}
@article{bb117759,
AUTHOR = "Ma, X.Y. and Yang, Y. and Bian, K.",
TITLE = "WiViHAR: A Deep Learning-Based Human Activity Recognition Method
Using WiFi and Vision Multimodal Fusion",
JOURNAL = HMS,
VOLUME = "56",
YEAR = "2026",
NUMBER = "3",
MONTH = "June",
PAGES = "582-591",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114344"}
@article{bb117760,
AUTHOR = "Yu, Y.Y. and Wang, T. and Qiang, Y. and Wang, X.Y. and Chen, X. and Qiu, W.",
TITLE = "Plane-Wave Image Reconstruction With Hy-PCF: A Novel Hybrid
CNN-Transformer for Progressive Cross-Domain Fusion",
JOURNAL = MedImg,
VOLUME = "45",
YEAR = "2026",
NUMBER = "6",
MONTH = "June",
PAGES = "3007-3020",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114345"}
@article{bb117761,
AUTHOR = "Hu, M.Q. and Sun, B. and Li, S.T. and Ma, J.Y.",
TITLE = "SPEN: Sub-Pixel Position Error Estimation Network for Multi-Modal
Image Matching",
JOURNAL = CirSysVideo,
VOLUME = "36",
YEAR = "2026",
NUMBER = "6",
MONTH = "June",
PAGES = "8006-8020",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114346"}
@article{bb117762,
AUTHOR = "Sun, Y.M. and Cui, X. and Wang, Z. and Cheng, H. and Dong, Y.F. and Zhu, P.F. and Li, K.",
TITLE = "TEDFuse: Task-Driven Equivariant Consistency Decomposition Network
for Multi-Modal Image Fusion",
JOURNAL = MultMed,
VOLUME = "28",
YEAR = "2026",
PAGES = "4332-4345",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114347"}
@article{bb117763,
AUTHOR = "Wang, M.Y. and Liu, Z.Y. and Li, K. and Wang, Y. and Wang, Y.W. and Wei, Y.Y. and Wang, F.",
TITLE = "Task-Generalized Adaptive Cross-Domain Learning for Multimodal Image
Fusion",
JOURNAL = MultMed,
VOLUME = "28",
YEAR = "2026",
PAGES = "4624-4637",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114348"}
@article{bb117764,
AUTHOR = "Yu, D. and Tang, Y. and Zhang, C.J. and Wang, W. and Yang, G.D. and Zheng, X.L. and Zhao, Y.",
TITLE = "IA2GNN: Imbalance-Aware Adaptive Graph Construction for Multi-Modal
Image Fusion",
JOURNAL = MultMed,
VOLUME = "28",
YEAR = "2026",
PAGES = "4747-4758",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114349"}
@article{bb117765,
AUTHOR = "Li, Y.L. and Li, L. and Zhao, X.W. and Wang, J.M.",
TITLE = "F&S-Net: A Dual Mission (Fusion and Super-Resolution) Framework Under
Various Input Resolution",
JOURNAL = MultMed,
VOLUME = "28",
YEAR = "2026",
PAGES = "4928-4941",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114350"}
@article{bb117766,
AUTHOR = "Qin, Y. and Feng, Y.L. and Sun, Y. and Peng, D.Z. and Peng, X. and Hu, P.",
TITLE = "Deep Information-Balanced Multimodal Learning",
JOURNAL = PAMI,
VOLUME = "48",
YEAR = "2026",
NUMBER = "8",
MONTH = "August",
PAGES = "9384-9396",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114351"}
@article{bb117767,
AUTHOR = "Huang, J.J. and Yan, P.X. and Liu, J. and Wu, J. and Wang, Z. and Wang, Y.T. and Wu, X.L. and Lin, L. and Li, G.B.",
TITLE = "DreamFuse: Toward Realistic and Seamless Image Fusion Across Diverse
Scenarios",
JOURNAL = PAMI,
VOLUME = "48",
YEAR = "2026",
NUMBER = "8",
MONTH = "August",
PAGES = "9502-9518",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114352"}
@article{bb117768,
AUTHOR = "Cao, K. and He, X. and Hu, T. and Xie, C.J. and Zhou, M. and Zhang, J.",
TITLE = "Shuffle Mamba: State Space Models With Random Shuffle for Multi-Modal
Image Fusion",
JOURNAL = CirSysVideo,
VOLUME = "36",
YEAR = "2026",
NUMBER = "7",
MONTH = "July",
PAGES = "9448-9461",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114353"}
@article{bb117769,
AUTHOR = "Zhu, A. and Hu, M. and Wang, X.H. and Ren, F.",
TITLE = "Beneficial Noise Learning for Robust Multimodal Fusion",
JOURNAL = MultMed,
VOLUME = "28",
YEAR = "2026",
PAGES = "5900-5911",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114354"}
@article{bb117770,
AUTHOR = "Gao, C.Z. and Li, W. and Weng, D. and Tao, R. and Xia, X.G. and Du, Q.",
TITLE = "HIMO: Cross-Arbitrary-Modality Image Invariant Feature Transform with
Hierarchical Intrinsic Major Orientation",
JOURNAL = PAMI,
VOLUME = "48",
YEAR = "2026",
NUMBER = "8",
MONTH = "August",
PAGES = "9001-9018",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114355"}
@article{bb117771,
AUTHOR = "Shi, X. and Zhang, R. and Liu, J.W. and Liu, Y.P. and Liang, Z. and Cheng, Q.K. and Lu, W.",
TITLE = "Modality Equilibrium Matters: Minor-Modality-Aware Adaptive
Alternating for Cross-Modal Memory Enhancement",
JOURNAL = PAMI,
VOLUME = "48",
YEAR = "2026",
NUMBER = "8",
MONTH = "August",
PAGES = "10176-10183",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114356"}
@article{bb117772,
AUTHOR = "Tian, W. and Du, Z.L. and Zhao, X.L. and Yu, Q.",
TITLE = "AMTFusion: Boosting 3D Object Detection by Adaptive Multi-Modal
Temporal Fusion and Augmentation",
JOURNAL = CirSysVideo,
VOLUME = "36",
YEAR = "2026",
NUMBER = "8",
MONTH = "August",
PAGES = "11561-11575",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114357"}
@article{bb117773,
AUTHOR = "Yang, B. and Jiang, Z.H. and Pan, D. and Lin, Z.P. and Gui, W.H.",
TITLE = "MOFM: A Multiple-in-One Flow Mamba for Unregistered Multi-Modal Image
Fusion",
JOURNAL = CirSysVideo,
VOLUME = "36",
YEAR = "2026",
NUMBER = "8",
MONTH = "August",
PAGES = "12296-12310",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114358"}
@inproceedings{bb117774,
AUTHOR = "Xue, F. and Elflein, S. and Leal Taixe, L. and Zhou, Q.",
TITLE = "MATCHA: Towards Matching Anything",
BOOKTITLE = CVPR25,
YEAR = "2025",
PAGES = "27081-27091",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114359"}
@inproceedings{bb117775,
AUTHOR = "Zhou, B. and Li, L. and Wang, Y.J. and Liu, H.F. and Yao, Y.Z. and Wang, W.G.",
TITLE = "UniAlign: Scaling Multimodal Alignment within One Unified Model",
BOOKTITLE = CVPR25,
YEAR = "2025",
PAGES = "29644-29655",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114360"}
@inproceedings{bb117776,
AUTHOR = "Hou, J.M. and Chen, X.Y. and Ran, R. and Cong, X.F. and Liu, X.Y. and You, J.W. and Deng, L.J.",
TITLE = "Binarized Neural Network for Multi-spectral Image Fusion",
BOOKTITLE = CVPR25,
YEAR = "2025",
PAGES = "2236-2245",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114361"}
@inproceedings{bb117777,
AUTHOR = "Li, Y. and Xing, Y.F. and Lan, X.Y. and Li, X. and Chen, H.F. and Jiang, D.M.",
TITLE = "AlignMamba: Enhancing Multimodal Mamba with Local and Global
Cross-Modal Alignment",
BOOKTITLE = CVPR25,
YEAR = "2025",
PAGES = "24774-24784",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114362"}
@inproceedings{bb117778,
AUTHOR = "Maniparambil, M. and Akshulakov, R. and Djilali, Y.A.D. and Narayan, S. and Singh, A. and O'Connor, N.E.",
TITLE = "Harnessing Frozen Unimodal Encoders for Flexible Multimodal Alignment",
BOOKTITLE = CVPR25,
YEAR = "2025",
PAGES = "29847-29857",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114363"}
@inproceedings{bb117779,
AUTHOR = "Li, H. and Hou, Y.N. and Xing, X.H. and Ma, Y.X. and Sun, X. and Zhang, Y.",
TITLE = "OccMamba: Semantic Occupancy Prediction with State Space Models",
BOOKTITLE = CVPR25,
YEAR = "2025",
PAGES = "11949-11959",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114364"}
@inproceedings{bb117780,
AUTHOR = "Wu, G.Y. and Liu, H.Y. and Fu, H.M. and Peng, Y.C. and Liu, J.Y. and Fan, X. and Liu, R.S.",
TITLE = "Every SAM Drop Counts: Embracing Semantic Priors for Multi-Modality
Image Fusion and Beyond",
BOOKTITLE = CVPR25,
YEAR = "2025",
PAGES = "17882-17891",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114365"}
@inproceedings{bb117781,
AUTHOR = "Tran, Q.H. and Ahmed, M. and Popattia, M. and Ahmed, M.H. and Konin, A. and Zia, M.Z.",
TITLE = "Learning by Aligning 2D Skeleton Sequences and Multi-Modality Fusion",
BOOKTITLE = ECCV24,
YEAR = "2024",
PAGES = "L: 141-161",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114366"}
@inproceedings{bb117782,
AUTHOR = "Li, C.X. and Liu, X.Y. and Wang, C. and Liu, Y.F. and Yu, W.H. and Shao, J. and Yuan, Y.X.",
TITLE = "GTP-4O: Modality-prompted Heterogeneous Graph Learning for Omni-modal
Biomedical Representation",
BOOKTITLE = ECCV24,
YEAR = "2024",
PAGES = "IV: 168-187",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114367"}
@inproceedings{bb117783,
AUTHOR = "Song, Z.Q. and Wang, L.F.",
TITLE = "Dual Multi-Modal Feature Fusion Network for the Evaluation of
Osteosarcoma",
BOOKTITLE = ICIP24,
YEAR = "2024",
PAGES = "2937-2943",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114368"}
@inproceedings{bb117784,
AUTHOR = "Gao, Z.X. and Jiang, X. and Xu, X. and Shen, F.M. and Li, Y.J. and Shen, H.T.",
TITLE = "Embracing Unimodal Aleatoric Uncertainty for Robust Multimodal Fusion",
BOOKTITLE = CVPR24,
YEAR = "2024",
PAGES = "26866-26875",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114369"}
@inproceedings{bb117785,
AUTHOR = "Jiang, H. and Karpur, A. and Cao, B. and Huang, Q.X. and Araujo, A.",
TITLE = "OmniGlue: Generalizable Feature Matching with Foundation Model
Guidance",
BOOKTITLE = CVPR24,
YEAR = "2024",
PAGES = "19865-19875",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114370"}
@inproceedings{bb117786,
AUTHOR = "Yi, X.P. and Xu, H. and Zhang, H. and Tang, L.F. and Ma, J.Y.",
TITLE = "Text-IF: Leveraging Semantic Text Guidance for Degradation-Aware and
Interactive Image Fusion",
BOOKTITLE = CVPR24,
YEAR = "2024",
PAGES = "27016-27025",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114371"}
@inproceedings{bb117787,
AUTHOR = "Vouitsis, N. and Liu, Z.Y. and Gorti, S.K. and Villecroze, V. and Cresswell, J.C. and Yu, G.W. and Loaiza Ganem, G. and Volkovs, M.",
TITLE = "Data-Efficient Multimodal Fusion on a Single GPU",
BOOKTITLE = CVPR24,
YEAR = "2024",
PAGES = "27229-27241",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114372"}
@inproceedings{bb117788,
AUTHOR = "Zhao, Z.X. and Bai, H.W. and Zhang, J.S. and Zhang, Y. and Zhang, K. and Xu, S. and Chen, D.D. and Timofte, R. and Van Gool, L.J.",
TITLE = "Equivariant Multi-Modality Image Fusion",
BOOKTITLE = CVPR24,
YEAR = "2024",
PAGES = "25912-25921",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114373"}
@inproceedings{bb117789,
AUTHOR = "Han, K.Y. and Cao, F.Z. and Shi, T.X. and Wang, P.",
TITLE = "A Dual Attention Network for Multimodal Remote Sensing Image Matching",
BOOKTITLE = CVIDL23,
YEAR = "2023",
PAGES = "128-134",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114374"}
@inproceedings{bb117790,
AUTHOR = "Liu, B. and Xu, Z.Q. and Bao, X.L. and Zhong, Z.",
TITLE = "MUNformer: A strong encoder that uses multi-level features extracted
by different feature extractors for fusion",
BOOKTITLE = CVIDL23,
YEAR = "2023",
PAGES = "291-295",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114375"}
@inproceedings{bb117791,
AUTHOR = "He, C.M. and Li, K. and Xu, G.X. and Zhang, Y. and Hu, R.Z. and Guo, Z.H. and Li, X.",
TITLE = "Degradation-Resistant Unfolding Network for Heterogeneous Image
Fusion",
BOOKTITLE = ICCV23,
YEAR = "2023",
PAGES = "12577-12587",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114376"}
@inproceedings{bb117792,
AUTHOR = "Liu, J.Y. and Liu, Z. and Wu, G.Y. and Ma, L. and Liu, R.S. and Zhong, W. and Luo, Z.X. and Fan, X.",
TITLE = "Multi-interactive Feature Learning and a Full-time Multi-modality
Benchmark for Image Fusion and Segmentation",
BOOKTITLE = ICCV23,
YEAR = "2023",
PAGES = "8081-8090",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114377"}
@inproceedings{bb117793,
AUTHOR = "Sippel, F. and Seiler, J. and Kaup, A.",
TITLE = "Cross Spectral Image Reconstruction Using a Deep Guided Neural
Network",
BOOKTITLE = ICIP23,
YEAR = "2023",
PAGES = "226-230",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114378"}
@inproceedings{bb117794,
AUTHOR = "Myers, A. and Kvinge, H. and Emerson, T.",
TITLE = "TopFusion: Using Topological Feature Space for Fusion and Imputation
in Multi-Modal Data",
BOOKTITLE = TAG-PRA23,
YEAR = "2023",
PAGES = "600-609",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114379"}
@inproceedings{bb117795,
AUTHOR = "Xue, Z. and Marculescu, R.",
TITLE = "Dynamic Multimodal Fusion",
BOOKTITLE = MULA23,
YEAR = "2023",
PAGES = "2575-2584",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114380"}
@inproceedings{bb117796,
AUTHOR = "Kong, L.K. and Qi, X.S. and Shen, Q.J. and Wang, J.C. and Zhang, J.Y. and Hu, Y. and Zhou, Q.C.",
TITLE = "Indescribable Multi-Modal Spatial Evaluator",
BOOKTITLE = CVPR23,
YEAR = "2023",
PAGES = "9853-9862",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114381"}
@inproceedings{bb117797,
AUTHOR = "Zhao, Z.X. and Bai, H.W. and Zhang, J.S. and Zhang, Y. and Xu, S. and Lin, Z. and Timofte, R. and Van Gool, L.J.",
TITLE = "CDDFuse: Correlation-Driven Dual-Branch Feature Decomposition for
Multi-Modality Image Fusion",
BOOKTITLE = CVPR23,
YEAR = "2023",
PAGES = "5906-5916",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114382"}
@inproceedings{bb117798,
AUTHOR = "Li, Y.W. and Quan, R.J. and Zhu, L.C. and Yang, Y.",
TITLE = "Efficient Multimodal Fusion via Interactive Prompting",
BOOKTITLE = CVPR23,
YEAR = "2023",
PAGES = "2604-2613",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114383"}
@inproceedings{bb117799,
AUTHOR = "Wetzer, E. and Lindblad, J. and Sladoje, N.",
TITLE = "Can Representation Learning for Multimodal Image Registration be
Improved by Supervision of Intermediate Layers?",
BOOKTITLE = IbPRIA23,
YEAR = "2023",
PAGES = "261-275",
BIBSOURCE = "http://www.visionbib.com/bibliography/match-pl502mmf1.html#TT114384"}
Last update:Aug 19, 2026 at 13:26:35