@misc{indiciae270577d1b907, title = {Better with Less: Tackling Heterogeneous Multi-Modal Image Joint Pretraining via Conditioned and Degraded Masked Autoencoder}, author = {Bowen Peng and Yongxiang Liu and Jie Zhou and Xiaodong Chen and Tianpeng Liu and Xiaogang Yu and Li Liu}, year = {2026}, url = {https://arxiv.org/abs/2604.16952}, note = {Source identifier: 2604.16952} }