@misc{indiciae8bd525976856, title = {MaD-Mix: Multi-Modal Data Mixtures via Latent Space Coupling for Vision-Language Model Training}, author = {Wanyun Xie and Francesco Tonin and Volkan Cevher}, year = {2026}, url = {https://arxiv.org/abs/2602.07790}, note = {Source identifier: 2602.07790} }