@misc{indiciae45615123ccb6, title = {Kernel-based Unsupervised Embedding Alignment for Enhanced Visual Representation in Vision-language Models}, author = {Shizhan Gong and Yankai Jiang and Qi Dou and Farzan Farnia}, year = {2025}, url = {https://arxiv.org/abs/2506.02557}, note = {Source identifier: 2506.02557} }