@misc{indiciaebc35875a82c4, title = {XM-ALIGN: Unified Cross-Modal Embedding Alignment for Face-Voice Association}, author = {Zhihua Fang and Shumei Tao and Junxu Wang and Liang He}, year = {2025}, doi = {10.1109/icassp55912.2026.11463876}, url = {https://arxiv.org/abs/2512.06757}, note = {Source identifier: 2512.06757} }