@misc{indiciaeaabc7863f488, title = {Face and Voice Cross-modal Association with Learning Convex Feature Embedding}, author = {Taewan Kim and Jiwoo Kang}, year = {2026}, doi = {10.1007/s00530-025-01872-9}, url = {https://arxiv.org/abs/2607.28129}, note = {Source identifier: 2607.28129} }