@misc{indiciae31f5f4de22a7, title = {Towards End-to-End Explainable Facial Action Unit Recognition via Vision-Language Joint Learning}, author = {Xuri Ge and Junchen Fu and Fuhai Chen and Shan An and Nicu Sebe and Joemon M. Jose}, year = {2024}, doi = {10.1145/3664647.3681443}, url = {https://arxiv.org/abs/2408.00644}, note = {Source identifier: 2408.00644} }