@misc{indiciae25ebdb580a5d, title = {Detecting Lip-Syncing Deepfakes: Vision Temporal Transformer for Analyzing Mouth Inconsistencies}, author = {Soumyya Kanti Datta and Shan Jia and Siwei Lyu}, year = {2025}, url = {https://arxiv.org/abs/2504.01470}, note = {Source identifier: 2504.01470} }