@misc{indiciae2d9df9843aee, title = {AVTENet: A Human-Cognition-Inspired Audio-Visual Transformer-Based Ensemble Network for Video Deepfake Detection}, author = {Ammarah Hashmi and Sahibzada Adil Shahzad and Chia-Wen Lin and Yu Tsao and Hsin-Min Wang}, year = {2025}, doi = {10.1109/tcds.2025.3554477}, url = {https://arxiv.org/abs/2310.13103}, note = {Source identifier: 2310.13103} }