@misc{indiciaedbe1af1c8341, title = {Evaluating Text-to-Image and Text-to-Video Synthesis with a Conditional Fréchet Distance}, author = {Jaywon Koo and Jefferson Hernandez and Moayed Haji-Ali and Ziyan Yang and Vicente Ordonez}, year = {2025}, url = {https://arxiv.org/abs/2503.21721}, note = {Source identifier: 2503.21721} }