@misc{indiciaeaa49a3ceebeb, title = {Video-ChatGPT: Towards Detailed Video Understanding via Large Vision and Language Models}, author = {Muhammad Maaz and Hanoona Rasheed and Salman Khan and Fahad Shahbaz Khan}, year = {2024}, url = {https://arxiv.org/abs/2306.05424}, note = {Source identifier: 2306.05424} }