@misc{indiciae135a0291bdb0, title = {LLM-Inference-Bench: Inference Benchmarking of Large Language Models on AI Accelerators}, author = {Krishna Teja Chitty-Venkata and Siddhisanket Raskar and Bharat Kale and Farah Ferdaus and Aditya Tanikanti and Ken Raffenetti and Valerie Taylor and Murali Emani and Venkatram Vishwanath}, year = {2024}, url = {https://arxiv.org/abs/2411.00136}, note = {Source identifier: 2411.00136} }