@inproceedings{2291769d067045cea8e875fd45d5e837,
title = "HPC Benchmarking: Scaling Right and Looking Beyond the Average",
abstract = "Designing a balanced HPC system requires an understanding of the dominant performance bottlenecks. There is as yet no well established methodology for a unified evaluation of HPC systems and workloads that quantifies the main performance bottlenecks. In this paper, we execute seven production HPC applications on a production HPC platform, and analyse the key performance bottlenecks: FLOPS performance and memory bandwidth congestion, and the implications on scaling out. We show that the results depend significantly on the number of execution processes and granularity of measurements. We therefore advocate for guidance in the application suites, on selecting the representative scale of the experiments. Also, we propose that the FLOPS performance and memory bandwidth should be represented in terms of the proportions of time with low, moderate and severe utilization. We show that this gives much more precise and actionable evidence than the average.",
keywords = "Bottlenecks, FLOPS, HPC applications, Memory bandwidth, Scaling-out",
author = "Milan Radulovic and Kazi Asifuzzaman and Paul Carpenter and Petar Radojkovi{\'c} and Eduard Ayguad{\'e}",
note = "Publisher Copyright: {\textcopyright} 2018, Springer International Publishing AG, part of Springer Nature.; 24th International European Conference on Parallel and Distributed Computing, Euro-Par 2018 ; Conference date: 27-08-2018 Through 31-08-2018",
year = "2018",
doi = "10.1007/978-3-319-96983-1_10",
language = "English",
isbn = "9783319969824",
series = "Lecture Notes in Computer Science (including subseries Lecture Notes in Artificial Intelligence and Lecture Notes in Bioinformatics)",
publisher = "Springer Verlag",
pages = "135--146",
editor = "Massimo Torquati and Marco Aldinucci and Luca Padovani",
booktitle = "Euro-Par 2018",
}