@inproceedings{3cd180c659b645d7847c58f1ffc0c948,
title = "heFFTe: Highly efficient fft for exascale",
abstract = "Exascale computing aspires to meet the increasing demands from large scientific applications. Software targeting exascale is typically designed for heterogeneous architectures; henceforth, it is not only important to develop well-designed software, but also make it aware of the hardware architecture and efficiently exploit its power. Currently, several and diverse applications, such as those part of the Exascale Computing Project (ECP) in the United States, rely on efficient computation of the Fast Fourier Transform (FFT). In this context, we present the design and implementation of heFFTe (Highly Efficient FFT for Exascale) library, which targets the upcoming exascale supercomputers. We provide highly (linearly) scalable GPU kernels that achieve more than 40× speedup with respect to local kernels from CPU state-of-the-art libraries, and over 2× speedup for the whole FFT computation. A communication model for parallel FFTs is also provided to analyze the bottleneck for large-scale problems. We show experiments obtained on Summit supercomputer at Oak Ridge National Laboratory, using up to 24,576 IBM Power9 cores and 6,144 NVIDIA V-100 GPUs.",
keywords = "Exascale, FFT, GPUs, Scalable algorithm",
author = "Alan Ayala and Stanimire Tomov and Azzam Haidar and Jack Dongarra",
note = "Publisher Copyright: {\textcopyright} Springer Nature Switzerland AG 2020.; 20th International Conference on Computational Science, ICCS 2020 ; Conference date: 03-06-2020 Through 05-06-2020",
year = "2020",
doi = "10.1007/978-3-030-50371-0_19",
language = "English",
isbn = "9783030503703",
series = "Lecture Notes in Computer Science (including subseries Lecture Notes in Artificial Intelligence and Lecture Notes in Bioinformatics)",
publisher = "Springer Science and Business Media Deutschland GmbH",
pages = "262--275",
editor = "Krzhizhanovskaya, {Valeria V.} and G{\'a}bor Z{\'a}vodszky and Lees, {Michael H.} and Sloot, {Peter M.A.} and Sloot, {Peter M.A.} and Sloot, {Peter M.A.} and Dongarra, {Jack J.} and S{\'e}rgio Brissos and Jo{\~a}o Teixeira",
booktitle = "Computational Science – ICCS 2020 - 20th International Conference, Proceedings",
}