diff options
author | Cedric Nugteren <web@cedricnugteren.nl> | 2016-03-02 21:18:01 +0100 |
---|---|---|
committer | Cedric Nugteren <web@cedricnugteren.nl> | 2016-03-02 21:18:01 +0100 |
commit | 60da54da5d8cb8dc763c13ba48ec6d8e557a609d (patch) | |
tree | 5c71017dd8280ddfaf7955d621bfd446d8578c1b /include/internal/routines | |
parent | fa79720557412cad605589301580ccda39edce6c (diff) |
Added preliminary support for xHER2 and xSYR2 routines
Diffstat (limited to 'include/internal/routines')
-rw-r--r-- | include/internal/routines/level2/xher2.h | 60 | ||||
-rw-r--r-- | include/internal/routines/level2/xsyr2.h | 46 |
2 files changed, 106 insertions, 0 deletions
diff --git a/include/internal/routines/level2/xher2.h b/include/internal/routines/level2/xher2.h new file mode 100644 index 00000000..26f69046 --- /dev/null +++ b/include/internal/routines/level2/xher2.h @@ -0,0 +1,60 @@ + +// ================================================================================================= +// This file is part of the CLBlast project. The project is licensed under Apache Version 2.0. This +// project loosely follows the Google C++ styleguide and uses a tab-size of two spaces and a max- +// width of 100 characters per line. +// +// Author(s): +// Cedric Nugteren <www.cedricnugteren.nl> +// +// This file implements the Xher2 routine. The precision is implemented using a template argument. +// +// ================================================================================================= + +#ifndef CLBLAST_ROUTINES_XHER2_H_ +#define CLBLAST_ROUTINES_XHER2_H_ + +#include "internal/routine.h" + +namespace clblast { +// ================================================================================================= + +// See comment at top of file for a description of the class +template <typename T> +class Xher2: public Routine<T> { + public: + + // Members and methods from the base class + using Routine<T>::db_; + using Routine<T>::source_string_; + using Routine<T>::queue_; + using Routine<T>::GetProgramFromCache; + using Routine<T>::TestVectorX; + using Routine<T>::TestVectorY; + using Routine<T>::TestMatrixA; + using Routine<T>::TestMatrixAP; + using Routine<T>::RunKernel; + using Routine<T>::ErrorIn; + + // Constructor + Xher2(Queue &queue, Event &event, const std::string &name = "HER2"); + + // Templated-precision implementation of the routine + StatusCode DoHer2(const Layout layout, const Triangle triangle, + const size_t n, + const T alpha, + const Buffer<T> &x_buffer, const size_t x_offset, const size_t x_inc, + const Buffer<T> &y_buffer, const size_t y_offset, const size_t y_inc, + const Buffer<T> &a_buffer, const size_t a_offset, const size_t a_ld, + const bool packed = false); + + private: + // Static variable to get the precision + const static Precision precision_; +}; + +// ================================================================================================= +} // namespace clblast + +// CLBLAST_ROUTINES_XHER2_H_ +#endif diff --git a/include/internal/routines/level2/xsyr2.h b/include/internal/routines/level2/xsyr2.h new file mode 100644 index 00000000..f4dc9375 --- /dev/null +++ b/include/internal/routines/level2/xsyr2.h @@ -0,0 +1,46 @@ + +// ================================================================================================= +// This file is part of the CLBlast project. The project is licensed under Apache Version 2.0. This +// project loosely follows the Google C++ styleguide and uses a tab-size of two spaces and a max- +// width of 100 characters per line. +// +// Author(s): +// Cedric Nugteren <www.cedricnugteren.nl> +// +// This file implements the Xsyr2 routine. The precision is implemented using a template argument. +// +// ================================================================================================= + +#ifndef CLBLAST_ROUTINES_XSYR2_H_ +#define CLBLAST_ROUTINES_XSYR2_H_ + +#include "internal/routines/level2/xher2.h" + +namespace clblast { +// ================================================================================================= + +// See comment at top of file for a description of the class +template <typename T> +class Xsyr2: public Xher2<T> { + public: + + // Uses the regular Xher2 routine + using Xher2<T>::DoHer2; + + // Constructor + Xsyr2(Queue &queue, Event &event, const std::string &name = "SYR2"); + + // Templated-precision implementation of the routine + StatusCode DoSyr2(const Layout layout, const Triangle triangle, + const size_t n, + const T alpha, + const Buffer<T> &x_buffer, const size_t x_offset, const size_t x_inc, + const Buffer<T> &y_buffer, const size_t y_offset, const size_t y_inc, + const Buffer<T> &a_buffer, const size_t a_offset, const size_t a_ld); +}; + +// ================================================================================================= +} // namespace clblast + +// CLBLAST_ROUTINES_XSYR2_H_ +#endif |