COMPMID-1446 : Add support for 3D output in NEGEMMLowpOutputStage

Change-Id: I61e7d39d09a9936b1128ec04038fa2d8dfe6a2c8 Reviewed-on: https://eu-gerrit-1.euhpc.arm.com/149211 Reviewed-by: Isabella Gottardi <isabella.gottardi@arm.com> Reviewed-by: Anthony Barbier <anthony.barbier@arm.com> Tested-by: bsgcomp <bsgcomp@arm.com>
author: Georgios Pinitas <georgios.pinitas@arm.com> 2018-09-18 18:38:37 +0100
committer: Anthony Barbier <anthony.barbier@arm.com> 2018-11-02 16:54:54 +0000
commit: 041f36d4dc1b6473d9f7136659a384d611fab0b6 (patch)
tree: 01a337d08f5c8f5382eddc29585608e612cc272d /arm_compute/core/NEON/kernels/NEGEMMLowpQuantizeDownInt32ToUint8ScaleByFixedPointKernel.h
parent: ceb889efc302464efd7fd20001d8a89a06c4e0bd (diff)
download: ComputeLibrary-041f36d4dc1b6473d9f7136659a384d611fab0b6.tar.gz
1 files changed, 13 insertions, 9 deletions
diff --git a/arm_compute/core/NEON/kernels/NEGEMMLowpQuantizeDownInt32ToUint8ScaleByFixedPointKernel.h b/arm_compute/core/NEON/kernels/NEGEMMLowpQuantizeDownInt32ToUint8ScaleByFixedPointKernel.h
index 030a0c766b..6ebb515af7 100644
--- a/arm_compute/core/NEON/kernels/NEGEMMLowpQuantizeDownInt32ToUint8ScaleByFixedPointKernel.h
+++ b/arm_compute/core/NEON/kernels/NEGEMMLowpQuantizeDownInt32ToUint8ScaleByFixedPointKernel.h
@@ -72,21 +72,24 @@ public:
      * @param[in]  min                          (Optional) Min value used to saturate down the output result before converting back to QASYMM8
      * @param[in]  max                          (Optional) Max value used to saturate up the output result before converting back to QASYMM8,
      *                                          Along with @p min, this value can be used to implement "rectified linear unit" activation functions
+     * @param[in]  gemm_3d_depth                (Optional)     Depth of GEMM 3D (Defaults to 1)
      */
-    void configure(const ITensor *input, const ITensor *bias, ITensor *output, int result_fixedpoint_multiplier, int result_shift, int result_offset_after_shift, int min = 0, int max = 0);
+    void configure(const ITensor *input, const ITensor *bias, ITensor *output, int result_fixedpoint_multiplier, int result_shift, int result_offset_after_shift,
+                   int min = 0, int max = 0, unsigned int gemm_3d_depth = 1);
     /** Static function to check if given info will lead to a valid configuration of @ref NEGEMMLowpQuantizeDownInt32ToUint8ScaleByFixedPointKernel
      *
-     * @param[in] input  Input tensor. Data type supported: S32
-     * @param[in] bias   Biases tensor. Only shared biases supported and it can be a nullptr if the biases addition is not required.
-     *                   Biases are 1D tensor with dimensions [OFM]. Data type supported: Same as @p input.
-     * @param[in] output Output tensor. Data type supported: Data type supported: QASYMM8
-     * @param[in] min    (Optional) Min value used to saturate down the output result before converting back to QASYMM8
-     * @param[in] max    (Optional) Max value used to saturate up the output result before converting back to QASYMM8,
-     *                   Along with @p min, this value can be used to implement "rectified linear unit" activation functions
+     * @param[in] input         Input tensor. Data type supported: S32
+     * @param[in] bias          Biases tensor. Only shared biases supported and it can be a nullptr if the biases addition is not required.
+     *                          Biases are 1D tensor with dimensions [OFM]. Data type supported: Same as @p input.
+     * @param[in] output        Output tensor. Data type supported: Data type supported: QASYMM8
+     * @param[in] min           (Optional) Min value used to saturate down the output result before converting back to QASYMM8
+     * @param[in] max           (Optional) Max value used to saturate up the output result before converting back to QASYMM8,
+     *                          Along with @p min, this value can be used to implement "rectified linear unit" activation functions
+     * @param[in] gemm_3d_depth (Optional)  Depth of GEMM 3D (Defaults to 1)
      *
      * @return a status
      */
-    static Status validate(const ITensorInfo *input, const ITensorInfo *bias, const ITensorInfo *output, int min = 0, int max = 0);
+    static Status validate(const ITensorInfo *input, const ITensorInfo *bias, const ITensorInfo *output, int min = 0, int max = 0, unsigned int gemm_3d_depth = 1);
 
     // Inherited methods overridden:
     void run(const Window &window, const ThreadInfo &info) override;
@@ -114,6 +117,7 @@ private:
     int                     _result_offset_after_shift;
     int                     _min;
     int                     _max;
+    unsigned int            _gemm_3d_depth;
 };
 } // namespace arm_compute
author	Georgios Pinitas <georgios.pinitas@arm.com>	2018-09-18 18:38:37 +0100
committer	Anthony Barbier <anthony.barbier@arm.com>	2018-11-02 16:54:54 +0000
commit	041f36d4dc1b6473d9f7136659a384d611fab0b6 (patch)
tree	01a337d08f5c8f5382eddc29585608e612cc272d /arm_compute/core/NEON/kernels/NEGEMMLowpQuantizeDownInt32ToUint8ScaleByFixedPointKernel.h
parent	ceb889efc302464efd7fd20001d8a89a06c4e0bd (diff)
download	ComputeLibrary-041f36d4dc1b6473d9f7136659a384d611fab0b6.tar.gz