// Copyright (c) 2019 PaddlePaddle Authors. All Rights Reserved. // // Licensed under the Apache License, Version 2.0 (the "License"); // you may not use this file except in compliance with the License. // You may obtain a copy of the License at // // http://www.apache.org/licenses/LICENSE-2.0 // // Unless required by applicable law or agreed to in writing, software // distributed under the License is distributed on an "AS IS" BASIS, // WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied. // See the License for the specific language governing permissions and // limitations under the License. #pragma once #include #include #include #include "lite/backends/opencl/cl_half.h" #include "lite/backends/opencl/cl_utility.h" #include "lite/core/kernel.h" #include "lite/kernels/opencl/image_helper.h" #include "lite/operators/op_params.h" #include "lite/utils/cp_logging.h" namespace paddle { namespace lite { namespace kernels { namespace opencl { class ElementwiseAddImageCompute : public KernelLite { public: using param_t = operators::ElementwiseParam; void PrepareForRun() override; void ReInitWhenNeeded() override; void GetGlobalWorkSize(); void Run() override; #ifdef LITE_WITH_PROFILE void SetProfileRuntimeKernelInfo(paddle::lite::profile::OpCharacter* ch) { ch->kernel_func_name = kernel_func_name_; ch->cl_event = event_; // `event_` defined in `kernel.h`, valid after kernel::Run } #endif std::string doc() const override { return "ElementwiseAdd using cl::Image2D, kFP16"; } protected: param_t* ele_param_{nullptr}; DDim last_x_dims_; DDim x_img_shape_ = DDim(std::vector( {static_cast(1), static_cast(1)})); DDim y_img_shape_ = DDim(std::vector( {static_cast(1), static_cast(1)})); DDim out_img_shape_ = DDim(std::vector( {static_cast(1), static_cast(1)})); std::string kernel_func_name_{"elementwise_add"}; std::string build_options_{"-DCL_DTYPE_half"}; std::string time_stamp_{GetTimeStamp()}; bool first_epoch_for_reinit_{true}; cl::Kernel kernel_; cl::NDRange global_work_size_ = cl::NDRange{ static_cast(1), static_cast(1), static_cast(1)}; }; } // namespace opencl } // namespace kernels } // namespace lite } // namespace paddle