Blame - src/runtime/CL/functions/CLConvolutionLayer.cpp - platform/external/ComputeLibrary

blob: 0014e717340cf3c5755dba451e169110554fc818 [file] [log] [blame]

Anthony Barbier	871448e	2017-03-24 14:54:29 +0000	[diff] [blame]	1	/*
Anthony Barbier	f45d5a9	2018-01-24 16:23:15 +0000	[diff] [blame]	2	* Copyright (c) 2017-2018 ARM Limited.
Anthony Barbier	871448e	2017-03-24 14:54:29 +0000	[diff] [blame]	3	*
				4	* SPDX-License-Identifier: MIT
				5	*
				6	* Permission is hereby granted, free of charge, to any person obtaining a copy
				7	* of this software and associated documentation files (the "Software"), to
				8	* deal in the Software without restriction, including without limitation the
				9	* rights to use, copy, modify, merge, publish, distribute, sublicense, and/or
				10	* sell copies of the Software, and to permit persons to whom the Software is
				11	* furnished to do so, subject to the following conditions:
				12	*
				13	* The above copyright notice and this permission notice shall be included in all
				14	* copies or substantial portions of the Software.
				15	*
				16	* THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND, EXPRESS OR
				17	* IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF MERCHANTABILITY,
				18	* FITNESS FOR A PARTICULAR PURPOSE AND NONINFRINGEMENT. IN NO EVENT SHALL THE
				19	* AUTHORS OR COPYRIGHT HOLDERS BE LIABLE FOR ANY CLAIM, DAMAGES OR OTHER
				20	* LIABILITY, WHETHER IN AN ACTION OF CONTRACT, TORT OR OTHERWISE, ARISING FROM,
				21	* OUT OF OR IN CONNECTION WITH THE SOFTWARE OR THE USE OR OTHER DEALINGS IN THE
				22	* SOFTWARE.
				23	*/
				24	#include "arm_compute/runtime/CL/functions/CLConvolutionLayer.h"
				25
				26	#include "arm_compute/core/PixelValue.h"
				27	#include "arm_compute/core/Utils.h"
				28	#include "arm_compute/core/Validate.h"
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	29	#include "arm_compute/core/utils/misc/ShapeCalculator.h"
Anthony Barbier	8140e1e	2017-12-14 23:48:46 +0000	[diff] [blame]	30	#include "arm_compute/core/utils/quantization/AsymmHelpers.h"
Anthony Barbier	871448e	2017-03-24 14:54:29 +0000	[diff] [blame]	31	#include "arm_compute/runtime/CL/CLScheduler.h"
				32
				33	#include <cmath>
Kaizen	8938bd3	2017-09-28 14:38:23 +0100	[diff] [blame]	34	#include <memory>
Anthony Barbier	871448e	2017-03-24 14:54:29 +0000	[diff] [blame]	35	#include <tuple>
				36
				37	using namespace arm_compute;
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	38	using namespace arm_compute::misc::shape_calculator;
Anthony Barbier	dbdab85	2017-06-23 15:42:00 +0100	[diff] [blame]	39
Kaizen	8938bd3	2017-09-28 14:38:23 +0100	[diff] [blame]	40	CLConvolutionLayer::CLConvolutionLayer(std::shared_ptr<IMemoryManager> memory_manager)
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	41	: _memory_manager(std::move(memory_manager)), _function()
Anthony Barbier	dbdab85	2017-06-23 15:42:00 +0100	[diff] [blame]	42	{
				43	}
				44
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	45	void CLConvolutionLayer::configure(ICLTensor input, const ICLTensor weights, const ICLTensor biases, ICLTensor output, const PadStrideInfo &conv_info, const WeightsInfo &weights_info,
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	46	const Size2D &dilation, const ActivationLayerInfo &act_info, bool enable_fast_math, unsigned int num_groups)
Anthony Barbier	8140e1e	2017-12-14 23:48:46 +0000	[diff] [blame]	47	{
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	48	ARM_COMPUTE_ERROR_ON_NULLPTR(input, weights, output);
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	49	ARM_COMPUTE_ERROR_THROW_ON(CLConvolutionLayer::validate(input->info(), weights->info(), ((biases != nullptr) ? biases->info() : nullptr), output->info(), conv_info, weights_info, dilation, act_info,
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	50	enable_fast_math, num_groups));
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	51
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	52	switch(CLConvolutionLayer::get_convolution_method(input->info(), weights->info(), output->info(), conv_info,
				53	weights_info, act_info, CLScheduler::get().target(), dilation, enable_fast_math))
Anthony Barbier	8140e1e	2017-12-14 23:48:46 +0000	[diff] [blame]	54	{
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	55	case ConvolutionMethod::WINOGRAD:
				56	{
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	57	ARM_COMPUTE_ERROR_ON(num_groups != 1);
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	58	auto f = arm_compute::support::cpp14::make_unique<CLWinogradConvolutionLayer>(_memory_manager);
				59	f->configure(input, weights, biases, output, conv_info, act_info, enable_fast_math);
				60	_function = std::move(f);
				61	break;
				62	}
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	63	case ConvolutionMethod::DIRECT:
Anthony Barbier	f45d5a9	2018-01-24 16:23:15 +0000	[diff] [blame]	64	{
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	65	ARM_COMPUTE_ERROR_ON(num_groups != 1);
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	66	auto f = arm_compute::support::cpp14::make_unique<CLDirectConvolutionLayer>();
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	67	f->configure(input, weights, biases, output, conv_info, act_info);
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	68	_function = std::move(f);
				69	break;
Anthony Barbier	f45d5a9	2018-01-24 16:23:15 +0000	[diff] [blame]	70	}
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	71	case ConvolutionMethod::GEMM:
Anthony Barbier	f45d5a9	2018-01-24 16:23:15 +0000	[diff] [blame]	72	{
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	73	auto f = arm_compute::support::cpp14::make_unique<CLGEMMConvolutionLayer>(_memory_manager);
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	74	f->configure(input, weights, biases, output, conv_info, weights_info, dilation, act_info, num_groups);
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	75	_function = std::move(f);
				76	break;
Anthony Barbier	f45d5a9	2018-01-24 16:23:15 +0000	[diff] [blame]	77	}
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	78	default:
				79	ARM_COMPUTE_ERROR("Not supported.");
				80	break;
Anthony Barbier	8140e1e	2017-12-14 23:48:46 +0000	[diff] [blame]	81	}
				82	}
				83
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	84	Status CLConvolutionLayer::validate(const ITensorInfo input, const ITensorInfo weights, const ITensorInfo biases, const ITensorInfo output, const PadStrideInfo &conv_info,
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	85	const WeightsInfo &weights_info, const Size2D &dilation, const ActivationLayerInfo &act_info, bool enable_fast_math, unsigned int num_groups)
Anthony Barbier	871448e	2017-03-24 14:54:29 +0000	[diff] [blame]	86	{
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	87	ARM_COMPUTE_RETURN_ERROR_ON_NULLPTR(input, weights, output);
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	88	ARM_COMPUTE_RETURN_ERROR_ON_MSG((num_groups != 1) && (input->data_layout() != DataLayout::NCHW), "Grouping (num_groups != 1) with NHWC data layout is not supported");
Anthony Barbier	8140e1e	2017-12-14 23:48:46 +0000	[diff] [blame]	89
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	90	const GPUTarget gpu_target = CLScheduler::get().target();
Anthony Barbier	871448e	2017-03-24 14:54:29 +0000	[diff] [blame]	91
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	92	switch(CLConvolutionLayer::get_convolution_method(input, weights, output, conv_info, weights_info, act_info, gpu_target, dilation, enable_fast_math))
Anthony Barbier	871448e	2017-03-24 14:54:29 +0000	[diff] [blame]	93	{
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	94	case ConvolutionMethod::WINOGRAD:
				95	{
				96	//Validate Winograd
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	97	ARM_COMPUTE_RETURN_ERROR_ON_MSG(num_groups != 1, "Grouping (num_groups != 1) with CLWinogradConvolutionLayer is not supported");
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	98	ARM_COMPUTE_RETURN_ON_ERROR(CLWinogradConvolutionLayer::validate(input, weights, biases, output, conv_info, act_info, enable_fast_math));
				99	break;
				100	}
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	101	case ConvolutionMethod::DIRECT:
Anthony Barbier	8140e1e	2017-12-14 23:48:46 +0000	[diff] [blame]	102	{
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	103	// Validate direct convolution layer
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	104	ARM_COMPUTE_RETURN_ERROR_ON_MSG(num_groups != 1, "Grouping (num_groups != 1) with CLDirectConvolutionLayer is not supported");
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	105	ARM_COMPUTE_RETURN_ON_ERROR(CLDirectConvolutionLayer::validate(input, weights, biases, output, conv_info, act_info));
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	106	break;
Anthony Barbier	8140e1e	2017-12-14 23:48:46 +0000	[diff] [blame]	107	}
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	108	case ConvolutionMethod::GEMM:
Anthony Barbier	8140e1e	2017-12-14 23:48:46 +0000	[diff] [blame]	109	{
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	110	// Validate gemm-based convolution layer
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	111	ARM_COMPUTE_RETURN_ON_ERROR(CLGEMMConvolutionLayer::validate(input, weights, biases, output, conv_info, weights_info, dilation, act_info, num_groups));
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	112	break;
Anthony Barbier	8140e1e	2017-12-14 23:48:46 +0000	[diff] [blame]	113	}
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	114	default:
				115	ARM_COMPUTE_ERROR("Not supported.");
				116	break;
Anthony Barbier	871448e	2017-03-24 14:54:29 +0000	[diff] [blame]	117	}
				118
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	119	return Status{};
				120	}
Kaizen	8938bd3	2017-09-28 14:38:23 +0100	[diff] [blame]	121
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	122	ConvolutionMethod CLConvolutionLayer::get_convolution_method(const ITensorInfo input, const ITensorInfo weights, const ITensorInfo *output, const PadStrideInfo &conv_info,
				123	const WeightsInfo &weights_info, const ActivationLayerInfo &act_info, const GPUTarget gpu_target, const Size2D &dilation, bool enable_fast_math)
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	124	{
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	125	ARM_COMPUTE_ERROR_ON_NULLPTR(input);
				126	ARM_COMPUTE_ERROR_ON_NULLPTR(output);
				127	ARM_COMPUTE_ERROR_ON_NULLPTR(weights);
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	128	ARM_COMPUTE_UNUSED(weights_info);
				129	ARM_COMPUTE_UNUSED(gpu_target);
Kaizen	8938bd3	2017-09-28 14:38:23 +0100	[diff] [blame]	130
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	131	const size_t idx_w = get_data_layout_dimension_index(input->data_layout(), DataLayoutDimension::WIDTH);
				132	const size_t idx_h = get_data_layout_dimension_index(input->data_layout(), DataLayoutDimension::HEIGHT);
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	133	const size_t idx_c = get_data_layout_dimension_index(input->data_layout(), DataLayoutDimension::CHANNEL);
				134
Jenkins	52ba29e	2018-08-29 15:32:11 +0000	[diff] [blame^]	135	/* Input spatial dims, kernel size, IFM/OFM, conv info*/
				136	using ConvolutionConfiguration = std::tuple<Size2D, Size2D, Size2D, PadStrideInfo, DataLayout>;
				137	using ConfigurationMethod = std::pair<ConvolutionConfiguration, ConvolutionMethod>;
				138
				139	const std::vector<ConfigurationMethod> known_configs =
				140	{
				141	// Alexnet
				142	ConfigurationMethod(ConvolutionConfiguration(Size2D(27U, 27U), Size2D(5U, 5U), Size2D(48U, 128U), PadStrideInfo(1U, 1U, 2U, 2U), DataLayout::NCHW), ConvolutionMethod::DIRECT),
				143	// VGG16 / VGG19
				144	ConfigurationMethod(ConvolutionConfiguration(Size2D(224U, 224U), Size2D(3U, 3U), Size2D(3U, 64U), PadStrideInfo(1U, 1U, 1U, 1U), DataLayout::NCHW), ConvolutionMethod::DIRECT),
				145	// Mobilenet 224
				146	ConfigurationMethod(ConvolutionConfiguration(Size2D(224U, 224U), Size2D(3U, 3U), Size2D(3U, 32U), PadStrideInfo(2U, 2U, 0U, 1U, 0U, 1U, DimensionRoundingType::FLOOR), DataLayout::NCHW), ConvolutionMethod::GEMM),
				147	// Mobilenet 160
				148	ConfigurationMethod(ConvolutionConfiguration(Size2D(160U, 160U), Size2D(3U, 3U), Size2D(3U, 24U), PadStrideInfo(2U, 2U, 0U, 1U, 0U, 1U, DimensionRoundingType::FLOOR), DataLayout::NCHW), ConvolutionMethod::GEMM),
				149	// Mobilenet 224
				150	ConfigurationMethod(ConvolutionConfiguration(Size2D(224U, 224U), Size2D(3U, 3U), Size2D(3U, 32U), PadStrideInfo(2U, 2U, 0U, 1U, 0U, 1U, DimensionRoundingType::FLOOR), DataLayout::NHWC), ConvolutionMethod::GEMM),
				151	// Mobilenet 160
				152	ConfigurationMethod(ConvolutionConfiguration(Size2D(160U, 160U), Size2D(3U, 3U), Size2D(3U, 24U), PadStrideInfo(2U, 2U, 0U, 1U, 0U, 1U, DimensionRoundingType::FLOOR), DataLayout::NHWC), ConvolutionMethod::GEMM),
				153	};
				154
				155	const auto find_config = [&](ConfigurationMethod c)
				156	{
				157	const ConvolutionConfiguration config = c.first;
				158	const PadStrideInfo info = std::get<3>(config);
				159	const DataLayout data_layout = std::get<4>(config);
				160
				161	return std::get<0>(config) == Size2D(input->dimension(idx_w), input->dimension(idx_h)) && std::get<1>(config) == Size2D(weights->dimension(idx_w), weights->dimension(idx_h))
				162	&& std::get<2>(config) == Size2D(weights->dimension(idx_c), weights->dimension(3)) && info.pad_top() == conv_info.pad_top() && info.pad_right() == conv_info.pad_right()
				163	&& info.pad_bottom() == conv_info.pad_bottom() && info.pad_left() == conv_info.pad_left() && info.stride() == conv_info.stride() && (data_layout == input->data_layout());
				164	};
				165
				166	std::vector<ConfigurationMethod>::const_iterator found;
				167	if((found = std::find_if(known_configs.begin(), known_configs.end(), find_config)) != known_configs.end())
				168	{
				169	return (*found).second;
				170	}
				171
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	172	if(dilation != Size2D(1U, 1U) \|\| (input->dimension(idx_c) < 16))
				173	{
				174	return ConvolutionMethod::GEMM;
				175	}
				176	else
				177	{
				178	return bool(CLWinogradConvolutionLayer::validate(input, weights, nullptr, output, conv_info, act_info, enable_fast_math)) ? ConvolutionMethod::WINOGRAD : ConvolutionMethod::GEMM;
				179	}
Anthony Barbier	871448e	2017-03-24 14:54:29 +0000	[diff] [blame]	180	}
				181
				182	void CLConvolutionLayer::run()
				183	{
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	184	prepare();
Anthony Barbier	06ea048	2018-02-22 15:45:35 +0000	[diff] [blame]	185	_function->run();
Anthony Barbier	871448e	2017-03-24 14:54:29 +0000	[diff] [blame]	186	}
Jenkins	b3a371b	2018-05-23 11:36:53 +0100	[diff] [blame]	187
				188	void CLConvolutionLayer::prepare()
				189	{
				190	_function->prepare();
				191	}