Commit 1e424367 authored by Steinar Midtskogen's avatar Steinar Midtskogen

Add unit tests for v256 intrinsics

Change-Id: I59f78e6fa53b794026edbba709e1c02af0f76a5f
parent 9edb72c9
/*
* Copyright (c) 2016, Alliance for Open Media. All rights reserved
*
* This source code is subject to the terms of the BSD 2 Clause License and
* the Alliance for Open Media Patent License 1.0. If the BSD 2 Clause License
* was not distributed with this source code in the LICENSE file, you can
* obtain it at www.aomedia.org/license/software. If the Alliance for Open
* Media Patent License 1.0 was not distributed with this source code in the
* PATENTS file, you can obtain it at www.aomedia.org/license/patent.
*/
#define ARCH AVX2
#define ARCH_POSTFIX(name) name##_avx2
#define SIMD_NAMESPACE simd_test_avx2
#include "./simd_impl.h"
/*
* Copyright (c) 2016, Alliance for Open Media. All rights reserved
*
* This source code is subject to the terms of the BSD 2 Clause License and
* the Alliance for Open Media Patent License 1.0. If the BSD 2 Clause License
* was not distributed with this source code in the LICENSE file, you can
* obtain it at www.aomedia.org/license/software. If the Alliance for Open
* Media Patent License 1.0 was not distributed with this source code in the
* PATENTS file, you can obtain it at www.aomedia.org/license/patent.
*/
#define ARCH AVX2
#define ARCH_POSTFIX(name) name##_avx2
#define SIMD_NAMESPACE simd_test_avx2
#include "./simd_cmp_impl.h"
This diff is collapsed.
......@@ -14,7 +14,7 @@
#include "test/clear_system_state.h"
#include "test/register_state_check.h"
#include "aom_dsp/aom_simd_inline.h"
#include "aom_dsp/simd/v128_intrinsics_c.h"
#include "aom_dsp/simd/v256_intrinsics_c.h"
namespace SIMD_NAMESPACE {
......@@ -67,6 +67,19 @@ TYPEDEF_SIMD(V128_V128V128);
TYPEDEF_SIMD(S64_V128V128);
TYPEDEF_SIMD(V128_V128U32);
TYPEDEF_SIMD(U32_V128V128);
TYPEDEF_SIMD(V256_V128);
TYPEDEF_SIMD(V256_V256);
TYPEDEF_SIMD(U64_V256);
TYPEDEF_SIMD(V256_V128V128);
TYPEDEF_SIMD(V256_V256V256);
TYPEDEF_SIMD(S64_V256V256);
TYPEDEF_SIMD(V256_V256U32);
TYPEDEF_SIMD(U32_V256V256);
TYPEDEF_SIMD(V256_U8);
TYPEDEF_SIMD(V256_U16);
TYPEDEF_SIMD(V256_U32);
TYPEDEF_SIMD(U32_V256);
TYPEDEF_SIMD(V64_V256);
// Google Test allows up to 50 tests per case, so split the largest
typedef ARCH_POSTFIX(V64_V64) ARCH_POSTFIX(V64_V64_Part2);
......@@ -74,6 +87,9 @@ typedef ARCH_POSTFIX(V64_V64V64) ARCH_POSTFIX(V64_V64V64_Part2);
typedef ARCH_POSTFIX(V128_V128) ARCH_POSTFIX(V128_V128_Part2);
typedef ARCH_POSTFIX(V128_V128) ARCH_POSTFIX(V128_V128_Part3);
typedef ARCH_POSTFIX(V128_V128V128) ARCH_POSTFIX(V128_V128V128_Part2);
typedef ARCH_POSTFIX(V256_V256) ARCH_POSTFIX(V256_V256_Part2);
typedef ARCH_POSTFIX(V256_V256) ARCH_POSTFIX(V256_V256_Part3);
typedef ARCH_POSTFIX(V256_V256V256) ARCH_POSTFIX(V256_V256V256_Part2);
// These functions are machine tuned located elsewhere
template <typename c_ret, typename c_arg>
......@@ -219,6 +235,70 @@ MY_TEST_P(ARCH_POSTFIX(V128_V128_Part3), TestIntrinsics) {
TestSimd1Arg<c_v128, c_v128>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(U64_V256), TestIntrinsics) {
TestSimd1Arg<uint64_t, c_v256>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V256_V256), TestIntrinsics) {
TestSimd1Arg<c_v256, c_v256>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V256_V128), TestIntrinsics) {
TestSimd1Arg<c_v256, c_v128>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V256_V256V256), TestIntrinsics) {
TestSimd2Args<c_v256, c_v256, c_v256>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V256_V128V128), TestIntrinsics) {
TestSimd2Args<c_v256, c_v128, c_v128>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(U32_V256V256), TestIntrinsics) {
TestSimd2Args<uint32_t, c_v256, c_v256>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(S64_V256V256), TestIntrinsics) {
TestSimd2Args<int64_t, c_v256, c_v256>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V256_V256V256_Part2), TestIntrinsics) {
TestSimd2Args<c_v256, c_v256, c_v256>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V256_V256U32), TestIntrinsics) {
TestSimd2Args<c_v256, c_v256, uint32_t>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V256_V256_Part2), TestIntrinsics) {
TestSimd1Arg<c_v256, c_v256>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V256_V256_Part3), TestIntrinsics) {
TestSimd1Arg<c_v256, c_v256>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V256_U8), TestIntrinsics) {
TestSimd1Arg<c_v256, uint8_t>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V256_U16), TestIntrinsics) {
TestSimd1Arg<c_v256, uint16_t>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V256_U32), TestIntrinsics) {
TestSimd1Arg<c_v256, uint32_t>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(U32_V256), TestIntrinsics) {
TestSimd1Arg<uint32_t, c_v256>(kIterations, mask, maskwidth, name);
}
MY_TEST_P(ARCH_POSTFIX(V64_V256), TestIntrinsics) {
TestSimd1Arg<c_v64, c_v256>(kIterations, mask, maskwidth, name);
}
// Add a macro layer since INSTANTIATE_TEST_CASE_P will quote the name
// so we need to expand it first with the prefix
#define INSTANTIATE(name, type, ...) \
......@@ -591,4 +671,252 @@ INSTANTIATE(ARCH, ARCH_POSTFIX(V128_U32), SIMD_TUPLE(v128_dup_32, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(S64_V128V128),
SIMD_TUPLE(v128_dotp_s16, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(U32_V256V256), SIMD_TUPLE(v256_sad_u8, 0U, 0U),
SIMD_TUPLE(v256_ssd_u8, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(U64_V256), SIMD_TUPLE(v256_hadd_u8, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(S64_V256V256),
SIMD_TUPLE(v256_dotp_s16, 0U, 0U));
INSTANTIATE(
ARCH, ARCH_POSTFIX(V256_V256V256), SIMD_TUPLE(v256_add_8, 0U, 0U),
SIMD_TUPLE(v256_add_16, 0U, 0U), SIMD_TUPLE(v256_sadd_s16, 0U, 0U),
SIMD_TUPLE(v256_add_32, 0U, 0U), SIMD_TUPLE(v256_sub_8, 0U, 0U),
SIMD_TUPLE(v256_ssub_u8, 0U, 0U), SIMD_TUPLE(v256_ssub_s8, 0U, 0U),
SIMD_TUPLE(v256_sub_16, 0U, 0U), SIMD_TUPLE(v256_ssub_s16, 0U, 0U),
SIMD_TUPLE(v256_ssub_u16, 0U, 0U), SIMD_TUPLE(v256_sub_32, 0U, 0U),
SIMD_TUPLE(v256_ziplo_8, 0U, 0U), SIMD_TUPLE(v256_ziphi_8, 0U, 0U),
SIMD_TUPLE(v256_ziplo_16, 0U, 0U), SIMD_TUPLE(v256_ziphi_16, 0U, 0U),
SIMD_TUPLE(v256_ziplo_32, 0U, 0U), SIMD_TUPLE(v256_ziphi_32, 0U, 0U),
SIMD_TUPLE(v256_ziplo_64, 0U, 0U), SIMD_TUPLE(v256_ziphi_64, 0U, 0U),
SIMD_TUPLE(v256_ziplo_128, 0U, 0U), SIMD_TUPLE(v256_ziphi_128, 0U, 0U),
SIMD_TUPLE(v256_unziphi_8, 0U, 0U), SIMD_TUPLE(v256_unziplo_8, 0U, 0U),
SIMD_TUPLE(v256_unziphi_16, 0U, 0U), SIMD_TUPLE(v256_unziplo_16, 0U, 0U),
SIMD_TUPLE(v256_unziphi_32, 0U, 0U), SIMD_TUPLE(v256_unziplo_32, 0U, 0U),
SIMD_TUPLE(v256_pack_s32_s16, 0U, 0U), SIMD_TUPLE(v256_pack_s16_u8, 0U, 0U),
SIMD_TUPLE(v256_pack_s16_s8, 0U, 0U), SIMD_TUPLE(v256_or, 0U, 0U),
SIMD_TUPLE(v256_xor, 0U, 0U), SIMD_TUPLE(v256_and, 0U, 0U),
SIMD_TUPLE(v256_andn, 0U, 0U), SIMD_TUPLE(v256_mullo_s16, 0U, 0U),
SIMD_TUPLE(v256_mulhi_s16, 0U, 0U), SIMD_TUPLE(v256_mullo_s32, 0U, 0U),
SIMD_TUPLE(v256_madd_s16, 0U, 0U), SIMD_TUPLE(v256_madd_us8, 0U, 0U),
SIMD_TUPLE(v256_avg_u8, 0U, 0U), SIMD_TUPLE(v256_rdavg_u8, 0U, 0U),
SIMD_TUPLE(v256_avg_u16, 0U, 0U), SIMD_TUPLE(v256_min_u8, 0U, 0U),
SIMD_TUPLE(v256_max_u8, 0U, 0U), SIMD_TUPLE(v256_min_s8, 0U, 0U),
SIMD_TUPLE(v256_max_s8, 0U, 0U), SIMD_TUPLE(v256_min_s16, 0U, 0U),
SIMD_TUPLE(v256_max_s16, 0U, 0U), SIMD_TUPLE(v256_cmpgt_s8, 0U, 0U),
SIMD_TUPLE(v256_cmplt_s8, 0U, 0U));
INSTANTIATE(
ARCH, ARCH_POSTFIX(V256_V256V256_Part2), SIMD_TUPLE(v256_cmpeq_8, 0U, 0U),
SIMD_TUPLE(v256_cmpgt_s16, 0U, 0U), SIMD_TUPLE(v256_cmplt_s16, 0U, 0U),
SIMD_TUPLE(v256_cmpeq_16, 0U, 0U), SIMD_TUPLE(v256_shuffle_8, 15U, 8U),
SIMD_TUPLE(v256_pshuffle_8, 15U, 8U), SIMD_TUPLE(imm_v256_align<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<2>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<3>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<5>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<6>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<7>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<8>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<9>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<10>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<11>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<12>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<13>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<14>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<15>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<16>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<17>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<18>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<19>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<20>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<21>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<22>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<23>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<24>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<25>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<26>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<27>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<28>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<29>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<30>, 0U, 0U),
SIMD_TUPLE(imm_v256_align<31>, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(V256_V128V128),
SIMD_TUPLE(v256_from_v128, 0U, 0U), SIMD_TUPLE(v256_zip_8, 0U, 0U),
SIMD_TUPLE(v256_zip_16, 0U, 0U), SIMD_TUPLE(v256_zip_32, 0U, 0U),
SIMD_TUPLE(v256_mul_s16, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(V256_V128),
SIMD_TUPLE(v256_unpack_u8_s16, 0U, 0U),
SIMD_TUPLE(v256_unpack_s8_s16, 0U, 0U),
SIMD_TUPLE(v256_unpack_u16_s32, 0U, 0U),
SIMD_TUPLE(v256_unpack_s16_s32, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(V256_V256U32), SIMD_TUPLE(v256_shl_8, 7U, 32U),
SIMD_TUPLE(v256_shr_u8, 7U, 32U), SIMD_TUPLE(v256_shr_s8, 7U, 32U),
SIMD_TUPLE(v256_shl_16, 15U, 32U),
SIMD_TUPLE(v256_shr_u16, 15U, 32U),
SIMD_TUPLE(v256_shr_s16, 15U, 32U),
SIMD_TUPLE(v256_shl_32, 31U, 32U),
SIMD_TUPLE(v256_shr_u32, 31U, 32U),
SIMD_TUPLE(v256_shr_s32, 31U, 32U));
INSTANTIATE(ARCH, ARCH_POSTFIX(V256_V256), SIMD_TUPLE(v256_abs_s8, 0U, 0U),
SIMD_TUPLE(v256_abs_s16, 0U, 0U), SIMD_TUPLE(v256_padd_s16, 0U, 0U),
SIMD_TUPLE(v256_unpacklo_u8_s16, 0U, 0U),
SIMD_TUPLE(v256_unpacklo_s8_s16, 0U, 0U),
SIMD_TUPLE(v256_unpacklo_u16_s32, 0U, 0U),
SIMD_TUPLE(v256_unpacklo_s16_s32, 0U, 0U),
SIMD_TUPLE(v256_unpackhi_u8_s16, 0U, 0U),
SIMD_TUPLE(v256_unpackhi_s8_s16, 0U, 0U),
SIMD_TUPLE(v256_unpackhi_u16_s32, 0U, 0U),
SIMD_TUPLE(v256_unpackhi_s16_s32, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<2>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<3>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<5>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<6>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<7>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<8>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<9>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<10>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<11>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<12>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<13>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<14>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<15>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<16>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<17>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<18>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<19>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<20>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<21>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<22>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<23>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<24>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<25>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<26>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<27>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<28>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<29>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<30>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_byte<31>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<2>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<3>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<5>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<6>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<7>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<8>, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(V256_V256_Part2),
SIMD_TUPLE(imm_v256_shl_n_byte<9>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<10>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<11>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<12>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<13>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<14>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<15>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<16>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<17>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<18>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<19>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<20>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<21>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<22>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<23>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<24>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<25>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<26>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<27>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<28>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<29>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<30>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_byte<31>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_8<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_8<2>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_8<3>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_8<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_8<5>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_8<6>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_8<7>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u8<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u8<2>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u8<3>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u8<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u8<5>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u8<6>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u8<7>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s8<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s8<2>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s8<3>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s8<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s8<5>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s8<6>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s8<7>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_16<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_16<2>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_16<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_16<6>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_16<8>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_16<10>, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(V256_V256_Part3),
SIMD_TUPLE(imm_v256_shl_n_16<12>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_16<14>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u16<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u16<2>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u16<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u16<6>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u16<8>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u16<10>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u16<12>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u16<14>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s16<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s16<2>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s16<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s16<6>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s16<8>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s16<10>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s16<12>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s16<14>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_32<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_32<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_32<8>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_32<12>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_32<16>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_32<20>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_32<24>, 0U, 0U),
SIMD_TUPLE(imm_v256_shl_n_32<28>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u32<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u32<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u32<8>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u32<12>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u32<16>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u32<20>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u32<24>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_u32<28>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s32<1>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s32<4>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s32<8>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s32<12>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s32<16>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s32<20>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s32<24>, 0U, 0U),
SIMD_TUPLE(imm_v256_shr_n_s32<28>, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(V256_U8), SIMD_TUPLE(v256_dup_8, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(V256_U16), SIMD_TUPLE(v256_dup_16, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(V256_U32), SIMD_TUPLE(v256_dup_32, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(U32_V256), SIMD_TUPLE(v256_low_u32, 0U, 0U));
INSTANTIATE(ARCH, ARCH_POSTFIX(V64_V256), SIMD_TUPLE(v256_low_v64, 0U, 0U));
} // namespace SIMD_NAMESPACE
......@@ -141,11 +141,13 @@ LIBAOM_TEST_SRCS-yes += simd_cmp_impl.h
LIBAOM_TEST_SRCS-$(HAVE_SSE2) += simd_cmp_sse2.cc
LIBAOM_TEST_SRCS-$(HAVE_SSSE3) += simd_cmp_ssse3.cc
LIBAOM_TEST_SRCS-$(HAVE_SSE4_1) += simd_cmp_sse4.cc
LIBAOM_TEST_SRCS-$(HAVE_AVX2) += simd_cmp_avx2.cc
LIBAOM_TEST_SRCS-$(HAVE_NEON) += simd_cmp_neon.cc
LIBAOM_TEST_SRCS-yes += simd_impl.h
LIBAOM_TEST_SRCS-$(HAVE_SSE2) += simd_sse2_test.cc
LIBAOM_TEST_SRCS-$(HAVE_SSSE3) += simd_ssse3_test.cc
LIBAOM_TEST_SRCS-$(HAVE_SSE4_1) += simd_sse4_test.cc
LIBAOM_TEST_SRCS-$(HAVE_AVX2) += simd_avx2_test.cc
LIBAOM_TEST_SRCS-$(HAVE_NEON) += simd_neon_test.cc
LIBAOM_TEST_SRCS-yes += intrapred_test.cc
LIBAOM_TEST_SRCS-$(CONFIG_INTRABC) += intrabc_test.cc
......
Markdown is supported
0% or .
You are about to add 0 people to the discussion. Proceed with caution.
Finish editing this message first!
Please register or to comment