2015-05-08 12:23:27 +05:30
|
|
|
/*
|
|
|
|
* Copyright (c) 2015 The WebM project authors. All Rights Reserved.
|
|
|
|
*
|
|
|
|
* Use of this source code is governed by a BSD-style license
|
|
|
|
* that can be found in the LICENSE file in the root of the source
|
|
|
|
* tree. An additional intellectual property rights grant can be found
|
|
|
|
* in the file PATENTS. All contributing project authors may
|
|
|
|
* be found in the AUTHORS file in the root of the source tree.
|
|
|
|
*/
|
|
|
|
|
2015-06-02 12:16:28 +05:30
|
|
|
#include <assert.h>
|
2015-07-31 10:53:25 -07:00
|
|
|
|
|
|
|
#include "vp9/common/vp9_enums.h"
|
2015-07-31 11:15:55 -07:00
|
|
|
#include "vpx_dsp/mips/inv_txfm_msa.h"
|
2015-05-08 12:23:27 +05:30
|
|
|
|
2015-06-01 09:19:01 +05:30
|
|
|
void vp9_iht8x8_64_add_msa(const int16_t *input, uint8_t *dst,
|
|
|
|
int32_t dst_stride, int32_t tx_type) {
|
2015-05-08 12:23:27 +05:30
|
|
|
v8i16 in0, in1, in2, in3, in4, in5, in6, in7;
|
|
|
|
|
|
|
|
/* load vector elements of 8x8 block */
|
2015-06-01 09:19:01 +05:30
|
|
|
LD_SH8(input, 8, in0, in1, in2, in3, in4, in5, in6, in7);
|
2015-05-08 12:23:27 +05:30
|
|
|
|
2015-06-01 09:19:01 +05:30
|
|
|
TRANSPOSE8x8_SH_SH(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
2015-05-08 12:23:27 +05:30
|
|
|
|
|
|
|
switch (tx_type) {
|
|
|
|
case DCT_DCT:
|
|
|
|
/* DCT in horizontal */
|
|
|
|
VP9_IDCT8x8_1D(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
|
|
|
/* DCT in vertical */
|
2015-06-01 09:19:01 +05:30
|
|
|
TRANSPOSE8x8_SH_SH(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
2015-05-08 12:23:27 +05:30
|
|
|
VP9_IDCT8x8_1D(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
|
|
|
break;
|
|
|
|
case ADST_DCT:
|
|
|
|
/* DCT in horizontal */
|
|
|
|
VP9_IDCT8x8_1D(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
|
|
|
/* ADST in vertical */
|
2015-06-01 09:19:01 +05:30
|
|
|
TRANSPOSE8x8_SH_SH(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
2015-05-08 12:23:27 +05:30
|
|
|
VP9_ADST8(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
|
|
|
break;
|
|
|
|
case DCT_ADST:
|
|
|
|
/* ADST in horizontal */
|
2015-06-01 09:19:01 +05:30
|
|
|
VP9_ADST8(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
2015-05-08 12:23:27 +05:30
|
|
|
/* DCT in vertical */
|
2015-06-01 09:19:01 +05:30
|
|
|
TRANSPOSE8x8_SH_SH(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
2015-05-08 12:23:27 +05:30
|
|
|
VP9_IDCT8x8_1D(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
|
|
|
break;
|
|
|
|
case ADST_ADST:
|
|
|
|
/* ADST in horizontal */
|
|
|
|
VP9_ADST8(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
|
|
|
/* ADST in vertical */
|
2015-06-01 09:19:01 +05:30
|
|
|
TRANSPOSE8x8_SH_SH(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
2015-05-08 12:23:27 +05:30
|
|
|
VP9_ADST8(in0, in1, in2, in3, in4, in5, in6, in7,
|
|
|
|
in0, in1, in2, in3, in4, in5, in6, in7);
|
|
|
|
break;
|
|
|
|
default:
|
|
|
|
assert(0);
|
|
|
|
break;
|
|
|
|
}
|
|
|
|
|
|
|
|
/* final rounding (add 2^4, divide by 2^5) and shift */
|
2015-06-01 09:19:01 +05:30
|
|
|
SRARI_H4_SH(in0, in1, in2, in3, 5);
|
|
|
|
SRARI_H4_SH(in4, in5, in6, in7, 5);
|
2015-05-08 12:23:27 +05:30
|
|
|
|
|
|
|
/* add block and store 8x8 */
|
2015-06-01 09:19:01 +05:30
|
|
|
VP9_ADDBLK_ST8x4_UB(dst, dst_stride, in0, in1, in2, in3);
|
|
|
|
dst += (4 * dst_stride);
|
|
|
|
VP9_ADDBLK_ST8x4_UB(dst, dst_stride, in4, in5, in6, in7);
|
2015-05-08 12:23:27 +05:30
|
|
|
}
|