You can not select more than 25 topics Topics must start with a letter or number, can include dashes ('-') and can be up to 35 characters long.

391 lines
21KB

  1. /*
  2. * Copyright (c) 2010 Jason Garrett-Glaser
  3. *
  4. * This file is part of Libav.
  5. *
  6. * Libav is free software; you can redistribute it and/or
  7. * modify it under the terms of the GNU Lesser General Public
  8. * License as published by the Free Software Foundation; either
  9. * version 2.1 of the License, or (at your option) any later version.
  10. *
  11. * Libav is distributed in the hope that it will be useful,
  12. * but WITHOUT ANY WARRANTY; without even the implied warranty of
  13. * MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the GNU
  14. * Lesser General Public License for more details.
  15. *
  16. * You should have received a copy of the GNU Lesser General Public
  17. * License along with Libav; if not, write to the Free Software
  18. * Foundation, Inc., 51 Franklin Street, Fifth Floor, Boston, MA 02110-1301 USA
  19. */
  20. #include "libavutil/cpu.h"
  21. #include "libavcodec/h264pred.h"
  22. #define PRED4x4(TYPE, DEPTH, OPT) \
  23. void ff_pred4x4_ ## TYPE ## _ ## DEPTH ## _ ## OPT (uint8_t *src, const uint8_t *topright, int stride);
  24. PRED4x4(dc, 10, mmxext)
  25. PRED4x4(down_left, 10, sse2)
  26. PRED4x4(down_left, 10, avx)
  27. PRED4x4(down_right, 10, sse2)
  28. PRED4x4(down_right, 10, ssse3)
  29. PRED4x4(down_right, 10, avx)
  30. PRED4x4(vertical_left, 10, sse2)
  31. PRED4x4(vertical_left, 10, avx)
  32. PRED4x4(vertical_right, 10, sse2)
  33. PRED4x4(vertical_right, 10, ssse3)
  34. PRED4x4(vertical_right, 10, avx)
  35. PRED4x4(horizontal_up, 10, mmxext)
  36. PRED4x4(horizontal_down, 10, sse2)
  37. PRED4x4(horizontal_down, 10, ssse3)
  38. PRED4x4(horizontal_down, 10, avx)
  39. #define PRED8x8(TYPE, DEPTH, OPT) \
  40. void ff_pred8x8_ ## TYPE ## _ ## DEPTH ## _ ## OPT (uint8_t *src, int stride);
  41. PRED8x8(dc, 10, mmxext)
  42. PRED8x8(dc, 10, sse2)
  43. PRED8x8(top_dc, 10, sse2)
  44. PRED8x8(plane, 10, sse2)
  45. PRED8x8(vertical, 10, sse2)
  46. PRED8x8(horizontal, 10, sse2)
  47. #define PRED8x8L(TYPE, DEPTH, OPT)\
  48. void ff_pred8x8l_ ## TYPE ## _ ## DEPTH ## _ ## OPT (uint8_t *src, int has_topleft, int has_topright, int stride);
  49. PRED8x8L(dc, 10, sse2)
  50. PRED8x8L(dc, 10, avx)
  51. PRED8x8L(128_dc, 10, mmxext)
  52. PRED8x8L(128_dc, 10, sse2)
  53. PRED8x8L(top_dc, 10, sse2)
  54. PRED8x8L(top_dc, 10, avx)
  55. PRED8x8L(vertical, 10, sse2)
  56. PRED8x8L(vertical, 10, avx)
  57. PRED8x8L(horizontal, 10, sse2)
  58. PRED8x8L(horizontal, 10, ssse3)
  59. PRED8x8L(horizontal, 10, avx)
  60. PRED8x8L(down_left, 10, sse2)
  61. PRED8x8L(down_left, 10, ssse3)
  62. PRED8x8L(down_left, 10, avx)
  63. PRED8x8L(down_right, 10, sse2)
  64. PRED8x8L(down_right, 10, ssse3)
  65. PRED8x8L(down_right, 10, avx)
  66. PRED8x8L(vertical_right, 10, sse2)
  67. PRED8x8L(vertical_right, 10, ssse3)
  68. PRED8x8L(vertical_right, 10, avx)
  69. PRED8x8L(horizontal_up, 10, sse2)
  70. PRED8x8L(horizontal_up, 10, ssse3)
  71. PRED8x8L(horizontal_up, 10, avx)
  72. #define PRED16x16(TYPE, DEPTH, OPT)\
  73. void ff_pred16x16_ ## TYPE ## _ ## DEPTH ## _ ## OPT (uint8_t *src, int stride);
  74. PRED16x16(dc, 10, mmxext)
  75. PRED16x16(dc, 10, sse2)
  76. PRED16x16(top_dc, 10, mmxext)
  77. PRED16x16(top_dc, 10, sse2)
  78. PRED16x16(128_dc, 10, mmxext)
  79. PRED16x16(128_dc, 10, sse2)
  80. PRED16x16(left_dc, 10, mmxext)
  81. PRED16x16(left_dc, 10, sse2)
  82. PRED16x16(vertical, 10, mmxext)
  83. PRED16x16(vertical, 10, sse2)
  84. PRED16x16(horizontal, 10, mmxext)
  85. PRED16x16(horizontal, 10, sse2)
  86. void ff_pred16x16_vertical_mmx (uint8_t *src, int stride);
  87. void ff_pred16x16_vertical_sse (uint8_t *src, int stride);
  88. void ff_pred16x16_horizontal_mmx (uint8_t *src, int stride);
  89. void ff_pred16x16_horizontal_mmx2 (uint8_t *src, int stride);
  90. void ff_pred16x16_horizontal_ssse3 (uint8_t *src, int stride);
  91. void ff_pred16x16_dc_mmx2 (uint8_t *src, int stride);
  92. void ff_pred16x16_dc_sse2 (uint8_t *src, int stride);
  93. void ff_pred16x16_dc_ssse3 (uint8_t *src, int stride);
  94. void ff_pred16x16_plane_h264_mmx (uint8_t *src, int stride);
  95. void ff_pred16x16_plane_h264_mmx2 (uint8_t *src, int stride);
  96. void ff_pred16x16_plane_h264_sse2 (uint8_t *src, int stride);
  97. void ff_pred16x16_plane_h264_ssse3 (uint8_t *src, int stride);
  98. void ff_pred16x16_plane_rv40_mmx (uint8_t *src, int stride);
  99. void ff_pred16x16_plane_rv40_mmx2 (uint8_t *src, int stride);
  100. void ff_pred16x16_plane_rv40_sse2 (uint8_t *src, int stride);
  101. void ff_pred16x16_plane_rv40_ssse3 (uint8_t *src, int stride);
  102. void ff_pred16x16_plane_svq3_mmx (uint8_t *src, int stride);
  103. void ff_pred16x16_plane_svq3_mmx2 (uint8_t *src, int stride);
  104. void ff_pred16x16_plane_svq3_sse2 (uint8_t *src, int stride);
  105. void ff_pred16x16_plane_svq3_ssse3 (uint8_t *src, int stride);
  106. void ff_pred16x16_tm_vp8_mmx (uint8_t *src, int stride);
  107. void ff_pred16x16_tm_vp8_mmx2 (uint8_t *src, int stride);
  108. void ff_pred16x16_tm_vp8_sse2 (uint8_t *src, int stride);
  109. void ff_pred8x8_top_dc_mmxext (uint8_t *src, int stride);
  110. void ff_pred8x8_dc_rv40_mmxext (uint8_t *src, int stride);
  111. void ff_pred8x8_dc_mmxext (uint8_t *src, int stride);
  112. void ff_pred8x8_vertical_mmx (uint8_t *src, int stride);
  113. void ff_pred8x8_horizontal_mmx (uint8_t *src, int stride);
  114. void ff_pred8x8_horizontal_mmx2 (uint8_t *src, int stride);
  115. void ff_pred8x8_horizontal_ssse3 (uint8_t *src, int stride);
  116. void ff_pred8x8_plane_mmx (uint8_t *src, int stride);
  117. void ff_pred8x8_plane_mmx2 (uint8_t *src, int stride);
  118. void ff_pred8x8_plane_sse2 (uint8_t *src, int stride);
  119. void ff_pred8x8_plane_ssse3 (uint8_t *src, int stride);
  120. void ff_pred8x8_tm_vp8_mmx (uint8_t *src, int stride);
  121. void ff_pred8x8_tm_vp8_mmx2 (uint8_t *src, int stride);
  122. void ff_pred8x8_tm_vp8_sse2 (uint8_t *src, int stride);
  123. void ff_pred8x8_tm_vp8_ssse3 (uint8_t *src, int stride);
  124. void ff_pred8x8l_top_dc_mmxext (uint8_t *src, int has_topleft, int has_topright, int stride);
  125. void ff_pred8x8l_top_dc_ssse3 (uint8_t *src, int has_topleft, int has_topright, int stride);
  126. void ff_pred8x8l_dc_mmxext (uint8_t *src, int has_topleft, int has_topright, int stride);
  127. void ff_pred8x8l_dc_ssse3 (uint8_t *src, int has_topleft, int has_topright, int stride);
  128. void ff_pred8x8l_horizontal_mmxext (uint8_t *src, int has_topleft, int has_topright, int stride);
  129. void ff_pred8x8l_horizontal_ssse3 (uint8_t *src, int has_topleft, int has_topright, int stride);
  130. void ff_pred8x8l_vertical_mmxext (uint8_t *src, int has_topleft, int has_topright, int stride);
  131. void ff_pred8x8l_vertical_ssse3 (uint8_t *src, int has_topleft, int has_topright, int stride);
  132. void ff_pred8x8l_down_left_mmxext (uint8_t *src, int has_topleft, int has_topright, int stride);
  133. void ff_pred8x8l_down_left_sse2 (uint8_t *src, int has_topleft, int has_topright, int stride);
  134. void ff_pred8x8l_down_left_ssse3 (uint8_t *src, int has_topleft, int has_topright, int stride);
  135. void ff_pred8x8l_down_right_mmxext (uint8_t *src, int has_topleft, int has_topright, int stride);
  136. void ff_pred8x8l_down_right_sse2 (uint8_t *src, int has_topleft, int has_topright, int stride);
  137. void ff_pred8x8l_down_right_ssse3 (uint8_t *src, int has_topleft, int has_topright, int stride);
  138. void ff_pred8x8l_vertical_right_mmxext(uint8_t *src, int has_topleft, int has_topright, int stride);
  139. void ff_pred8x8l_vertical_right_sse2(uint8_t *src, int has_topleft, int has_topright, int stride);
  140. void ff_pred8x8l_vertical_right_ssse3(uint8_t *src, int has_topleft, int has_topright, int stride);
  141. void ff_pred8x8l_vertical_left_sse2(uint8_t *src, int has_topleft, int has_topright, int stride);
  142. void ff_pred8x8l_vertical_left_ssse3(uint8_t *src, int has_topleft, int has_topright, int stride);
  143. void ff_pred8x8l_horizontal_up_mmxext(uint8_t *src, int has_topleft, int has_topright, int stride);
  144. void ff_pred8x8l_horizontal_up_ssse3(uint8_t *src, int has_topleft, int has_topright, int stride);
  145. void ff_pred8x8l_horizontal_down_mmxext(uint8_t *src, int has_topleft, int has_topright, int stride);
  146. void ff_pred8x8l_horizontal_down_sse2(uint8_t *src, int has_topleft, int has_topright, int stride);
  147. void ff_pred8x8l_horizontal_down_ssse3(uint8_t *src, int has_topleft, int has_topright, int stride);
  148. void ff_pred4x4_dc_mmxext (uint8_t *src, const uint8_t *topright, int stride);
  149. void ff_pred4x4_down_left_mmxext (uint8_t *src, const uint8_t *topright, int stride);
  150. void ff_pred4x4_down_right_mmxext (uint8_t *src, const uint8_t *topright, int stride);
  151. void ff_pred4x4_vertical_left_mmxext(uint8_t *src, const uint8_t *topright, int stride);
  152. void ff_pred4x4_vertical_right_mmxext(uint8_t *src, const uint8_t *topright, int stride);
  153. void ff_pred4x4_horizontal_up_mmxext(uint8_t *src, const uint8_t *topright, int stride);
  154. void ff_pred4x4_horizontal_down_mmxext(uint8_t *src, const uint8_t *topright, int stride);
  155. void ff_pred4x4_tm_vp8_mmx (uint8_t *src, const uint8_t *topright, int stride);
  156. void ff_pred4x4_tm_vp8_mmx2 (uint8_t *src, const uint8_t *topright, int stride);
  157. void ff_pred4x4_tm_vp8_ssse3 (uint8_t *src, const uint8_t *topright, int stride);
  158. void ff_pred4x4_vertical_vp8_mmxext(uint8_t *src, const uint8_t *topright, int stride);
  159. void ff_h264_pred_init_x86(H264PredContext *h, int codec_id, const int bit_depth, const int chroma_format_idc)
  160. {
  161. #if HAVE_YASM
  162. int mm_flags = av_get_cpu_flags();
  163. if (bit_depth == 8) {
  164. if (mm_flags & AV_CPU_FLAG_MMX) {
  165. h->pred16x16[VERT_PRED8x8 ] = ff_pred16x16_vertical_mmx;
  166. h->pred16x16[HOR_PRED8x8 ] = ff_pred16x16_horizontal_mmx;
  167. if (chroma_format_idc == 1) {
  168. h->pred8x8 [VERT_PRED8x8 ] = ff_pred8x8_vertical_mmx;
  169. h->pred8x8 [HOR_PRED8x8 ] = ff_pred8x8_horizontal_mmx;
  170. }
  171. if (codec_id == CODEC_ID_VP8) {
  172. h->pred16x16[PLANE_PRED8x8 ] = ff_pred16x16_tm_vp8_mmx;
  173. h->pred8x8 [PLANE_PRED8x8 ] = ff_pred8x8_tm_vp8_mmx;
  174. h->pred4x4 [TM_VP8_PRED ] = ff_pred4x4_tm_vp8_mmx;
  175. } else {
  176. if (chroma_format_idc == 1)
  177. h->pred8x8 [PLANE_PRED8x8] = ff_pred8x8_plane_mmx;
  178. if (codec_id == CODEC_ID_SVQ3) {
  179. if (mm_flags & AV_CPU_FLAG_CMOV)
  180. h->pred16x16[PLANE_PRED8x8] = ff_pred16x16_plane_svq3_mmx;
  181. } else if (codec_id == CODEC_ID_RV40) {
  182. h->pred16x16[PLANE_PRED8x8] = ff_pred16x16_plane_rv40_mmx;
  183. } else {
  184. h->pred16x16[PLANE_PRED8x8] = ff_pred16x16_plane_h264_mmx;
  185. }
  186. }
  187. }
  188. if (mm_flags & AV_CPU_FLAG_MMX2) {
  189. h->pred16x16[HOR_PRED8x8 ] = ff_pred16x16_horizontal_mmx2;
  190. h->pred16x16[DC_PRED8x8 ] = ff_pred16x16_dc_mmx2;
  191. if (chroma_format_idc == 1)
  192. h->pred8x8[HOR_PRED8x8 ] = ff_pred8x8_horizontal_mmx2;
  193. h->pred8x8l [TOP_DC_PRED ] = ff_pred8x8l_top_dc_mmxext;
  194. h->pred8x8l [DC_PRED ] = ff_pred8x8l_dc_mmxext;
  195. h->pred8x8l [HOR_PRED ] = ff_pred8x8l_horizontal_mmxext;
  196. h->pred8x8l [VERT_PRED ] = ff_pred8x8l_vertical_mmxext;
  197. h->pred8x8l [DIAG_DOWN_RIGHT_PRED ] = ff_pred8x8l_down_right_mmxext;
  198. h->pred8x8l [VERT_RIGHT_PRED ] = ff_pred8x8l_vertical_right_mmxext;
  199. h->pred8x8l [HOR_UP_PRED ] = ff_pred8x8l_horizontal_up_mmxext;
  200. h->pred8x8l [DIAG_DOWN_LEFT_PRED ] = ff_pred8x8l_down_left_mmxext;
  201. h->pred8x8l [HOR_DOWN_PRED ] = ff_pred8x8l_horizontal_down_mmxext;
  202. h->pred4x4 [DIAG_DOWN_RIGHT_PRED ] = ff_pred4x4_down_right_mmxext;
  203. h->pred4x4 [VERT_RIGHT_PRED ] = ff_pred4x4_vertical_right_mmxext;
  204. h->pred4x4 [HOR_DOWN_PRED ] = ff_pred4x4_horizontal_down_mmxext;
  205. h->pred4x4 [DC_PRED ] = ff_pred4x4_dc_mmxext;
  206. if (codec_id == CODEC_ID_VP8 || codec_id == CODEC_ID_H264) {
  207. h->pred4x4 [DIAG_DOWN_LEFT_PRED] = ff_pred4x4_down_left_mmxext;
  208. }
  209. if (codec_id == CODEC_ID_SVQ3 || codec_id == CODEC_ID_H264) {
  210. h->pred4x4 [VERT_LEFT_PRED ] = ff_pred4x4_vertical_left_mmxext;
  211. }
  212. if (codec_id != CODEC_ID_RV40) {
  213. h->pred4x4 [HOR_UP_PRED ] = ff_pred4x4_horizontal_up_mmxext;
  214. }
  215. if (codec_id == CODEC_ID_SVQ3 || codec_id == CODEC_ID_H264) {
  216. if (chroma_format_idc == 1) {
  217. h->pred8x8[TOP_DC_PRED8x8 ] = ff_pred8x8_top_dc_mmxext;
  218. h->pred8x8[DC_PRED8x8 ] = ff_pred8x8_dc_mmxext;
  219. }
  220. }
  221. if (codec_id == CODEC_ID_VP8) {
  222. h->pred16x16[PLANE_PRED8x8 ] = ff_pred16x16_tm_vp8_mmx2;
  223. h->pred8x8 [DC_PRED8x8 ] = ff_pred8x8_dc_rv40_mmxext;
  224. h->pred8x8 [PLANE_PRED8x8 ] = ff_pred8x8_tm_vp8_mmx2;
  225. h->pred4x4 [TM_VP8_PRED ] = ff_pred4x4_tm_vp8_mmx2;
  226. h->pred4x4 [VERT_PRED ] = ff_pred4x4_vertical_vp8_mmxext;
  227. } else {
  228. if (chroma_format_idc == 1)
  229. h->pred8x8 [PLANE_PRED8x8] = ff_pred8x8_plane_mmx2;
  230. if (codec_id == CODEC_ID_SVQ3) {
  231. h->pred16x16[PLANE_PRED8x8 ] = ff_pred16x16_plane_svq3_mmx2;
  232. } else if (codec_id == CODEC_ID_RV40) {
  233. h->pred16x16[PLANE_PRED8x8 ] = ff_pred16x16_plane_rv40_mmx2;
  234. } else {
  235. h->pred16x16[PLANE_PRED8x8 ] = ff_pred16x16_plane_h264_mmx2;
  236. }
  237. }
  238. }
  239. if (mm_flags & AV_CPU_FLAG_SSE) {
  240. h->pred16x16[VERT_PRED8x8] = ff_pred16x16_vertical_sse;
  241. }
  242. if (mm_flags & AV_CPU_FLAG_SSE2) {
  243. h->pred16x16[DC_PRED8x8 ] = ff_pred16x16_dc_sse2;
  244. h->pred8x8l [DIAG_DOWN_LEFT_PRED ] = ff_pred8x8l_down_left_sse2;
  245. h->pred8x8l [DIAG_DOWN_RIGHT_PRED ] = ff_pred8x8l_down_right_sse2;
  246. h->pred8x8l [VERT_RIGHT_PRED ] = ff_pred8x8l_vertical_right_sse2;
  247. h->pred8x8l [VERT_LEFT_PRED ] = ff_pred8x8l_vertical_left_sse2;
  248. h->pred8x8l [HOR_DOWN_PRED ] = ff_pred8x8l_horizontal_down_sse2;
  249. if (codec_id == CODEC_ID_VP8) {
  250. h->pred16x16[PLANE_PRED8x8 ] = ff_pred16x16_tm_vp8_sse2;
  251. h->pred8x8 [PLANE_PRED8x8 ] = ff_pred8x8_tm_vp8_sse2;
  252. } else {
  253. if (chroma_format_idc == 1)
  254. h->pred8x8 [PLANE_PRED8x8] = ff_pred8x8_plane_sse2;
  255. if (codec_id == CODEC_ID_SVQ3) {
  256. h->pred16x16[PLANE_PRED8x8] = ff_pred16x16_plane_svq3_sse2;
  257. } else if (codec_id == CODEC_ID_RV40) {
  258. h->pred16x16[PLANE_PRED8x8] = ff_pred16x16_plane_rv40_sse2;
  259. } else {
  260. h->pred16x16[PLANE_PRED8x8] = ff_pred16x16_plane_h264_sse2;
  261. }
  262. }
  263. }
  264. if (mm_flags & AV_CPU_FLAG_SSSE3) {
  265. h->pred16x16[HOR_PRED8x8 ] = ff_pred16x16_horizontal_ssse3;
  266. h->pred16x16[DC_PRED8x8 ] = ff_pred16x16_dc_ssse3;
  267. if (chroma_format_idc == 1)
  268. h->pred8x8 [HOR_PRED8x8 ] = ff_pred8x8_horizontal_ssse3;
  269. h->pred8x8l [TOP_DC_PRED ] = ff_pred8x8l_top_dc_ssse3;
  270. h->pred8x8l [DC_PRED ] = ff_pred8x8l_dc_ssse3;
  271. h->pred8x8l [HOR_PRED ] = ff_pred8x8l_horizontal_ssse3;
  272. h->pred8x8l [VERT_PRED ] = ff_pred8x8l_vertical_ssse3;
  273. h->pred8x8l [DIAG_DOWN_LEFT_PRED ] = ff_pred8x8l_down_left_ssse3;
  274. h->pred8x8l [DIAG_DOWN_RIGHT_PRED ] = ff_pred8x8l_down_right_ssse3;
  275. h->pred8x8l [VERT_RIGHT_PRED ] = ff_pred8x8l_vertical_right_ssse3;
  276. h->pred8x8l [VERT_LEFT_PRED ] = ff_pred8x8l_vertical_left_ssse3;
  277. h->pred8x8l [HOR_UP_PRED ] = ff_pred8x8l_horizontal_up_ssse3;
  278. h->pred8x8l [HOR_DOWN_PRED ] = ff_pred8x8l_horizontal_down_ssse3;
  279. if (codec_id == CODEC_ID_VP8) {
  280. h->pred8x8 [PLANE_PRED8x8 ] = ff_pred8x8_tm_vp8_ssse3;
  281. h->pred4x4 [TM_VP8_PRED ] = ff_pred4x4_tm_vp8_ssse3;
  282. } else {
  283. if (chroma_format_idc == 1)
  284. h->pred8x8 [PLANE_PRED8x8] = ff_pred8x8_plane_ssse3;
  285. if (codec_id == CODEC_ID_SVQ3) {
  286. h->pred16x16[PLANE_PRED8x8] = ff_pred16x16_plane_svq3_ssse3;
  287. } else if (codec_id == CODEC_ID_RV40) {
  288. h->pred16x16[PLANE_PRED8x8] = ff_pred16x16_plane_rv40_ssse3;
  289. } else {
  290. h->pred16x16[PLANE_PRED8x8] = ff_pred16x16_plane_h264_ssse3;
  291. }
  292. }
  293. }
  294. } else if (bit_depth == 10) {
  295. if (mm_flags & AV_CPU_FLAG_MMX2) {
  296. h->pred4x4[DC_PRED ] = ff_pred4x4_dc_10_mmxext;
  297. h->pred4x4[HOR_UP_PRED ] = ff_pred4x4_horizontal_up_10_mmxext;
  298. if (chroma_format_idc == 1)
  299. h->pred8x8[DC_PRED8x8 ] = ff_pred8x8_dc_10_mmxext;
  300. h->pred8x8l[DC_128_PRED ] = ff_pred8x8l_128_dc_10_mmxext;
  301. h->pred16x16[DC_PRED8x8 ] = ff_pred16x16_dc_10_mmxext;
  302. h->pred16x16[TOP_DC_PRED8x8 ] = ff_pred16x16_top_dc_10_mmxext;
  303. h->pred16x16[DC_128_PRED8x8 ] = ff_pred16x16_128_dc_10_mmxext;
  304. h->pred16x16[LEFT_DC_PRED8x8 ] = ff_pred16x16_left_dc_10_mmxext;
  305. h->pred16x16[VERT_PRED8x8 ] = ff_pred16x16_vertical_10_mmxext;
  306. h->pred16x16[HOR_PRED8x8 ] = ff_pred16x16_horizontal_10_mmxext;
  307. }
  308. if (mm_flags & AV_CPU_FLAG_SSE2) {
  309. h->pred4x4[DIAG_DOWN_LEFT_PRED ] = ff_pred4x4_down_left_10_sse2;
  310. h->pred4x4[DIAG_DOWN_RIGHT_PRED] = ff_pred4x4_down_right_10_sse2;
  311. h->pred4x4[VERT_LEFT_PRED ] = ff_pred4x4_vertical_left_10_sse2;
  312. h->pred4x4[VERT_RIGHT_PRED ] = ff_pred4x4_vertical_right_10_sse2;
  313. h->pred4x4[HOR_DOWN_PRED ] = ff_pred4x4_horizontal_down_10_sse2;
  314. if (chroma_format_idc == 1) {
  315. h->pred8x8[DC_PRED8x8 ] = ff_pred8x8_dc_10_sse2;
  316. h->pred8x8[TOP_DC_PRED8x8 ] = ff_pred8x8_top_dc_10_sse2;
  317. h->pred8x8[PLANE_PRED8x8 ] = ff_pred8x8_plane_10_sse2;
  318. h->pred8x8[VERT_PRED8x8 ] = ff_pred8x8_vertical_10_sse2;
  319. h->pred8x8[HOR_PRED8x8 ] = ff_pred8x8_horizontal_10_sse2;
  320. }
  321. h->pred8x8l[VERT_PRED ] = ff_pred8x8l_vertical_10_sse2;
  322. h->pred8x8l[HOR_PRED ] = ff_pred8x8l_horizontal_10_sse2;
  323. h->pred8x8l[DC_PRED ] = ff_pred8x8l_dc_10_sse2;
  324. h->pred8x8l[DC_128_PRED ] = ff_pred8x8l_128_dc_10_sse2;
  325. h->pred8x8l[TOP_DC_PRED ] = ff_pred8x8l_top_dc_10_sse2;
  326. h->pred8x8l[DIAG_DOWN_LEFT_PRED ] = ff_pred8x8l_down_left_10_sse2;
  327. h->pred8x8l[DIAG_DOWN_RIGHT_PRED] = ff_pred8x8l_down_right_10_sse2;
  328. h->pred8x8l[VERT_RIGHT_PRED ] = ff_pred8x8l_vertical_right_10_sse2;
  329. h->pred8x8l[HOR_UP_PRED ] = ff_pred8x8l_horizontal_up_10_sse2;
  330. h->pred16x16[DC_PRED8x8 ] = ff_pred16x16_dc_10_sse2;
  331. h->pred16x16[TOP_DC_PRED8x8 ] = ff_pred16x16_top_dc_10_sse2;
  332. h->pred16x16[DC_128_PRED8x8 ] = ff_pred16x16_128_dc_10_sse2;
  333. h->pred16x16[LEFT_DC_PRED8x8 ] = ff_pred16x16_left_dc_10_sse2;
  334. h->pred16x16[VERT_PRED8x8 ] = ff_pred16x16_vertical_10_sse2;
  335. h->pred16x16[HOR_PRED8x8 ] = ff_pred16x16_horizontal_10_sse2;
  336. }
  337. if (mm_flags & AV_CPU_FLAG_SSSE3) {
  338. h->pred4x4[DIAG_DOWN_RIGHT_PRED] = ff_pred4x4_down_right_10_ssse3;
  339. h->pred4x4[VERT_RIGHT_PRED ] = ff_pred4x4_vertical_right_10_ssse3;
  340. h->pred4x4[HOR_DOWN_PRED ] = ff_pred4x4_horizontal_down_10_ssse3;
  341. h->pred8x8l[HOR_PRED ] = ff_pred8x8l_horizontal_10_ssse3;
  342. h->pred8x8l[DIAG_DOWN_LEFT_PRED ] = ff_pred8x8l_down_left_10_ssse3;
  343. h->pred8x8l[DIAG_DOWN_RIGHT_PRED] = ff_pred8x8l_down_right_10_ssse3;
  344. h->pred8x8l[VERT_RIGHT_PRED ] = ff_pred8x8l_vertical_right_10_ssse3;
  345. h->pred8x8l[HOR_UP_PRED ] = ff_pred8x8l_horizontal_up_10_ssse3;
  346. }
  347. #if HAVE_AVX
  348. if (mm_flags & AV_CPU_FLAG_AVX) {
  349. h->pred4x4[DIAG_DOWN_LEFT_PRED ] = ff_pred4x4_down_left_10_avx;
  350. h->pred4x4[DIAG_DOWN_RIGHT_PRED] = ff_pred4x4_down_right_10_avx;
  351. h->pred4x4[VERT_LEFT_PRED ] = ff_pred4x4_vertical_left_10_avx;
  352. h->pred4x4[VERT_RIGHT_PRED ] = ff_pred4x4_vertical_right_10_avx;
  353. h->pred4x4[HOR_DOWN_PRED ] = ff_pred4x4_horizontal_down_10_avx;
  354. h->pred8x8l[VERT_PRED ] = ff_pred8x8l_vertical_10_avx;
  355. h->pred8x8l[HOR_PRED ] = ff_pred8x8l_horizontal_10_avx;
  356. h->pred8x8l[DC_PRED ] = ff_pred8x8l_dc_10_avx;
  357. h->pred8x8l[TOP_DC_PRED ] = ff_pred8x8l_top_dc_10_avx;
  358. h->pred8x8l[DIAG_DOWN_RIGHT_PRED] = ff_pred8x8l_down_right_10_avx;
  359. h->pred8x8l[DIAG_DOWN_LEFT_PRED ] = ff_pred8x8l_down_left_10_avx;
  360. h->pred8x8l[VERT_RIGHT_PRED ] = ff_pred8x8l_vertical_right_10_avx;
  361. h->pred8x8l[HOR_UP_PRED ] = ff_pred8x8l_horizontal_up_10_avx;
  362. }
  363. #endif /* HAVE_AVX */
  364. }
  365. #endif /* HAVE_YASM */
  366. }