1 /*
2 * Copyright (C) 2004 Michael Niedermayer <michaelni@gmx.at>
3 *
4 * This file is part of FFmpeg.
5 *
6 * FFmpeg is free software; you can redistribute it and/or
7 * modify it under the terms of the GNU Lesser General Public
8 * License as published by the Free Software Foundation; either
9 * version 2.1 of the License, or (at your option) any later version.
10 *
11 * FFmpeg is distributed in the hope that it will be useful,
12 * but WITHOUT ANY WARRANTY; without even the implied warranty of
13 * MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the GNU
14 * Lesser General Public License for more details.
15 *
16 * You should have received a copy of the GNU Lesser General Public
17 * License along with FFmpeg; if not, write to the Free Software
18 * Foundation, Inc., 51 Franklin Street, Fifth Floor, Boston, MA 02110-1301 USA
19 */
20
30
34
35
40 for(y=0; y<b_h; y++){
41 //FIXME ugly misuse of obmc_stride
42 const uint8_t *obmc1= obmc + y*obmc_stride;
43 const uint8_t *obmc2= obmc1+ (obmc_stride>>1);
44 const uint8_t *obmc3= obmc1+ obmc_stride*(obmc_stride>>1);
47 for(x=0; x<b_w; x++){
48 int v= obmc1[x] * block[3][x + y*src_stride]
49 +obmc2[x] * block[2][x + y*src_stride]
50 +obmc3[x] * block[1][x + y*src_stride]
51 +obmc4[x] * block[0][x + y*src_stride];
52
56 }
57 if(add){
58 v += dst[x + src_x];
60 if(v&(~255)) v= ~(v>>31);
61 dst8[x + y*src_stride] =
v;
62 }else{
64 }
65 }
66 }
67 }
68
70 {
73
76 if (edges_needed) {
79 }
82 if (edges_needed) {
83 for (i = 0; frame->
data[i]; i++) {
88 }
91 }
92
93 return 0;
94 }
95
97 int plane_index,
level, orientation;
98
99 for(plane_index=0; plane_index<3; plane_index++){
101 for(orientation=level ? 1:0; orientation<4; orientation++){
103 }
104 }
105 }
108 }
109
113
116
121
122 return 0;
123 }
124
126 int i;
128
129 for(i=0; i<
QROOT; i++){
131 v *= pow(2, 1.0 / QROOT);
132 }
133 }
136 8,7,6,5,4,3,2,1,
137 7,7,0,0,0,0,0,1,
138 6,0,6,0,0,0,2,0,
139 5,0,0,5,0,3,0,0,
140 4,0,0,0,4,0,0,0,
141 3,0,0,5,0,3,0,0,
142 2,0,6,0,0,0,2,0,
143 1,7,0,0,0,0,0,1,
144 };
145
146 static const uint8_t brane[256]={
147 0x00,0x01,0x01,0x01,0x01,0x01,0x01,0x01,0x11,0x12,0x12,0x12,0x12,0x12,0x12,0x12,
148 0x04,0x05,0xcc,0xcc,0xcc,0xcc,0xcc,0x41,0x15,0x16,0xcc,0xcc,0xcc,0xcc,0xcc,0x52,
149 0x04,0xcc,0x05,0xcc,0xcc,0xcc,0x41,0xcc,0x15,0xcc,0x16,0xcc,0xcc,0xcc,0x52,0xcc,
150 0x04,0xcc,0xcc,0x05,0xcc,0x41,0xcc,0xcc,0x15,0xcc,0xcc,0x16,0xcc,0x52,0xcc,0xcc,
151 0x04,0xcc,0xcc,0xcc,0x41,0xcc,0xcc,0xcc,0x15,0xcc,0xcc,0xcc,0x16,0xcc,0xcc,0xcc,
152 0x04,0xcc,0xcc,0x41,0xcc,0x05,0xcc,0xcc,0x15,0xcc,0xcc,0x52,0xcc,0x16,0xcc,0xcc,
153 0x04,0xcc,0x41,0xcc,0xcc,0xcc,0x05,0xcc,0x15,0xcc,0x52,0xcc,0xcc,0xcc,0x16,0xcc,
154 0x04,0x41,0xcc,0xcc,0xcc,0xcc,0xcc,0x05,0x15,0x52,0xcc,0xcc,0xcc,0xcc,0xcc,0x16,
155 0x44,0x45,0x45,0x45,0x45,0x45,0x45,0x45,0x55,0x56,0x56,0x56,0x56,0x56,0x56,0x56,
156 0x48,0x49,0xcc,0xcc,0xcc,0xcc,0xcc,0x85,0x59,0x5A,0xcc,0xcc,0xcc,0xcc,0xcc,0x96,
157 0x48,0xcc,0x49,0xcc,0xcc,0xcc,0x85,0xcc,0x59,0xcc,0x5A,0xcc,0xcc,0xcc,0x96,0xcc,
158 0x48,0xcc,0xcc,0x49,0xcc,0x85,0xcc,0xcc,0x59,0xcc,0xcc,0x5A,0xcc,0x96,0xcc,0xcc,
159 0x48,0xcc,0xcc,0xcc,0x49,0xcc,0xcc,0xcc,0x59,0xcc,0xcc,0xcc,0x96,0xcc,0xcc,0xcc,
160 0x48,0xcc,0xcc,0x85,0xcc,0x49,0xcc,0xcc,0x59,0xcc,0xcc,0x96,0xcc,0x5A,0xcc,0xcc,
161 0x48,0xcc,0x85,0xcc,0xcc,0xcc,0x49,0xcc,0x59,0xcc,0x96,0xcc,0xcc,0xcc,0x5A,0xcc,
162 0x48,0x85,0xcc,0xcc,0xcc,0xcc,0xcc,0x49,0x59,0x96,0xcc,0xcc,0xcc,0xcc,0xcc,0x5A,
163 };
164
165 static const uint8_t needs[16]={
166 0,1,0,0,
167 2,4,2,0,
168 0,1,0,0,
169 15
170 };
171
175 int16_t *tmpI= tmpIt;
179 r= brane[dx + 16*dy]&15;
180 l= brane[dx + 16*dy]>>4;
181
182 b= needs[l] | needs[
r];
184 b= 15;
185
186 if(b&5){
188 for(x=0; x < b_w; x++){
189 int a_1=src[x + HTAPS_MAX/2-4];
190 int a0= src[x + HTAPS_MAX/2-3];
191 int a1= src[x + HTAPS_MAX/2-2];
192 int a2= src[x + HTAPS_MAX/2-1];
193 int a3= src[x + HTAPS_MAX/2+0];
194 int a4= src[x + HTAPS_MAX/2+1];
195 int a5= src[x + HTAPS_MAX/2+2];
196 int a6= src[x + HTAPS_MAX/2+3];
197 int am=0;
199 am= 20*(a2+
a3) - 5*(a1+a4) + (a0+
a5);
200 tmpI[x]= am;
201 am= (am+16)>>5;
202 }else{
204 tmpI[x]= am;
205 am= (am+32)>>6;
206 }
207
208 if(am&(~255)) am= ~(am>>31);
209 tmp2[x]= am;
210 }
211 tmpI+= 64;
212 tmp2+= 64;
213 src += stride;
214 }
216 }
218 tmp2= tmp2t[1];
219
220 if(b&2){
221 for(y=0; y < b_h; y++){
222 for(x=0; x < b_w+1; x++){
231 int am=0;
233 am= (20*(a2+
a3) - 5*(a1+a4) + (a0+
a5) + 16)>>5;
234 else
236
237 if(am&(~255)) am= ~(am>>31);
238 tmp2[x]= am;
239 }
240 src += stride;
241 tmp2+= 64;
242 }
244 }
246 tmp2= tmp2t[2];
247 tmpI= tmpIt;
248 if(b&4){
249 for(y=0; y < b_h; y++){
250 for(x=0; x < b_w; x++){
259 int am=0;
261 am= (20*(a2+
a3) - 5*(a1+a4) + (a0+
a5) + 512)>>10;
262 else
264 if(am&(~255)) am= ~(am>>31);
265 tmp2[x]= am;
266 }
267 tmpI+= 64;
268 tmp2+= 64;
269 }
270 }
271
274 hpel[ 2]= src + 1;
275
276 hpel[ 4]= tmp2t[1];
277 hpel[ 5]= tmp2t[2];
278 hpel[ 6]= tmp2t[1] + 1;
279
280 hpel[ 8]= src + stride;
281 hpel[ 9]= hpel[1] + 64;
282 hpel[10]= hpel[8] + 1;
283
284 #define MC_STRIDE(x) (needs[x] ? 64 : stride)
285
286 if(b==15){
287 int dxy = dx / 8 + dy / 8 * 4;
288 const uint8_t *src1 = hpel[dxy ];
289 const uint8_t *src2 = hpel[dxy + 1];
290 const uint8_t *src3 = hpel[dxy + 4];
291 const uint8_t *src4 = hpel[dxy + 5];
296 dx&=7;
297 dy&=7;
298 for(y=0; y < b_h; y++){
299 for(x=0; x < b_w; x++){
300 dst[x]= ((8-dx)*(8-dy)*src1[x] + dx*(8-dy)*src2[x]+
301 (8-dx)* dy *src3[x] + dx* dy *src4[x]+32)>>6;
302 }
303 src1+=stride1;
304 src2+=stride2;
305 src3+=stride3;
306 src4+=stride4;
307 dst +=stride;
308 }
309 }else{
314 int a= weight[((dx&7) + (8*(dy&7)))];
316 for(y=0; y < b_h; y++){
317 for(x=0; x < b_w; x++){
318 dst[x]= (a*src1[x] + b*src2[x] + 4)>>3;
319 }
320 src1+=stride1;
321 src2+=stride2;
322 dst +=stride;
323 }
324 }
325 }
326
327 void ff_snow_pred_block(
SnowContext *
s,
uint8_t *dst,
uint8_t *tmp, ptrdiff_t
stride,
int sx,
int sy,
int b_w,
int b_h,
BlockNode *
block,
int plane_index,
int w,
int h){
330 const unsigned color = block->
color[plane_index];
331 const unsigned color4 = color*0x01010101;
332 if(b_w==32){
333 for(y=0; y < b_h; y++){
334 *(uint32_t*)&dst[0 + y*stride]= color4;
335 *(uint32_t*)&dst[4 + y*stride]= color4;
336 *(uint32_t*)&dst[8 + y*stride]= color4;
337 *(uint32_t*)&dst[12+ y*stride]= color4;
338 *(uint32_t*)&dst[16+ y*stride]= color4;
339 *(uint32_t*)&dst[20+ y*stride]= color4;
340 *(uint32_t*)&dst[24+ y*stride]= color4;
341 *(uint32_t*)&dst[28+ y*stride]= color4;
342 }
343 }else if(b_w==16){
344 for(y=0; y < b_h; y++){
345 *(uint32_t*)&dst[0 + y*stride]= color4;
346 *(uint32_t*)&dst[4 + y*stride]= color4;
347 *(uint32_t*)&dst[8 + y*stride]= color4;
348 *(uint32_t*)&dst[12+ y*stride]= color4;
349 }
350 }else if(b_w==8){
351 for(y=0; y < b_h; y++){
352 *(uint32_t*)&dst[0 + y*stride]= color4;
353 *(uint32_t*)&dst[4 + y*stride]= color4;
354 }
355 }else if(b_w==4){
356 for(y=0; y < b_h; y++){
357 *(uint32_t*)&dst[0 + y*stride]= color4;
358 }
359 }else{
360 for(y=0; y < b_h; y++){
361 for(x=0; x < b_w; x++){
362 dst[x + y*stride]=
color;
363 }
364 }
365 }
366 }else{
369 int mx= block->
mx*scale;
370 int my= block->
my*scale;
371 const int dx= mx&15;
372 const int dy= my&15;
373 const int tab_index= 3 - (b_w>>2) + (b_w>>4);
376 src += sx + sy*stride;
380 stride, stride,
382 sx, sy, w, h);
384 }
385
387
388 av_assert2((tab_index>=0 && tab_index<4) || b_w==32);
389 if( (dx&3) || (dy&3)
390 || !(b_w == b_h || 2*b_w == b_h || b_w == 2*b_h)
391 || (b_w&(b_w-1))
392 || b_w == 1
393 || b_h == 1
395 mc_block(&s->
plane[plane_index], dst, src, stride, b_w, b_h, dx, dy);
396 else if(b_w==32){
398 for(y=0; y<b_h; y+=16){
401 }
402 }else if(b_w==b_h)
404 else if(b_w==2*b_h){
407 }else{
411 }
412 }
413 }
414
415 #define mca(dx,dy,b_w)\
416 static void mc_block_hpel ## dx ## dy ## b_w(uint8_t *dst, const uint8_t *src, ptrdiff_t stride, int h){\
417 av_assert2(h==b_w);\
418 mc_block(NULL, dst, src-(HTAPS_MAX/2-1)-(HTAPS_MAX/2-1)*stride, stride, b_w, b_w, dx, dy);\
419 }
420
429
433 int i, j;
434
436 s->
max_ref_frames=1;
//just make sure it's not an invalid value in case of no initial keyframe
437
443
444 #define mcf(dx,dy)\
445 s->qdsp.put_qpel_pixels_tab [0][dy+dx/4]=\
446 s->qdsp.put_no_rnd_qpel_pixels_tab[0][dy+dx/4]=\
447 s->h264qpel.put_h264_qpel_pixels_tab[0][dy+dx/4];\
448 s->qdsp.put_qpel_pixels_tab [1][dy+dx/4]=\
449 s->qdsp.put_no_rnd_qpel_pixels_tab[1][dy+dx/4]=\
450 s->h264qpel.put_h264_qpel_pixels_tab[1][dy+dx/4];
451
468
469 #define mcfh(dx,dy)\
470 s->hdsp.put_pixels_tab [0][dy/4+dx/8]=\
471 s->hdsp.put_no_rnd_pixels_tab[0][dy/4+dx/8]=\
472 mc_block_hpel ## dx ## dy ## 16;\
473 s->hdsp.put_pixels_tab [1][dy/4+dx/8]=\
474 s->hdsp.put_no_rnd_pixels_tab[1][dy/4+dx/8]=\
475 mc_block_hpel ## dx ## dy ## 8;
476
481
483
484 // dec += FFMAX(s->chroma_h_shift, s->chroma_v_shift);
485
488
494
500 goto fail;
501 }
502
506 goto fail;
507
508 return 0;
509 fail:
511 }
512
515 int plane_index,
level, orientation;
516 int ret, emu_buf_size;
517
525 }
526
530 }
531
532 for(plane_index=0; plane_index < s->
nb_planes; plane_index++){
535
536 if(plane_index){
539 }
542
544 for(orientation=level ? 1 : 0; orientation<4; orientation++){
546
550 b->
width = (w + !(orientation&1))>>1;
551 b->
height= (h + !(orientation>1))>>1;
552
556
557 if(orientation&1){
560 }
561 if(orientation>1){
564 }
566
567 if(level)
569 //FIXME avoid this realloc
573 goto fail;
574 }
575 w= (w+1)>>1;
576 h= (h+1)>>1;
577 }
578 }
579
580 return 0;
581 fail:
583 }
584
585 #define USE_HALFPEL_PLANE 0
586
589
591 int is_chroma= !!p;
596
600 if (!halfpel[1][p] || !halfpel[2][p] || !halfpel[3][p])
602
604 for(y=0; y<h; y++){
605 for(x=0; x<w; x++){
606 int i= y*ls + x;
607
608 halfpel[1][p][i]= (20*(src[i] + src[i+1]) - 5*(src[i-1] + src[i+2]) + (src[i-2] + src[i+3]) + 16 )>>5;
609 }
610 }
611 for(y=0; y<h; y++){
612 for(x=0; x<w; x++){
613 int i= y*ls + x;
614
615 halfpel[2][p][i]= (20*(src[i] + src[i+ls]) - 5*(src[i-ls] + src[i+2*ls]) + (src[i-2*ls] + src[i+3*ls]) + 16 )>>5;
616 }
617 }
618 src= halfpel[1][p];
619 for(y=0; y<h; y++){
620 for(x=0; x<w; x++){
621 int i= y*ls + x;
622
623 halfpel[3][p][i]= (20*(src[i] + src[i+ls]) - 5*(src[i-ls] + src[i+2*ls]) + (src[i-2*ls] + src[i+3*ls]) + 16 )>>5;
624 }
625 }
626
627 //FIXME border!
628 }
629 return 0;
630 }
631
633 {
635 int i;
636
639 for(i=0; i<9; i++)
642 }
643 }
644
648
650
658 }
661
664 }else{
665 int i;
668 break;
672 return -1;
673 }
674 }
677
679
680 return 0;
681 }
682
684 {
685 int plane_index,
level, orientation, i;
686
692
698
702
708 }
710 }
711
712 for(plane_index=0; plane_index < s->
nb_planes; plane_index++){
714 for(orientation=level ? 1 : 0; orientation<4; orientation++){
716
718 }
719 }
720 }
723 }