[svn] / branches / release-1_3-branch / xvidcore / src / dct / idct.c Repository:
ViewVC logotype

Annotation of /branches/release-1_3-branch/xvidcore/src/dct/idct.c

Parent Directory Parent Directory | Revision Log Revision Log


Revision 1653 - (view) (download)
Original Path: trunk/xvidcore/src/dct/idct.c

1 : edgomez 1382 /*****************************************************************************
2 :     *
3 :     * XVID MPEG-4 VIDEO CODEC
4 :     * - Inverse DCT -
5 :     *
6 :     * These routines are from Independent JPEG Group's free JPEG software
7 :     * Copyright (C) 1991-1998, Thomas G. Lane (see the file README.IJG)
8 :     *
9 :     * This program is free software ; you can redistribute it and/or modify
10 :     * it under the terms of the GNU General Public License as published by
11 :     * the Free Software Foundation ; either version 2 of the License, or
12 :     * (at your option) any later version.
13 :     *
14 :     * This program is distributed in the hope that it will be useful,
15 :     * but WITHOUT ANY WARRANTY ; without even the implied warranty of
16 :     * MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the
17 :     * GNU General Public License for more details.
18 :     *
19 :     * You should have received a copy of the GNU General Public License
20 :     * along with this program ; if not, write to the Free Software
21 :     * Foundation, Inc., 59 Temple Place, Suite 330, Boston, MA 02111-1307 USA
22 :     *
23 : suxen_drol 1653 * $Id: idct.c,v 1.9 2005-11-22 10:23:01 suxen_drol Exp $
24 : edgomez 1382 *
25 :     ****************************************************************************/
26 : edgomez 851
27 :     /* Copyright (C) 1996, MPEG Software Simulation Group. All Rights Reserved. */
28 :    
29 :     /*
30 :     * Disclaimer of Warranty
31 : Isibaar 3 *
32 : edgomez 851 * These software programs are available to the user without any license fee or
33 :     * royalty on an "as is" basis. The MPEG Software Simulation Group disclaims
34 :     * any and all warranties, whether express, implied, or statuary, including any
35 :     * implied warranties or merchantability or of fitness for a particular
36 :     * purpose. In no event shall the copyright-holder be liable for any
37 :     * incidental, punitive, or consequential damages of any kind whatsoever
38 :     * arising from the use of these programs.
39 : Isibaar 3 *
40 : edgomez 851 * This disclaimer of warranty extends to the user of these programs and user's
41 :     * customers, employees, agents, transferees, successors, and assigns.
42 : Isibaar 3 *
43 : edgomez 851 * The MPEG Software Simulation Group does not represent or warrant that the
44 :     * programs furnished hereunder are free of infringement of any third-party
45 :     * patents.
46 : Isibaar 3 *
47 : edgomez 851 * Commercial implementations of MPEG-1 and MPEG-2 video, including shareware,
48 :     * are subject to royalty fees to patent holders. Many of these patents are
49 :     * general enough such that they are unavoidable regardless of implementation
50 :     * design.
51 : Isibaar 3 *
52 : edgomez 851 * MPEG2AVI
53 :     * --------
54 :     * v0.16B33 renamed the initialization function to init_idct_int32()
55 :     * v0.16B32 removed the unused idct_row() and idct_col() functions
56 :     * v0.16B3 changed var declarations to static, to enforce data align
57 :     * v0.16B22 idct_FAST() renamed to idct_int32()
58 :     * also merged idct_FAST() into a single function, to help VC++
59 :     * optimize it.
60 : edgomez 1382 *
61 : edgomez 851 * v0.14 changed int to long, to avoid confusion when compiling on x86
62 :     * platform ( in VC++ "int" -> 32bits )
63 :     */
64 : Isibaar 3
65 :     /**********************************************************/
66 :     /* inverse two dimensional DCT, Chen-Wang algorithm */
67 :     /* (cf. IEEE ASSP-32, pp. 803-816, Aug. 1984) */
68 :     /* 32-bit integer arithmetic (8 bit coefficients) */
69 :     /* 11 mults, 29 adds per DCT */
70 :     /* sE, 18.8.91 */
71 :     /**********************************************************/
72 :     /* coefficients extended to 12 bit for IEEE1180-1990 */
73 :     /* compliance sE, 2.1.94 */
74 :     /**********************************************************/
75 :    
76 :     /* this code assumes >> to be a two's-complement arithmetic */
77 :     /* right shift: (-2)>>1 == -1 , (-3)>>1 == -2 */
78 :    
79 :     #include "idct.h"
80 :    
81 : edgomez 195 #define W1 2841 /* 2048*sqrt(2)*cos(1*pi/16) */
82 :     #define W2 2676 /* 2048*sqrt(2)*cos(2*pi/16) */
83 :     #define W3 2408 /* 2048*sqrt(2)*cos(3*pi/16) */
84 :     #define W5 1609 /* 2048*sqrt(2)*cos(5*pi/16) */
85 :     #define W6 1108 /* 2048*sqrt(2)*cos(6*pi/16) */
86 :     #define W7 565 /* 2048*sqrt(2)*cos(7*pi/16) */
87 : Isibaar 3
88 : edgomez 1537 /* private data
89 :     * Initialized by idct_int32_init so it's mostly RO data,
90 :     * doesn't hurt thread safety */
91 : edgomez 195 static short iclip[1024]; /* clipping table */
92 : Isibaar 3 static short *iclp;
93 :    
94 :     /* private prototypes */
95 :    
96 :     /* row (horizontal) IDCT
97 :     *
98 :     * 7 pi 1
99 :     * dst[k] = sum c[l] * src[l] * cos( -- * ( k + - ) * l )
100 :     * l=0 8 2
101 :     *
102 :     * where: c[0] = 128
103 :     * c[1..7] = 128*sqrt(2)
104 :     */
105 :    
106 : edgomez 1382 #if 0
107 : Isibaar 3 static void idctrow(blk)
108 :     short *blk;
109 :     {
110 :     int X0, X1, X2, X3, X4, X5, X6, X7, X8;
111 :    
112 : edgomez 1382 /* shortcut */
113 : Isibaar 3 if (!((X1 = blk[4]<<11) | (X2 = blk[6]) | (X3 = blk[2]) |
114 :     (X4 = blk[1]) | (X5 = blk[7]) | (X6 = blk[5]) | (X7 = blk[3])))
115 :     {
116 :     blk[0]=blk[1]=blk[2]=blk[3]=blk[4]=blk[5]=blk[6]=blk[7]=blk[0]<<3;
117 :     return;
118 :     }
119 :    
120 : edgomez 1382 X0 = (blk[0]<<11) + 128; /* for proper rounding in the fourth stage */
121 : Isibaar 3
122 : edgomez 1382 /* first stage */
123 : Isibaar 3 X8 = W7*(X4+X5);
124 :     X4 = X8 + (W1-W7)*X4;
125 :     X5 = X8 - (W1+W7)*X5;
126 :     X8 = W3*(X6+X7);
127 :     X6 = X8 - (W3-W5)*X6;
128 :     X7 = X8 - (W3+W5)*X7;
129 : edgomez 1382
130 :     /* second stage */
131 : Isibaar 3 X8 = X0 + X1;
132 :     X0 -= X1;
133 :     X1 = W6*(X3+X2);
134 :     X2 = X1 - (W2+W6)*X2;
135 :     X3 = X1 + (W2-W6)*X3;
136 :     X1 = X4 + X6;
137 :     X4 -= X6;
138 :     X6 = X5 + X7;
139 :     X5 -= X7;
140 : edgomez 1382
141 :     /* third stage */
142 : Isibaar 3 X7 = X8 + X3;
143 :     X8 -= X3;
144 :     X3 = X0 + X2;
145 :     X0 -= X2;
146 :     X2 = (181*(X4+X5)+128)>>8;
147 :     X4 = (181*(X4-X5)+128)>>8;
148 : edgomez 1382
149 :     /* fourth stage */
150 : Isibaar 3 blk[0] = (X7+X1)>>8;
151 :     blk[1] = (X3+X2)>>8;
152 :     blk[2] = (X0+X4)>>8;
153 :     blk[3] = (X8+X6)>>8;
154 :     blk[4] = (X8-X6)>>8;
155 :     blk[5] = (X0-X4)>>8;
156 :     blk[6] = (X3-X2)>>8;
157 :     blk[7] = (X7-X1)>>8;
158 : edgomez 1382 }
159 :     #endif
160 : Isibaar 3
161 :     /* column (vertical) IDCT
162 :     *
163 :     * 7 pi 1
164 :     * dst[8*k] = sum c[l] * src[8*l] * cos( -- * ( k + - ) * l )
165 :     * l=0 8 2
166 :     *
167 :     * where: c[0] = 1/1024
168 :     * c[1..7] = (1/1024)*sqrt(2)
169 :     */
170 : edgomez 1382
171 :     #if 0
172 : Isibaar 3 static void idctcol(blk)
173 :     short *blk;
174 :     {
175 :     int X0, X1, X2, X3, X4, X5, X6, X7, X8;
176 :    
177 : edgomez 1382 /* shortcut */
178 : Isibaar 3 if (!((X1 = (blk[8*4]<<8)) | (X2 = blk[8*6]) | (X3 = blk[8*2]) |
179 :     (X4 = blk[8*1]) | (X5 = blk[8*7]) | (X6 = blk[8*5]) | (X7 = blk[8*3])))
180 :     {
181 :     blk[8*0]=blk[8*1]=blk[8*2]=blk[8*3]=blk[8*4]=blk[8*5]=blk[8*6]=blk[8*7]=
182 :     iclp[(blk[8*0]+32)>>6];
183 :     return;
184 :     }
185 :    
186 :     X0 = (blk[8*0]<<8) + 8192;
187 :    
188 : edgomez 1382 /* first stage */
189 : Isibaar 3 X8 = W7*(X4+X5) + 4;
190 :     X4 = (X8+(W1-W7)*X4)>>3;
191 :     X5 = (X8-(W1+W7)*X5)>>3;
192 :     X8 = W3*(X6+X7) + 4;
193 :     X6 = (X8-(W3-W5)*X6)>>3;
194 :     X7 = (X8-(W3+W5)*X7)>>3;
195 : edgomez 1382
196 :     /* second stage */
197 : Isibaar 3 X8 = X0 + X1;
198 :     X0 -= X1;
199 :     X1 = W6*(X3+X2) + 4;
200 :     X2 = (X1-(W2+W6)*X2)>>3;
201 :     X3 = (X1+(W2-W6)*X3)>>3;
202 :     X1 = X4 + X6;
203 :     X4 -= X6;
204 :     X6 = X5 + X7;
205 :     X5 -= X7;
206 : edgomez 1382
207 :     /* third stage */
208 : Isibaar 3 X7 = X8 + X3;
209 :     X8 -= X3;
210 :     X3 = X0 + X2;
211 :     X0 -= X2;
212 :     X2 = (181*(X4+X5)+128)>>8;
213 :     X4 = (181*(X4-X5)+128)>>8;
214 : edgomez 1382
215 :     /* fourth stage */
216 : Isibaar 3 blk[8*0] = iclp[(X7+X1)>>14];
217 :     blk[8*1] = iclp[(X3+X2)>>14];
218 :     blk[8*2] = iclp[(X0+X4)>>14];
219 :     blk[8*3] = iclp[(X8+X6)>>14];
220 :     blk[8*4] = iclp[(X8-X6)>>14];
221 :     blk[8*5] = iclp[(X0-X4)>>14];
222 :     blk[8*6] = iclp[(X3-X2)>>14];
223 :     blk[8*7] = iclp[(X7-X1)>>14];
224 : edgomez 1382 }
225 :     #endif
226 : Isibaar 3
227 : edgomez 1382 /* function pointer */
228 : Isibaar 3 idctFuncPtr idct;
229 :    
230 :     /* two dimensional inverse discrete cosine transform */
231 : edgomez 195 void
232 :     idct_int32(short *const block)
233 : Isibaar 3 {
234 :    
235 : edgomez 1382 /*
236 :     * idct_int32_init() must be called before the first call to this
237 :     * function!
238 :     */
239 : Isibaar 3
240 :    
241 : edgomez 1382 #if 0
242 :     int i;
243 :     long i;
244 : Isibaar 3
245 : edgomez 1382 for (i=0; i<8; i++)
246 :     idctrow(block+8*i);
247 : Isibaar 3
248 : edgomez 1382 for (i=0; i<8; i++)
249 :     idctcol(block+i);
250 :     #endif
251 :    
252 : edgomez 1537 short *blk;
253 :     long i;
254 :     long X0, X1, X2, X3, X4, X5, X6, X7, X8;
255 : Isibaar 3
256 :    
257 : edgomez 1382 for (i = 0; i < 8; i++) /* idct rows */
258 : Isibaar 3 {
259 : edgomez 195 blk = block + (i << 3);
260 :     if (!
261 :     ((X1 = blk[4] << 11) | (X2 = blk[6]) | (X3 = blk[2]) | (X4 =
262 :     blk[1]) |
263 :     (X5 = blk[7]) | (X6 = blk[5]) | (X7 = blk[3]))) {
264 :     blk[0] = blk[1] = blk[2] = blk[3] = blk[4] = blk[5] = blk[6] =
265 :     blk[7] = blk[0] << 3;
266 :     continue;
267 :     }
268 : Isibaar 3
269 : edgomez 1382 X0 = (blk[0] << 11) + 128; /* for proper rounding in the fourth stage */
270 : Isibaar 3
271 : edgomez 1382 /* first stage */
272 : edgomez 195 X8 = W7 * (X4 + X5);
273 :     X4 = X8 + (W1 - W7) * X4;
274 :     X5 = X8 - (W1 + W7) * X5;
275 :     X8 = W3 * (X6 + X7);
276 :     X6 = X8 - (W3 - W5) * X6;
277 :     X7 = X8 - (W3 + W5) * X7;
278 : Isibaar 3
279 : edgomez 1382 /* second stage */
280 : edgomez 195 X8 = X0 + X1;
281 :     X0 -= X1;
282 :     X1 = W6 * (X3 + X2);
283 :     X2 = X1 - (W2 + W6) * X2;
284 :     X3 = X1 + (W2 - W6) * X3;
285 :     X1 = X4 + X6;
286 :     X4 -= X6;
287 :     X6 = X5 + X7;
288 :     X5 -= X7;
289 : Isibaar 3
290 : edgomez 1382 /* third stage */
291 : edgomez 195 X7 = X8 + X3;
292 :     X8 -= X3;
293 :     X3 = X0 + X2;
294 :     X0 -= X2;
295 :     X2 = (181 * (X4 + X5) + 128) >> 8;
296 :     X4 = (181 * (X4 - X5) + 128) >> 8;
297 : Isibaar 3
298 : edgomez 1382 /* fourth stage */
299 : Isibaar 3
300 : edgomez 195 blk[0] = (short) ((X7 + X1) >> 8);
301 :     blk[1] = (short) ((X3 + X2) >> 8);
302 :     blk[2] = (short) ((X0 + X4) >> 8);
303 :     blk[3] = (short) ((X8 + X6) >> 8);
304 :     blk[4] = (short) ((X8 - X6) >> 8);
305 :     blk[5] = (short) ((X0 - X4) >> 8);
306 :     blk[6] = (short) ((X3 - X2) >> 8);
307 :     blk[7] = (short) ((X7 - X1) >> 8);
308 :    
309 : edgomez 1382 } /* end for ( i = 0; i < 8; ++i ) IDCT-rows */
310 : edgomez 195
311 :    
312 :    
313 : edgomez 1382 for (i = 0; i < 8; i++) /* idct columns */
314 : Isibaar 3 {
315 : edgomez 195 blk = block + i;
316 : edgomez 1382 /* shortcut */
317 : edgomez 195 if (!
318 :     ((X1 = (blk[8 * 4] << 8)) | (X2 = blk[8 * 6]) | (X3 =
319 :     blk[8 *
320 :     2]) | (X4 =
321 :     blk[8 *
322 :     1])
323 :     | (X5 = blk[8 * 7]) | (X6 = blk[8 * 5]) | (X7 = blk[8 * 3]))) {
324 :     blk[8 * 0] = blk[8 * 1] = blk[8 * 2] = blk[8 * 3] = blk[8 * 4] =
325 :     blk[8 * 5] = blk[8 * 6] = blk[8 * 7] =
326 :     iclp[(blk[8 * 0] + 32) >> 6];
327 :     continue;
328 :     }
329 :    
330 :     X0 = (blk[8 * 0] << 8) + 8192;
331 :    
332 : edgomez 1382 /* first stage */
333 : edgomez 195 X8 = W7 * (X4 + X5) + 4;
334 :     X4 = (X8 + (W1 - W7) * X4) >> 3;
335 :     X5 = (X8 - (W1 + W7) * X5) >> 3;
336 :     X8 = W3 * (X6 + X7) + 4;
337 :     X6 = (X8 - (W3 - W5) * X6) >> 3;
338 :     X7 = (X8 - (W3 + W5) * X7) >> 3;
339 :    
340 : edgomez 1382 /* second stage */
341 : edgomez 195 X8 = X0 + X1;
342 :     X0 -= X1;
343 :     X1 = W6 * (X3 + X2) + 4;
344 :     X2 = (X1 - (W2 + W6) * X2) >> 3;
345 :     X3 = (X1 + (W2 - W6) * X3) >> 3;
346 :     X1 = X4 + X6;
347 :     X4 -= X6;
348 :     X6 = X5 + X7;
349 :     X5 -= X7;
350 :    
351 : edgomez 1382 /* third stage */
352 : edgomez 195 X7 = X8 + X3;
353 :     X8 -= X3;
354 :     X3 = X0 + X2;
355 :     X0 -= X2;
356 :     X2 = (181 * (X4 + X5) + 128) >> 8;
357 :     X4 = (181 * (X4 - X5) + 128) >> 8;
358 :    
359 : edgomez 1382 /* fourth stage */
360 : edgomez 195 blk[8 * 0] = iclp[(X7 + X1) >> 14];
361 :     blk[8 * 1] = iclp[(X3 + X2) >> 14];
362 :     blk[8 * 2] = iclp[(X0 + X4) >> 14];
363 :     blk[8 * 3] = iclp[(X8 + X6) >> 14];
364 :     blk[8 * 4] = iclp[(X8 - X6) >> 14];
365 :     blk[8 * 5] = iclp[(X0 - X4) >> 14];
366 :     blk[8 * 6] = iclp[(X3 - X2) >> 14];
367 :     blk[8 * 7] = iclp[(X7 - X1) >> 14];
368 : Isibaar 3 }
369 :    
370 : edgomez 1382 } /* end function idct_int32(block) */
371 : Isibaar 3
372 :    
373 : edgomez 195 void
374 : suxen_drol 1653 idct_int32_init(void)
375 : Isibaar 3 {
376 : edgomez 195 int i;
377 : Isibaar 3
378 : edgomez 195 iclp = iclip + 512;
379 :     for (i = -512; i < 512; i++)
380 :     iclp[i] = (i < -256) ? -256 : ((i > 255) ? 255 : i);
381 : Isibaar 3 }

No admin address has been configured
ViewVC Help
Powered by ViewVC 1.0.4