mirror of
https://github.com/FFmpeg/FFmpeg.git
synced 2024-12-23 12:43:46 +02:00
simplify, speedup and reduce needed headroom by 2 bits in the 3rd
vertical lifting step Originally committed as revision 10162 to svn://svn.ffmpeg.org/ffmpeg/trunk
This commit is contained in:
parent
dd30437bbe
commit
4bf1790421
@ -590,18 +590,17 @@ void ff_snow_vertical_compose97i_mmx(DWTELEM *b0, DWTELEM *b1, DWTELEM *b2, DWTE
|
||||
snow_vertical_compose_mmx_sub("mm1","mm3","mm5","mm7","mm0","mm2","mm4","mm6")
|
||||
snow_vertical_compose_mmx_store(REG_S,"mm0","mm2","mm4","mm6")
|
||||
"mov %2, %%"REG_a" \n\t"
|
||||
snow_vertical_compose_mmx_load(REG_c,"mm1","mm3","mm5","mm7")
|
||||
snow_vertical_compose_mmx_add(REG_a,"mm0","mm2","mm4","mm6")
|
||||
snow_vertical_compose_mmx_sll("2","mm1","mm3","mm5","mm7")
|
||||
snow_vertical_compose_mmx_r2r_add("mm1","mm3","mm5","mm7","mm0","mm2","mm4","mm6")
|
||||
snow_vertical_compose_mmx_sra("2","mm0","mm2","mm4","mm6")
|
||||
snow_vertical_compose_mmx_add(REG_c,"mm0","mm2","mm4","mm6")
|
||||
|
||||
"pcmpeqd %%mm1, %%mm1 \n\t"
|
||||
"pslld $31, %%mm1 \n\t"
|
||||
"psrld $28, %%mm1 \n\t"
|
||||
"psrld $30, %%mm1 \n\t"
|
||||
"mov %1, %%"REG_S" \n\t"
|
||||
|
||||
snow_vertical_compose_mmx_r2r_add("mm1","mm1","mm1","mm1","mm0","mm2","mm4","mm6")
|
||||
snow_vertical_compose_mmx_sra("4","mm0","mm2","mm4","mm6")
|
||||
snow_vertical_compose_mmx_sra("2","mm0","mm2","mm4","mm6")
|
||||
snow_vertical_compose_mmx_add(REG_c,"mm0","mm2","mm4","mm6")
|
||||
snow_vertical_compose_mmx_store(REG_c,"mm0","mm2","mm4","mm6")
|
||||
snow_vertical_compose_mmx_add(REG_S,"mm0","mm2","mm4","mm6")
|
||||
|
Loading…
Reference in New Issue
Block a user