mirror of https://github.com/PCSX2/pcsx2.git
146 lines
4.7 KiB
C++
146 lines
4.7 KiB
C++
/*
|
|
* Copyright (C) 2007-2009 Gabest
|
|
* http://www.gabest.org
|
|
*
|
|
* This Program is free software; you can redistribute it and/or modify
|
|
* it under the terms of the GNU General Public License as published by
|
|
* the Free Software Foundation; either version 2, or (at your option)
|
|
* any later version.
|
|
*
|
|
* This Program is distributed in the hope that it will be useful,
|
|
* but WITHOUT ANY WARRANTY; without even the implied warranty of
|
|
* MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the
|
|
* GNU General Public License for more details.
|
|
*
|
|
* You should have received a copy of the GNU General Public License
|
|
* along with GNU Make; see the file COPYING. If not, write to
|
|
* the Free Software Foundation, Inc., 51 Franklin Street, Fifth Floor, Boston, MA 02110-1301, USA USA.
|
|
* http://www.gnu.org/copyleft/gpl.html
|
|
*
|
|
*/
|
|
|
|
#pragma once
|
|
|
|
#include "GSScanlineEnvironment.h"
|
|
#include "GSFunctionMap.h"
|
|
|
|
using namespace Xbyak;
|
|
|
|
class GSDrawScanlineCodeGenerator : public GSCodeGenerator
|
|
{
|
|
void operator = (const GSDrawScanlineCodeGenerator&);
|
|
|
|
GSScanlineSelector m_sel;
|
|
GSScanlineLocalData& m_local;
|
|
|
|
void Generate();
|
|
|
|
#if _M_SSE >= 0x501
|
|
|
|
void Init();
|
|
void Step();
|
|
void TestZ(const Ymm& temp1, const Ymm& temp2);
|
|
void SampleTexture();
|
|
void Wrap(const Ymm& uv0);
|
|
void Wrap(const Ymm& uv0, const Ymm& uv1);
|
|
void SampleTextureLOD();
|
|
void WrapLOD(const Ymm& uv0);
|
|
void WrapLOD(const Ymm& uv0, const Ymm& uv1);
|
|
void AlphaTFX();
|
|
void ReadMask();
|
|
void TestAlpha();
|
|
void ColorTFX();
|
|
void Fog();
|
|
void ReadFrame();
|
|
void TestDestAlpha();
|
|
void WriteMask();
|
|
void WriteZBuf();
|
|
void AlphaBlend();
|
|
void WriteFrame();
|
|
|
|
#if defined(_M_AMD64) || defined(_WIN64)
|
|
void ReadPixel(const Ymm& dst, const Ymm& temp, const Reg64& addr);
|
|
void WritePixel(const Ymm& src, const Ymm& temp, const Reg64& addr, const Reg32& mask, bool fast, int psm, int fz);
|
|
void WritePixel(const Xmm& src, const Reg64& addr, uint8 i, uint8 j, int psm);
|
|
#else
|
|
void ReadPixel(const Ymm& dst, const Ymm& temp, const Reg32& addr);
|
|
void WritePixel(const Ymm& src, const Ymm& temp, const Reg32& addr, const Reg32& mask, bool fast, int psm, int fz);
|
|
void WritePixel(const Xmm& src, const Reg32& addr, uint8 i, uint8 j, int psm);
|
|
#endif
|
|
|
|
void ReadTexel(int pixels, int mip_offset = 0);
|
|
void ReadTexel(const Ymm& dst, const Ymm& addr, uint8 i);
|
|
|
|
void modulate16(const Ymm& a, const Operand& f, int shift);
|
|
void lerp16(const Ymm& a, const Ymm& b, const Ymm& f, int shift);
|
|
void lerp16_4(const Ymm& a, const Ymm& b, const Ymm& f);
|
|
void mix16(const Ymm& a, const Ymm& b, const Ymm& temp);
|
|
void clamp16(const Ymm& a, const Ymm& temp);
|
|
void alltrue();
|
|
void blend(const Ymm& a, const Ymm& b, const Ymm& mask);
|
|
void blendr(const Ymm& b, const Ymm& a, const Ymm& mask);
|
|
void blend8(const Ymm& a, const Ymm& b);
|
|
void blend8r(const Ymm& b, const Ymm& a);
|
|
|
|
#else
|
|
|
|
void Init();
|
|
void Step();
|
|
void TestZ(const Xmm& temp1, const Xmm& temp2);
|
|
void SampleTexture();
|
|
void Wrap(const Xmm& uv0);
|
|
void Wrap(const Xmm& uv0, const Xmm& uv1);
|
|
void SampleTextureLOD();
|
|
void WrapLOD(const Xmm& uv0);
|
|
void WrapLOD(const Xmm& uv0, const Xmm& uv1);
|
|
void AlphaTFX();
|
|
void ReadMask();
|
|
void TestAlpha();
|
|
void ColorTFX();
|
|
void Fog();
|
|
void ReadFrame();
|
|
void TestDestAlpha();
|
|
void WriteMask();
|
|
void WriteZBuf();
|
|
void AlphaBlend();
|
|
void WriteFrame();
|
|
|
|
#if defined(_M_AMD64) || defined(_WIN64)
|
|
void ReadPixel(const Xmm& dst, const Reg64& addr);
|
|
void WritePixel(const Xmm& src, const Reg64& addr, const Reg8& mask, bool fast, int psm, int fz);
|
|
void WritePixel(const Xmm& src, const Reg64& addr, uint8 i, int psm);
|
|
#else
|
|
void ReadPixel(const Xmm& dst, const Reg32& addr);
|
|
void WritePixel(const Xmm& src, const Reg32& addr, const Reg8& mask, bool fast, int psm, int fz);
|
|
void WritePixel(const Xmm& src, const Reg32& addr, uint8 i, int psm);
|
|
#endif
|
|
|
|
void ReadTexel(int pixels, int mip_offset = 0);
|
|
void ReadTexel(const Xmm& dst, const Xmm& addr, uint8 i);
|
|
|
|
void modulate16(const Xmm& a, const Operand& f, int shift);
|
|
void lerp16(const Xmm& a, const Xmm& b, const Xmm& f, int shift);
|
|
void lerp16_4(const Xmm& a, const Xmm& b, const Xmm& f);
|
|
void mix16(const Xmm& a, const Xmm& b, const Xmm& temp);
|
|
void clamp16(const Xmm& a, const Xmm& temp);
|
|
void alltrue();
|
|
void blend(const Xmm& a, const Xmm& b, const Xmm& mask);
|
|
void blendr(const Xmm& b, const Xmm& a, const Xmm& mask);
|
|
void blend8(const Xmm& a, const Xmm& b);
|
|
void blend8r(const Xmm& b, const Xmm& a);
|
|
|
|
#endif
|
|
|
|
public:
|
|
GSDrawScanlineCodeGenerator(void* param, uint64 key, void* code, size_t maxsize);
|
|
|
|
#if _M_SSE >= 0x501
|
|
alignas(8) static const uint8 m_test[16][8];
|
|
static const GSVector8 m_log2_coef[4];
|
|
#else
|
|
static const GSVector4i m_test[8];
|
|
static const GSVector4 m_log2_coef[4];
|
|
#endif
|
|
|
|
};
|