libm/src/s_exp2f.c

2fe8fb19SBen Gras/*-
2fe8fb19SBen Gras * Copyright (c) 2005 David Schultz <das@FreeBSD.ORG>
2fe8fb19SBen Gras * All rights reserved.
2fe8fb19SBen Gras *
2fe8fb19SBen Gras * Redistribution and use in source and binary forms, with or without
2fe8fb19SBen Gras * modification, are permitted provided that the following conditions
2fe8fb19SBen Gras * are met:
2fe8fb19SBen Gras * 1. Redistributions of source code must retain the above copyright
2fe8fb19SBen Gras *    notice, this list of conditions and the following disclaimer.
2fe8fb19SBen Gras * 2. Redistributions in binary form must reproduce the above copyright
2fe8fb19SBen Gras *    notice, this list of conditions and the following disclaimer in the
2fe8fb19SBen Gras *    documentation and/or other materials provided with the distribution.
2fe8fb19SBen Gras *
2fe8fb19SBen Gras * THIS SOFTWARE IS PROVIDED BY THE AUTHOR AND CONTRIBUTORS ``AS IS'' AND
2fe8fb19SBen Gras * ANY EXPRESS OR IMPLIED WARRANTIES, INCLUDING, BUT NOT LIMITED TO, THE
2fe8fb19SBen Gras * IMPLIED WARRANTIES OF MERCHANTABILITY AND FITNESS FOR A PARTICULAR PURPOSE
2fe8fb19SBen Gras * ARE DISCLAIMED.  IN NO EVENT SHALL THE AUTHOR OR CONTRIBUTORS BE LIABLE
2fe8fb19SBen Gras * FOR ANY DIRECT, INDIRECT, INCIDENTAL, SPECIAL, EXEMPLARY, OR CONSEQUENTIAL
2fe8fb19SBen Gras * DAMAGES (INCLUDING, BUT NOT LIMITED TO, PROCUREMENT OF SUBSTITUTE GOODS
2fe8fb19SBen Gras * OR SERVICES; LOSS OF USE, DATA, OR PROFITS; OR BUSINESS INTERRUPTION)
2fe8fb19SBen Gras * HOWEVER CAUSED AND ON ANY THEORY OF LIABILITY, WHETHER IN CONTRACT, STRICT
2fe8fb19SBen Gras * LIABILITY, OR TORT (INCLUDING NEGLIGENCE OR OTHERWISE) ARISING IN ANY WAY
2fe8fb19SBen Gras * OUT OF THE USE OF THIS SOFTWARE, EVEN IF ADVISED OF THE POSSIBILITY OF
2fe8fb19SBen Gras * SUCH DAMAGE.
2fe8fb19SBen Gras */
2fe8fb19SBen Gras
2fe8fb19SBen Gras#include <sys/cdefs.h>
*0a6a1f1dSLionel Sambuc__RCSID("$NetBSD: s_exp2f.c,v 1.2 2014/03/16 22:30:43 dsl Exp $");
2fe8fb19SBen Gras#ifdef __FBSDID
2fe8fb19SBen Gras__FBSDID("$FreeBSD: src/lib/msun/src/s_exp2f.c,v 1.9 2008/02/22 02:27:34 das Exp $");
2fe8fb19SBen Gras#endif
2fe8fb19SBen Gras
2fe8fb19SBen Gras#include <float.h>
2fe8fb19SBen Gras
2fe8fb19SBen Gras#include "math.h"
2fe8fb19SBen Gras#include "math_private.h"
2fe8fb19SBen Gras
2fe8fb19SBen Gras#define	TBLBITS	4
2fe8fb19SBen Gras#define	TBLSIZE	(1 << TBLBITS)
2fe8fb19SBen Gras
2fe8fb19SBen Grasstatic const float
2fe8fb19SBen Gras    redux   = 0x1.8p23f / TBLSIZE,
2fe8fb19SBen Gras    P1	    = 0x1.62e430p-1f,
2fe8fb19SBen Gras    P2	    = 0x1.ebfbe0p-3f,
2fe8fb19SBen Gras    P3	    = 0x1.c6b348p-5f,
2fe8fb19SBen Gras    P4	    = 0x1.3b2c9cp-7f;
2fe8fb19SBen Gras
*0a6a1f1dSLionel Sambuc/*
*0a6a1f1dSLionel Sambuc * For out of range values we need to generate the appropriate
*0a6a1f1dSLionel Sambuc * underflow or overflow trap as well as generating infinity or zero.
*0a6a1f1dSLionel Sambuc * This means we have to get the fpu to execute an instruction that
*0a6a1f1dSLionel Sambuc * will generate the trap (and not have the compiler optimise it away).
*0a6a1f1dSLionel Sambuc * This is normally done by calculating 'huge * huge' or 'tiny * tiny'.
*0a6a1f1dSLionel Sambuc *
*0a6a1f1dSLionel Sambuc * i386 is particularly problematic.
*0a6a1f1dSLionel Sambuc * The 'float' result is returned on the x87 stack, so is 'long double'.
*0a6a1f1dSLionel Sambuc * If we just multiply two 'float' values the caller will see 0x1p+/-200
*0a6a1f1dSLionel Sambuc * (not 0 or infinity).
*0a6a1f1dSLionel Sambuc * If we use 'double' the compiler does a store-load which will convert the
*0a6a1f1dSLionel Sambuc * value and generate the required exception.
*0a6a1f1dSLionel Sambuc */
*0a6a1f1dSLionel Sambuc#ifdef __i386__
*0a6a1f1dSLionel Sambucstatic volatile double overflow = 0x1p+1000;
*0a6a1f1dSLionel Sambucstatic volatile double underflow = 0x1p-1000;
*0a6a1f1dSLionel Sambuc#else
*0a6a1f1dSLionel Sambucstatic volatile float huge = 0x1p+100;
*0a6a1f1dSLionel Sambucstatic volatile float tiny = 0x1p-100;
*0a6a1f1dSLionel Sambuc#define overflow (huge * huge)
*0a6a1f1dSLionel Sambuc#define underflow (tiny * tiny)
*0a6a1f1dSLionel Sambuc#endif
2fe8fb19SBen Gras
2fe8fb19SBen Grasstatic const double exp2ft[TBLSIZE] = {
2fe8fb19SBen Gras	0x1.6a09e667f3bcdp-1,
2fe8fb19SBen Gras	0x1.7a11473eb0187p-1,
2fe8fb19SBen Gras	0x1.8ace5422aa0dbp-1,
2fe8fb19SBen Gras	0x1.9c49182a3f090p-1,
2fe8fb19SBen Gras	0x1.ae89f995ad3adp-1,
2fe8fb19SBen Gras	0x1.c199bdd85529cp-1,
2fe8fb19SBen Gras	0x1.d5818dcfba487p-1,
2fe8fb19SBen Gras	0x1.ea4afa2a490dap-1,
2fe8fb19SBen Gras	0x1.0000000000000p+0,
2fe8fb19SBen Gras	0x1.0b5586cf9890fp+0,
2fe8fb19SBen Gras	0x1.172b83c7d517bp+0,
2fe8fb19SBen Gras	0x1.2387a6e756238p+0,
2fe8fb19SBen Gras	0x1.306fe0a31b715p+0,
2fe8fb19SBen Gras	0x1.3dea64c123422p+0,
2fe8fb19SBen Gras	0x1.4bfdad5362a27p+0,
2fe8fb19SBen Gras	0x1.5ab07dd485429p+0,
2fe8fb19SBen Gras};
2fe8fb19SBen Gras
2fe8fb19SBen Gras/*
2fe8fb19SBen Gras * exp2f(x): compute the base 2 exponential of x
2fe8fb19SBen Gras *
2fe8fb19SBen Gras * Accuracy: Peak error < 0.501 ulp; location of peak: -0.030110927.
2fe8fb19SBen Gras *
2fe8fb19SBen Gras * Method: (equally-spaced tables)
2fe8fb19SBen Gras *
2fe8fb19SBen Gras *   Reduce x:
2fe8fb19SBen Gras *     x = 2**k + y, for integer k and |y| <= 1/2.
2fe8fb19SBen Gras *     Thus we have exp2f(x) = 2**k * exp2(y).
2fe8fb19SBen Gras *
2fe8fb19SBen Gras *   Reduce y:
2fe8fb19SBen Gras *     y = i/TBLSIZE + z for integer i near y * TBLSIZE.
2fe8fb19SBen Gras *     Thus we have exp2(y) = exp2(i/TBLSIZE) * exp2(z),
2fe8fb19SBen Gras *     with |z| <= 2**-(TBLSIZE+1).
2fe8fb19SBen Gras *
2fe8fb19SBen Gras *   We compute exp2(i/TBLSIZE) via table lookup and exp2(z) via a
2fe8fb19SBen Gras *   degree-4 minimax polynomial with maximum error under 1.4 * 2**-33.
2fe8fb19SBen Gras *   Using double precision for everything except the reduction makes
2fe8fb19SBen Gras *   roundoff error insignificant and simplifies the scaling step.
2fe8fb19SBen Gras *
2fe8fb19SBen Gras *   This method is due to Tang, but I do not use his suggested parameters:
2fe8fb19SBen Gras *
2fe8fb19SBen Gras *	Tang, P.  Table-driven Implementation of the Exponential Function
2fe8fb19SBen Gras *	in IEEE Floating-Point Arithmetic.  TOMS 15(2), 144-157 (1989).
2fe8fb19SBen Gras */
2fe8fb19SBen Grasfloat
2fe8fb19SBen Grasexp2f(float x)
2fe8fb19SBen Gras{
2fe8fb19SBen Gras	double tv, twopk, u, z;
2fe8fb19SBen Gras	float t;
2fe8fb19SBen Gras	uint32_t hx, ix, i0;
2fe8fb19SBen Gras	int32_t k;
2fe8fb19SBen Gras
2fe8fb19SBen Gras	/* Filter out exceptional cases. */
2fe8fb19SBen Gras	GET_FLOAT_WORD(hx, x);
2fe8fb19SBen Gras	ix = hx & 0x7fffffff;		/* high word of |x| */
2fe8fb19SBen Gras	if(ix >= 0x43000000) {			/* |x| >= 128 */
2fe8fb19SBen Gras		if(ix >= 0x7f800000) {
2fe8fb19SBen Gras			if ((ix & 0x7fffff) != 0 || (hx & 0x80000000) == 0)
2fe8fb19SBen Gras				return (x + x);	/* x is NaN or +Inf */
2fe8fb19SBen Gras			else
2fe8fb19SBen Gras				return (0.0);	/* x is -Inf */
2fe8fb19SBen Gras		}
2fe8fb19SBen Gras		if(x >= 0x1.0p7f)
*0a6a1f1dSLionel Sambuc			return overflow;	/* +infinity with overflow */
2fe8fb19SBen Gras		if(x <= -0x1.2cp7f)
*0a6a1f1dSLionel Sambuc			return underflow;	/* zero with underflow */
2fe8fb19SBen Gras	} else if (ix <= 0x33000000) {		/* |x| <= 0x1p-25 */
2fe8fb19SBen Gras		return (1.0f + x);
2fe8fb19SBen Gras	}
2fe8fb19SBen Gras
2fe8fb19SBen Gras	/* Reduce x, computing z, i0, and k. */
2fe8fb19SBen Gras	STRICT_ASSIGN(float, t, x + redux);
2fe8fb19SBen Gras	GET_FLOAT_WORD(i0, t);
2fe8fb19SBen Gras	i0 += TBLSIZE / 2;
2fe8fb19SBen Gras	k = (i0 >> TBLBITS) << 20;
2fe8fb19SBen Gras	i0 &= TBLSIZE - 1;
2fe8fb19SBen Gras	t -= redux;
2fe8fb19SBen Gras	z = x - t;
2fe8fb19SBen Gras	INSERT_WORDS(twopk, 0x3ff00000 + k, 0);
2fe8fb19SBen Gras
2fe8fb19SBen Gras	/* Compute r = exp2(y) = exp2ft[i0] * p(z). */
2fe8fb19SBen Gras	tv = exp2ft[i0];
2fe8fb19SBen Gras	u = tv * z;
2fe8fb19SBen Gras	tv = tv + u * (P1 + z * P2) + u * (z * z) * (P3 + z * P4);
2fe8fb19SBen Gras
2fe8fb19SBen Gras	/* Scale by 2**(k>>20). */
2fe8fb19SBen Gras	return (tv * twopk);
2fe8fb19SBen Gras}