14a2e7687Sjeremylt // Copyright (c) 2017-2018, Lawrence Livermore National Security, LLC. 24a2e7687Sjeremylt // Produced at the Lawrence Livermore National Laboratory. LLNL-CODE-734707. 34a2e7687Sjeremylt // All Rights reserved. See files LICENSE and NOTICE for details. 44a2e7687Sjeremylt // 54a2e7687Sjeremylt // This file is part of CEED, a collection of benchmarks, miniapps, software 64a2e7687Sjeremylt // libraries and APIs for efficient high-order finite element and spectral 74a2e7687Sjeremylt // element discretizations for exascale applications. For more information and 84a2e7687Sjeremylt // source code availability see http://github.com/ceed. 94a2e7687Sjeremylt // 104a2e7687Sjeremylt // The CEED research is supported by the Exascale Computing Project 17-SC-20-SC, 114a2e7687Sjeremylt // a collaborative effort of two U.S. Department of Energy organizations (Office 124a2e7687Sjeremylt // of Science and the National Nuclear Security Administration) responsible for 134a2e7687Sjeremylt // the planning and preparation of a capable exascale ecosystem, including 144a2e7687Sjeremylt // software, applications, hardware, advanced system engineering and early 154a2e7687Sjeremylt // testbed platforms, in support of the nation's exascale computing imperative. 164a2e7687Sjeremylt 174a2e7687Sjeremylt #include <string.h> 184a2e7687Sjeremylt #include "ceed-blocked.h" 194a2e7687Sjeremylt #include "../ref/ceed-ref.h" 204a2e7687Sjeremylt 214a2e7687Sjeremylt static int CeedOperatorDestroy_Blocked(CeedOperator op) { 224a2e7687Sjeremylt int ierr; 234ce2993fSjeremylt CeedOperator_Blocked *impl; 244ce2993fSjeremylt ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr); 254a2e7687Sjeremylt 264a2e7687Sjeremylt for (CeedInt i=0; i<impl->numein+impl->numeout; i++) { 274a2e7687Sjeremylt ierr = CeedElemRestrictionDestroy(&impl->blkrestr[i]); CeedChk(ierr); 284a2e7687Sjeremylt ierr = CeedVectorDestroy(&impl->evecs[i]); CeedChk(ierr); 294a2e7687Sjeremylt } 304a2e7687Sjeremylt ierr = CeedFree(&impl->blkrestr); CeedChk(ierr); 314a2e7687Sjeremylt ierr = CeedFree(&impl->evecs); CeedChk(ierr); 324a2e7687Sjeremylt ierr = CeedFree(&impl->edata); CeedChk(ierr); 3316c359e6Sjeremylt ierr = CeedFree(&impl->inputstate); CeedChk(ierr); 344a2e7687Sjeremylt 35aedaa0e5Sjeremylt for (CeedInt i=0; i<impl->numein; i++) { 3691703d3fSjeremylt ierr = CeedVectorDestroy(&impl->evecsin[i]); CeedChk(ierr); 37aedaa0e5Sjeremylt ierr = CeedVectorDestroy(&impl->qvecsin[i]); CeedChk(ierr); 384a2e7687Sjeremylt } 3991703d3fSjeremylt ierr = CeedFree(&impl->evecsin); CeedChk(ierr); 40aedaa0e5Sjeremylt ierr = CeedFree(&impl->qvecsin); CeedChk(ierr); 414a2e7687Sjeremylt 42aedaa0e5Sjeremylt for (CeedInt i=0; i<impl->numeout; i++) { 4391703d3fSjeremylt ierr = CeedVectorDestroy(&impl->evecsout[i]); CeedChk(ierr); 44aedaa0e5Sjeremylt ierr = CeedVectorDestroy(&impl->qvecsout[i]); CeedChk(ierr); 45aedaa0e5Sjeremylt } 4691703d3fSjeremylt ierr = CeedFree(&impl->evecsout); CeedChk(ierr); 47aedaa0e5Sjeremylt ierr = CeedFree(&impl->qvecsout); CeedChk(ierr); 484a2e7687Sjeremylt 49fe2413ffSjeremylt ierr = CeedFree(&impl); CeedChk(ierr); 504a2e7687Sjeremylt return 0; 514a2e7687Sjeremylt } 524a2e7687Sjeremylt 534a2e7687Sjeremylt /* 544a2e7687Sjeremylt Setup infields or outfields 554a2e7687Sjeremylt */ 56fe2413ffSjeremylt static int CeedOperatorSetupFields_Blocked(CeedQFunction qf, CeedOperator op, 57fe2413ffSjeremylt bool inOrOut, 584a2e7687Sjeremylt CeedElemRestriction *blkrestr, 5991703d3fSjeremylt CeedVector *fullevecs, CeedVector *evecs, 60aedaa0e5Sjeremylt CeedVector *qvecs, CeedInt starte, 614a2e7687Sjeremylt CeedInt numfields, CeedInt Q) { 624d1cd9fcSJeremy L Thompson CeedInt dim, ierr, ncomp, P; 63aedaa0e5Sjeremylt Ceed ceed; 64aedaa0e5Sjeremylt ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr); 65d1bcdac9Sjeremylt CeedBasis basis; 66d1bcdac9Sjeremylt CeedElemRestriction r; 67aedaa0e5Sjeremylt CeedOperatorField *opfields; 68aedaa0e5Sjeremylt CeedQFunctionField *qffields; 69fe2413ffSjeremylt if (inOrOut) { 70aedaa0e5Sjeremylt ierr = CeedOperatorGetFields(op, NULL, &opfields); 71fe2413ffSjeremylt CeedChk(ierr); 72aedaa0e5Sjeremylt ierr = CeedQFunctionGetFields(qf, NULL, &qffields); 73fe2413ffSjeremylt CeedChk(ierr); 74fe2413ffSjeremylt } else { 75aedaa0e5Sjeremylt ierr = CeedOperatorGetFields(op, &opfields, NULL); 76fe2413ffSjeremylt CeedChk(ierr); 77aedaa0e5Sjeremylt ierr = CeedQFunctionGetFields(qf, &qffields, NULL); 78fe2413ffSjeremylt CeedChk(ierr); 79fe2413ffSjeremylt } 804a2e7687Sjeremylt const CeedInt blksize = 8; 814a2e7687Sjeremylt 824a2e7687Sjeremylt // Loop over fields 834a2e7687Sjeremylt for (CeedInt i=0; i<numfields; i++) { 84d1bcdac9Sjeremylt CeedEvalMode emode; 85aedaa0e5Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qffields[i], &emode); CeedChk(ierr); 864a2e7687Sjeremylt 874a2e7687Sjeremylt if (emode != CEED_EVAL_WEIGHT) { 88aedaa0e5Sjeremylt ierr = CeedOperatorFieldGetElemRestriction(opfields[i], &r); 89d1bcdac9Sjeremylt CeedChk(ierr); 90fe2413ffSjeremylt CeedElemRestriction_Ref *data; 91*de686571SJeremy L Thompson ierr = CeedElemRestrictionGetData(r, (void *)&data); CeedChk(ierr); 924ce2993fSjeremylt Ceed ceed; 934ce2993fSjeremylt ierr = CeedElemRestrictionGetCeed(r, &ceed); CeedChk(ierr); 944d1cd9fcSJeremy L Thompson CeedInt nelem, elemsize, ndof; 954ce2993fSjeremylt ierr = CeedElemRestrictionGetNumElements(r, &nelem); CeedChk(ierr); 964ce2993fSjeremylt ierr = CeedElemRestrictionGetElementSize(r, &elemsize); CeedChk(ierr); 974ce2993fSjeremylt ierr = CeedElemRestrictionGetNumDoF(r, &ndof); CeedChk(ierr); 984ce2993fSjeremylt ierr = CeedElemRestrictionGetNumComponents(r, &ncomp); CeedChk(ierr); 994ce2993fSjeremylt ierr = CeedElemRestrictionCreateBlocked(ceed, nelem, elemsize, 1004ce2993fSjeremylt blksize, ndof, ncomp, 1014a2e7687Sjeremylt CEED_MEM_HOST, CEED_COPY_VALUES, 102aedaa0e5Sjeremylt data->indices, &blkrestr[i+starte]); 1034a2e7687Sjeremylt CeedChk(ierr); 104aedaa0e5Sjeremylt ierr = CeedElemRestrictionCreateVector(blkrestr[i+starte], NULL, 10591703d3fSjeremylt &fullevecs[i+starte]); 1064a2e7687Sjeremylt CeedChk(ierr); 1074a2e7687Sjeremylt } 1084a2e7687Sjeremylt 1094a2e7687Sjeremylt switch(emode) { 1104a2e7687Sjeremylt case CEED_EVAL_NONE: 111aedaa0e5Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp); 112d1bcdac9Sjeremylt CeedChk(ierr); 11391703d3fSjeremylt ierr = CeedVectorCreate(ceed, Q*ncomp*blksize, &qvecs[i]); CeedChk(ierr); 114aedaa0e5Sjeremylt break; 115aedaa0e5Sjeremylt case CEED_EVAL_INTERP: 116aedaa0e5Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp); 117aedaa0e5Sjeremylt CeedChk(ierr); 1184d1cd9fcSJeremy L Thompson ierr = CeedElemRestrictionGetElementSize(r, &P); 1194d1cd9fcSJeremy L Thompson CeedChk(ierr); 1204d1cd9fcSJeremy L Thompson ierr = CeedVectorCreate(ceed, P*ncomp*blksize, &evecs[i]); CeedChk(ierr); 121aedaa0e5Sjeremylt ierr = CeedVectorCreate(ceed, Q*ncomp*blksize, &qvecs[i]); CeedChk(ierr); 1224a2e7687Sjeremylt break; 1234a2e7687Sjeremylt case CEED_EVAL_GRAD: 124aedaa0e5Sjeremylt ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr); 125aedaa0e5Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qffields[i], &ncomp); 126*de686571SJeremy L Thompson CeedChk(ierr); 127d1bcdac9Sjeremylt ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr); 1284d1cd9fcSJeremy L Thompson ierr = CeedElemRestrictionGetElementSize(r, &P); 1294d1cd9fcSJeremy L Thompson CeedChk(ierr); 1304d1cd9fcSJeremy L Thompson ierr = CeedVectorCreate(ceed, P*ncomp*blksize, &evecs[i]); CeedChk(ierr); 131aedaa0e5Sjeremylt ierr = CeedVectorCreate(ceed, Q*ncomp*dim*blksize, &qvecs[i]); CeedChk(ierr); 1324a2e7687Sjeremylt break; 1334a2e7687Sjeremylt case CEED_EVAL_WEIGHT: // Only on input fields 134aedaa0e5Sjeremylt ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr); 135aedaa0e5Sjeremylt ierr = CeedVectorCreate(ceed, Q*blksize, &qvecs[i]); CeedChk(ierr); 136d1bcdac9Sjeremylt ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE, 137aedaa0e5Sjeremylt CEED_EVAL_WEIGHT, NULL, qvecs[i]); CeedChk(ierr); 138aedaa0e5Sjeremylt 1394a2e7687Sjeremylt break; 1404a2e7687Sjeremylt case CEED_EVAL_DIV: 1414a2e7687Sjeremylt break; // Not implimented 1424a2e7687Sjeremylt case CEED_EVAL_CURL: 1434a2e7687Sjeremylt break; // Not implimented 1444a2e7687Sjeremylt } 1454a2e7687Sjeremylt } 1464a2e7687Sjeremylt return 0; 1474a2e7687Sjeremylt } 1484a2e7687Sjeremylt 1494a2e7687Sjeremylt /* 1504a2e7687Sjeremylt CeedOperator needs to connect all the named fields (be they active or passive) 1514a2e7687Sjeremylt to the named inputs and outputs of its CeedQFunction. 1524a2e7687Sjeremylt */ 1534a2e7687Sjeremylt static int CeedOperatorSetup_Blocked(CeedOperator op) { 1544a2e7687Sjeremylt int ierr; 1554ce2993fSjeremylt bool setupdone; 1564ce2993fSjeremylt ierr = CeedOperatorGetSetupStatus(op, &setupdone); CeedChk(ierr); 1574ce2993fSjeremylt if (setupdone) return 0; 158aedaa0e5Sjeremylt Ceed ceed; 159aedaa0e5Sjeremylt ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr); 1604ce2993fSjeremylt CeedOperator_Blocked *impl; 1614ce2993fSjeremylt ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr); 1624ce2993fSjeremylt CeedQFunction qf; 1634ce2993fSjeremylt ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr); 1644ce2993fSjeremylt CeedInt Q, numinputfields, numoutputfields; 1654ce2993fSjeremylt ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr); 1664a2e7687Sjeremylt ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields); 1674a2e7687Sjeremylt CeedChk(ierr); 168d1bcdac9Sjeremylt CeedOperatorField *opinputfields, *opoutputfields; 169d1bcdac9Sjeremylt ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields); 170d1bcdac9Sjeremylt CeedChk(ierr); 171d1bcdac9Sjeremylt CeedQFunctionField *qfinputfields, *qfoutputfields; 172d1bcdac9Sjeremylt ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields); 173d1bcdac9Sjeremylt CeedChk(ierr); 1744a2e7687Sjeremylt 1754a2e7687Sjeremylt // Allocate 176aedaa0e5Sjeremylt ierr = CeedCalloc(numinputfields + numoutputfields, &impl->blkrestr); 1774a2e7687Sjeremylt CeedChk(ierr); 178aedaa0e5Sjeremylt ierr = CeedCalloc(numinputfields + numoutputfields, &impl->evecs); 1794a2e7687Sjeremylt CeedChk(ierr); 180aedaa0e5Sjeremylt ierr = CeedCalloc(numinputfields + numoutputfields, &impl->edata); 1814a2e7687Sjeremylt CeedChk(ierr); 1824a2e7687Sjeremylt 18316c359e6Sjeremylt ierr = CeedCalloc(16, &impl->inputstate); CeedChk(ierr); 18491703d3fSjeremylt ierr = CeedCalloc(16, &impl->evecsin); CeedChk(ierr); 18591703d3fSjeremylt ierr = CeedCalloc(16, &impl->evecsout); CeedChk(ierr); 186aedaa0e5Sjeremylt ierr = CeedCalloc(16, &impl->qvecsin); CeedChk(ierr); 187aedaa0e5Sjeremylt ierr = CeedCalloc(16, &impl->qvecsout); CeedChk(ierr); 1884a2e7687Sjeremylt 189aedaa0e5Sjeremylt impl->numein = numinputfields; impl->numeout = numoutputfields; 190aedaa0e5Sjeremylt 1914a2e7687Sjeremylt // Set up infield and outfield pointer arrays 1924a2e7687Sjeremylt // Infields 193aedaa0e5Sjeremylt ierr = CeedOperatorSetupFields_Blocked(qf, op, 0, impl->blkrestr, 19491703d3fSjeremylt impl->evecs, impl->evecsin, 19591703d3fSjeremylt impl->qvecsin, 0, 196aedaa0e5Sjeremylt numinputfields, Q); 1974a2e7687Sjeremylt CeedChk(ierr); 1984a2e7687Sjeremylt // Outfields 199aedaa0e5Sjeremylt ierr = CeedOperatorSetupFields_Blocked(qf, op, 1, impl->blkrestr, 20091703d3fSjeremylt impl->evecs, impl->evecsout, 20191703d3fSjeremylt impl->qvecsout, numinputfields, 20291703d3fSjeremylt numoutputfields, Q); 2034a2e7687Sjeremylt CeedChk(ierr); 204aedaa0e5Sjeremylt 2054ce2993fSjeremylt ierr = CeedOperatorSetSetupDone(op); CeedChk(ierr); 2064a2e7687Sjeremylt 2074a2e7687Sjeremylt return 0; 2084a2e7687Sjeremylt } 2094a2e7687Sjeremylt 2104a2e7687Sjeremylt static int CeedOperatorApply_Blocked(CeedOperator op, CeedVector invec, 2114a2e7687Sjeremylt CeedVector outvec, CeedRequest *request) { 2124a2e7687Sjeremylt int ierr; 2134ce2993fSjeremylt CeedOperator_Blocked *impl; 2144ce2993fSjeremylt ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr); 2154ce2993fSjeremylt const CeedInt blksize = 8; 216d1bcdac9Sjeremylt CeedInt Q, elemsize, numinputfields, numoutputfields, numelements, ncomp; 2174ce2993fSjeremylt ierr = CeedOperatorGetNumElements(op, &numelements); CeedChk(ierr); 2184ce2993fSjeremylt ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr); 2194ce2993fSjeremylt CeedInt nblks = (numelements/blksize) + !!(numelements%blksize); 2204ce2993fSjeremylt CeedQFunction qf; 2214ce2993fSjeremylt ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr); 2224ce2993fSjeremylt ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields); 2234ce2993fSjeremylt CeedChk(ierr); 2244dccadb6Sjeremylt CeedTransposeMode lmode; 225d1bcdac9Sjeremylt CeedOperatorField *opinputfields, *opoutputfields; 226d1bcdac9Sjeremylt ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields); 227d1bcdac9Sjeremylt CeedChk(ierr); 228d1bcdac9Sjeremylt CeedQFunctionField *qfinputfields, *qfoutputfields; 229d1bcdac9Sjeremylt ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields); 230d1bcdac9Sjeremylt CeedChk(ierr); 231d1bcdac9Sjeremylt CeedEvalMode emode; 232d1bcdac9Sjeremylt CeedVector vec; 233d1bcdac9Sjeremylt CeedBasis basis; 234d1bcdac9Sjeremylt CeedElemRestriction Erestrict; 23516c359e6Sjeremylt uint64_t state; 2364a2e7687Sjeremylt 2374a2e7687Sjeremylt // Setup 2384a2e7687Sjeremylt ierr = CeedOperatorSetup_Blocked(op); CeedChk(ierr); 2394a2e7687Sjeremylt 2404a2e7687Sjeremylt // Input Evecs and Restriction 2414a2e7687Sjeremylt for (CeedInt i=0; i<numinputfields; i++) { 242d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode); 243d1bcdac9Sjeremylt CeedChk(ierr); 2444a2e7687Sjeremylt if (emode == CEED_EVAL_WEIGHT) { // Skip 2454a2e7687Sjeremylt } else { 246d1bcdac9Sjeremylt // Get input vector 247d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr); 248d1bcdac9Sjeremylt if (vec == CEED_VECTOR_ACTIVE) 249d1bcdac9Sjeremylt vec = invec; 2504a2e7687Sjeremylt // Restrict 25116c359e6Sjeremylt ierr = CeedVectorGetState(vec, &state); CeedChk(ierr); 2528d713cf6Sjeremylt if (state != impl->inputstate[i] || vec == invec) { 2534dccadb6Sjeremylt ierr = CeedOperatorFieldGetLMode(opinputfields[i], &lmode); CeedChk(ierr); 2544a2e7687Sjeremylt ierr = CeedElemRestrictionApply(impl->blkrestr[i], CEED_NOTRANSPOSE, 255d1bcdac9Sjeremylt lmode, vec, impl->evecs[i], 2564a2e7687Sjeremylt request); CeedChk(ierr); CeedChk(ierr); 25716c359e6Sjeremylt impl->inputstate[i] = state; 25816c359e6Sjeremylt } 2594a2e7687Sjeremylt // Get evec 2604a2e7687Sjeremylt ierr = CeedVectorGetArrayRead(impl->evecs[i], CEED_MEM_HOST, 2614a2e7687Sjeremylt (const CeedScalar **) &impl->edata[i]); 2624a2e7687Sjeremylt CeedChk(ierr); 2634a2e7687Sjeremylt } 2644a2e7687Sjeremylt } 2654a2e7687Sjeremylt 2664a2e7687Sjeremylt // Output Evecs 2674a2e7687Sjeremylt for (CeedInt i=0; i<numoutputfields; i++) { 2684a2e7687Sjeremylt ierr = CeedVectorGetArray(impl->evecs[i+impl->numein], CEED_MEM_HOST, 2694a2e7687Sjeremylt &impl->edata[i + numinputfields]); CeedChk(ierr); 2704a2e7687Sjeremylt } 2714a2e7687Sjeremylt 2724a2e7687Sjeremylt // Loop through elements 2734a2e7687Sjeremylt for (CeedInt e=0; e<nblks*blksize; e+=blksize) { 2744a2e7687Sjeremylt // Input basis apply if needed 2754a2e7687Sjeremylt for (CeedInt i=0; i<numinputfields; i++) { 2764a2e7687Sjeremylt // Get elemsize, emode, ncomp 277d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetElemRestriction(opinputfields[i], &Erestrict); 278d1bcdac9Sjeremylt CeedChk(ierr); 279d1bcdac9Sjeremylt ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize); 280d1bcdac9Sjeremylt CeedChk(ierr); 281d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode); 282d1bcdac9Sjeremylt CeedChk(ierr); 283d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qfinputfields[i], &ncomp); 284d1bcdac9Sjeremylt CeedChk(ierr); 2854a2e7687Sjeremylt // Basis action 2864a2e7687Sjeremylt switch(emode) { 2874a2e7687Sjeremylt case CEED_EVAL_NONE: 288aedaa0e5Sjeremylt ierr = CeedVectorSetArray(impl->qvecsin[i], CEED_MEM_HOST, 289aedaa0e5Sjeremylt CEED_USE_POINTER, 290aedaa0e5Sjeremylt &impl->edata[i][e*Q*ncomp]); CeedChk(ierr); 2914a2e7687Sjeremylt break; 2924a2e7687Sjeremylt case CEED_EVAL_INTERP: 293aedaa0e5Sjeremylt ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr); 29491703d3fSjeremylt ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST, 295aedaa0e5Sjeremylt CEED_USE_POINTER, 296aedaa0e5Sjeremylt &impl->edata[i][e*elemsize*ncomp]); 2974a2e7687Sjeremylt CeedChk(ierr); 298d1bcdac9Sjeremylt ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE, 29991703d3fSjeremylt CEED_EVAL_INTERP, impl->evecsin[i], 300aedaa0e5Sjeremylt impl->qvecsin[i]); CeedChk(ierr); 3014a2e7687Sjeremylt break; 3024a2e7687Sjeremylt case CEED_EVAL_GRAD: 303aedaa0e5Sjeremylt ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr); 30491703d3fSjeremylt ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST, 305aedaa0e5Sjeremylt CEED_USE_POINTER, 306aedaa0e5Sjeremylt &impl->edata[i][e*elemsize*ncomp]); 3074a2e7687Sjeremylt CeedChk(ierr); 308d1bcdac9Sjeremylt ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE, 30991703d3fSjeremylt CEED_EVAL_GRAD, impl->evecsin[i], 310aedaa0e5Sjeremylt impl->qvecsin[i]); CeedChk(ierr); 3114a2e7687Sjeremylt break; 3124a2e7687Sjeremylt case CEED_EVAL_WEIGHT: 3134a2e7687Sjeremylt break; // No action 3144a2e7687Sjeremylt case CEED_EVAL_DIV: 3154a2e7687Sjeremylt break; // Not implimented 3164a2e7687Sjeremylt case CEED_EVAL_CURL: 3174a2e7687Sjeremylt break; // Not implimented 3184a2e7687Sjeremylt } 3194a2e7687Sjeremylt } 3204a2e7687Sjeremylt 3214a2e7687Sjeremylt // Output pointers 3224a2e7687Sjeremylt for (CeedInt i=0; i<numoutputfields; i++) { 323d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode); 324d1bcdac9Sjeremylt CeedChk(ierr); 3254a2e7687Sjeremylt if (emode == CEED_EVAL_NONE) { 326d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qfoutputfields[i], &ncomp); 327d1bcdac9Sjeremylt CeedChk(ierr); 328aedaa0e5Sjeremylt ierr = CeedVectorSetArray(impl->qvecsout[i], CEED_MEM_HOST, 329aedaa0e5Sjeremylt CEED_USE_POINTER, 330aedaa0e5Sjeremylt &impl->edata[i + numinputfields][e*Q*ncomp]); 331aedaa0e5Sjeremylt CeedChk(ierr); 3324a2e7687Sjeremylt } 3334a2e7687Sjeremylt } 3344a2e7687Sjeremylt // Q function 335aedaa0e5Sjeremylt ierr = CeedQFunctionApply(qf, Q*blksize, impl->qvecsin, impl->qvecsout); 336aedaa0e5Sjeremylt CeedChk(ierr); 3374a2e7687Sjeremylt 3384a2e7687Sjeremylt // Output basis apply if needed 3394a2e7687Sjeremylt for (CeedInt i=0; i<numoutputfields; i++) { 3404a2e7687Sjeremylt // Get elemsize, emode, ncomp 341d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetElemRestriction(opoutputfields[i], &Erestrict); 342d1bcdac9Sjeremylt CeedChk(ierr); 343d1bcdac9Sjeremylt ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize); 344d1bcdac9Sjeremylt CeedChk(ierr); 345d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode); 346d1bcdac9Sjeremylt CeedChk(ierr); 347d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetNumComponents(qfoutputfields[i], &ncomp); 348d1bcdac9Sjeremylt CeedChk(ierr); 3494a2e7687Sjeremylt // Basis action 3504a2e7687Sjeremylt switch(emode) { 3514a2e7687Sjeremylt case CEED_EVAL_NONE: 3524a2e7687Sjeremylt break; // No action 3534a2e7687Sjeremylt case CEED_EVAL_INTERP: 354d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis); 355d1bcdac9Sjeremylt CeedChk(ierr); 35691703d3fSjeremylt ierr = CeedVectorSetArray(impl->evecsout[i], CEED_MEM_HOST, 357aedaa0e5Sjeremylt CEED_USE_POINTER, 3584a2e7687Sjeremylt &impl->edata[i + numinputfields][e*elemsize*ncomp]); 359*de686571SJeremy L Thompson CeedChk(ierr); 360aedaa0e5Sjeremylt ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE, 361aedaa0e5Sjeremylt CEED_EVAL_INTERP, impl->qvecsout[i], 36291703d3fSjeremylt impl->evecsout[i]); CeedChk(ierr); 3634a2e7687Sjeremylt break; 3644a2e7687Sjeremylt case CEED_EVAL_GRAD: 365d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis); 366d1bcdac9Sjeremylt CeedChk(ierr); 36791703d3fSjeremylt ierr = CeedVectorSetArray(impl->evecsout[i], CEED_MEM_HOST, 368aedaa0e5Sjeremylt CEED_USE_POINTER, 369aedaa0e5Sjeremylt &impl->edata[i + numinputfields][e*elemsize*ncomp]); 370*de686571SJeremy L Thompson CeedChk(ierr); 371d1bcdac9Sjeremylt ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE, 372aedaa0e5Sjeremylt CEED_EVAL_GRAD, impl->qvecsout[i], 37391703d3fSjeremylt impl->evecsout[i]); CeedChk(ierr); 3744a2e7687Sjeremylt break; 3754ce2993fSjeremylt case CEED_EVAL_WEIGHT: { 3764ce2993fSjeremylt Ceed ceed; 3774ce2993fSjeremylt ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr); 3784ce2993fSjeremylt return CeedError(ceed, 1, 3794a2e7687Sjeremylt "CEED_EVAL_WEIGHT cannot be an output evaluation mode"); 3804a2e7687Sjeremylt break; // Should not occur 3814ce2993fSjeremylt } 3824a2e7687Sjeremylt case CEED_EVAL_DIV: 3834a2e7687Sjeremylt break; // Not implimented 3844a2e7687Sjeremylt case CEED_EVAL_CURL: 3854a2e7687Sjeremylt break; // Not implimented 3864a2e7687Sjeremylt } 3874a2e7687Sjeremylt } 3884a2e7687Sjeremylt } 3894a2e7687Sjeremylt 3904a2e7687Sjeremylt // Zero lvecs 391d1bcdac9Sjeremylt for (CeedInt i=0; i<numoutputfields; i++) { 392d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr); 39352d6035fSJeremy L Thompson if (vec == CEED_VECTOR_ACTIVE) { 39452d6035fSJeremy L Thompson if (!impl->add) { 395d1bcdac9Sjeremylt vec = outvec; 396d1bcdac9Sjeremylt ierr = CeedVectorSetValue(vec, 0.0); CeedChk(ierr); 3974a2e7687Sjeremylt } 39852d6035fSJeremy L Thompson } else { 39952d6035fSJeremy L Thompson ierr = CeedVectorSetValue(vec, 0.0); CeedChk(ierr); 40052d6035fSJeremy L Thompson } 40152d6035fSJeremy L Thompson } 40252d6035fSJeremy L Thompson impl->add = false; 4034a2e7687Sjeremylt 4044a2e7687Sjeremylt // Output restriction 4054a2e7687Sjeremylt for (CeedInt i=0; i<numoutputfields; i++) { 4064a2e7687Sjeremylt // Restore evec 4074a2e7687Sjeremylt ierr = CeedVectorRestoreArray(impl->evecs[i+impl->numein], 4084a2e7687Sjeremylt &impl->edata[i + numinputfields]); CeedChk(ierr); 409d1bcdac9Sjeremylt // Get output vector 410d1bcdac9Sjeremylt ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr); 4114a2e7687Sjeremylt // Active 412d1bcdac9Sjeremylt if (vec == CEED_VECTOR_ACTIVE) 413d1bcdac9Sjeremylt vec = outvec; 4144a2e7687Sjeremylt // Restrict 4154dccadb6Sjeremylt ierr = CeedOperatorFieldGetLMode(opoutputfields[i], &lmode); CeedChk(ierr); 4164a2e7687Sjeremylt ierr = CeedElemRestrictionApply(impl->blkrestr[i+impl->numein], CEED_TRANSPOSE, 417d1bcdac9Sjeremylt lmode, impl->evecs[i+impl->numein], vec, 4184a2e7687Sjeremylt request); CeedChk(ierr); 419d1bcdac9Sjeremylt 4204a2e7687Sjeremylt } 4214a2e7687Sjeremylt 4224a2e7687Sjeremylt // Restore input arrays 4234a2e7687Sjeremylt for (CeedInt i=0; i<numinputfields; i++) { 424d1bcdac9Sjeremylt ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode); 425d1bcdac9Sjeremylt CeedChk(ierr); 4264a2e7687Sjeremylt if (emode == CEED_EVAL_WEIGHT) { // Skip 4274a2e7687Sjeremylt } else { 4284a2e7687Sjeremylt ierr = CeedVectorRestoreArrayRead(impl->evecs[i], 429d1bcdac9Sjeremylt (const CeedScalar **) &impl->edata[i]); 430d1bcdac9Sjeremylt CeedChk(ierr); 4314a2e7687Sjeremylt } 4324a2e7687Sjeremylt } 4334a2e7687Sjeremylt 4344a2e7687Sjeremylt return 0; 4354a2e7687Sjeremylt } 4364a2e7687Sjeremylt 4374a2e7687Sjeremylt int CeedOperatorCreate_Blocked(CeedOperator op) { 4384a2e7687Sjeremylt int ierr; 439fe2413ffSjeremylt Ceed ceed; 440fe2413ffSjeremylt ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr); 4414ce2993fSjeremylt CeedOperator_Blocked *impl; 4424a2e7687Sjeremylt 4434a2e7687Sjeremylt ierr = CeedCalloc(1, &impl); CeedChk(ierr); 444*de686571SJeremy L Thompson ierr = CeedOperatorSetData(op, (void *)&impl); CeedChk(ierr); 445fe2413ffSjeremylt 446fe2413ffSjeremylt ierr = CeedSetBackendFunction(ceed, "Operator", op, "Apply", 447fe2413ffSjeremylt CeedOperatorApply_Blocked); CeedChk(ierr); 448fe2413ffSjeremylt ierr = CeedSetBackendFunction(ceed, "Operator", op, "Destroy", 449fe2413ffSjeremylt CeedOperatorDestroy_Blocked); CeedChk(ierr); 4504a2e7687Sjeremylt return 0; 4514a2e7687Sjeremylt } 452